{"source":"BenchGecko","url":"https://benchgecko.ai/model/kimi-k2-thinking","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/kimi-k2-thinking","cite":"Kimi K2 Thinking · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/kimi-k2-thinking","slug":"kimi-k2-thinking","name":"Kimi K2 Thinking","provider":{"name":"moonshotai","slug":"moonshotai"},"model_type":"text","release_date":"2025-11-06","is_open_source":true,"status":"active","context_window":262144,"description":"Kimi K2 Thinking is Moonshot AI’s most advanced open reasoning model to date, extending the K2 series into agentic, long-horizon reasoning. Built on the trillion-parameter Mixture-of-Experts (MoE) architecture introduced in...","benchgecko_score":{"value":58.3,"rank":100,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":58.3,"list_price":{"input_usd_per_m":0.6,"output_usd_per_m":2.5,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":0.6,"output":2.5},"providers":[{"provider":"Google","endpoint":"google-vertex","quantization":null,"context_tokens":262144,"input_usd_per_m":0.6,"output_usd_per_m":2.5,"cache_read_usd_per_m":null,"uptime_1d_pct":99.92,"as_of":"2026-10-05"},{"provider":"Novita","endpoint":"novita/bf16","quantization":"bf16","context_tokens":262144,"input_usd_per_m":0.6,"output_usd_per_m":2.5,"cache_read_usd_per_m":0.15,"uptime_1d_pct":99.94,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"Google","endpoint":"google-vertex","quantization":null,"context_tokens":262144,"input_usd_per_m":0.6,"output_usd_per_m":2.5,"cache_read_usd_per_m":null,"uptime_1d_pct":99.92,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"OpenCompass — AIME2025","benchmark_slug":"oc-aime2025","score":94.1,"unit":"%","category":"math","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-aime2025"},{"benchmark":"OpenCompass — IFEval","benchmark_slug":"oc-ifeval","score":92.4,"unit":"%","category":"language","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-ifeval"},{"benchmark":"OpenCompass — MMLU-Pro","benchmark_slug":"oc-mmlu-pro","score":84.3,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-mmlu-pro"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":83.04,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"OpenCompass — GPQA-Diamond","benchmark_slug":"oc-gpqa-diamond","score":82.7,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-gpqa-diamond"},{"benchmark":"LiveBench — Mathematics","benchmark_slug":"livebench-mathematics","score":81.1,"unit":"%","category":"math","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-mathematics"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":78.96,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"OpenCompass — LiveCodeBenchV6","benchmark_slug":"oc-livecodebenchv6","score":77.1,"unit":"%","category":"coding","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-livecodebenchv6"},{"benchmark":"LiveBench — Coding","benchmark_slug":"livebench-coding","score":67.44,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-coding"},{"benchmark":"LiveBench — Language","benchmark_slug":"livebench-language","score":66.45,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-language"},{"benchmark":"LiveBench — Reasoning","benchmark_slug":"livebench-reasoning","score":63.49,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-reasoning"},{"benchmark":"SWE-Bench Verified (Bash Only)","benchmark_slug":"swe-bench-verified-bash-only","score":63.4,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/swe-bench-verified-bash-only"},{"benchmark":"LiveBench — If","benchmark_slug":"livebench-if","score":62.03,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-if"},{"benchmark":"LiveBench — Overall","benchmark_slug":"livebench-overall","score":61.59,"unit":"%","category":"knowledge","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-overall"},{"benchmark":"Metr Time Horizons","benchmark_slug":"metr-time-horizons","score":59.19,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/metr-time-horizons"},{"benchmark":"LiveBench — Data Analysis","benchmark_slug":"livebench-data-analysis","score":52.29,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-data-analysis"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":42.79,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"LiveBench — Agentic Coding","benchmark_slug":"livebench-agentic-coding","score":38.33,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-agentic-coding"},{"benchmark":"Terminal Bench","benchmark_slug":"terminal-bench","score":35.7,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/terminal-bench"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":31.6,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":21.4,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"OpenCompass — HLE","benchmark_slug":"oc-hle","score":21.3,"unit":"%","category":"knowledge","source":"OpenCompass","source_url":"https://rank.opencompass.org.cn","benchmark_url":"https://benchgecko.ai/benchmark/oc-hle"},{"benchmark":"Chess Puzzles","benchmark_slug":"chess-puzzles","score":20,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/chess-puzzles"},{"benchmark":"Cl Bench","benchmark_slug":"cl-bench","score":17.6,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/cl-bench"},{"benchmark":"PostTrainBench","benchmark_slug":"posttrainbench","score":7.25,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/posttrainbench"},{"benchmark":"APEX-Agents","benchmark_slug":"apex-agents","score":4,"unit":"%","category":"agentic","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/apex-agents"},{"benchmark":"FrontierMath-Tier-4-2025-07-01-Private","benchmark_slug":"frontiermath-tier-4-2025-07-01-private","score":0,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-2025-07-01-private"}],"gecko_tests":{"gecko_score":null,"rank":null,"tests_taken":2,"tests":[{"test":"who-are-you","name":"Who Are You","grade":"B","value":100,"verdict":"Knows who made it","url":"https://benchgecko.ai/gecko-tests/who-are-you"},{"test":"tokenizer-tax","name":"Tokenizer Tax","grade":"E","value":7.7,"verdict":"92% more tokens outside English","url":"https://benchgecko.ai/gecko-tests/tokenizer-tax"}],"scorecard_url":"https://benchgecko.ai/gecko-tests/scorecard"},"links":{"page":"https://benchgecko.ai/model/kimi-k2-thinking","json":"https://benchgecko.ai/api/v1/models/kimi-k2-thinking","price_history":"https://benchgecko.ai/api/v1/price-history/kimi-k2-thinking","provider_prices":"https://benchgecko.ai/pricing/arbitrage/kimi-k2-thinking"}}