{"source":"BenchGecko","url":"https://benchgecko.ai/model/deepseek-chat","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/deepseek-chat","cite":"DeepSeek V3 · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/deepseek-chat","slug":"deepseek-chat","name":"DeepSeek V3","provider":{"name":"DeepSeek","slug":"deepseek"},"model_type":"text","release_date":"2024-12-26","is_open_source":true,"status":"active","context_window":163840,"description":"DeepSeek-V3 is the latest model from the DeepSeek team, building upon the instruction following and coding abilities of the previous versions. Pre-trained on nearly 15 trillion tokens, the reported evaluations...","benchgecko_score":{"value":55.1,"rank":117,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":55.1,"list_price":{"input_usd_per_m":0.2574,"output_usd_per_m":1.0287,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":0.2574,"output":1.0287},"providers":[{"provider":"StreamLake","endpoint":"streamlake","quantization":null,"context_tokens":128000,"input_usd_per_m":0.2574,"output_usd_per_m":1.0287,"cache_read_usd_per_m":null,"uptime_1d_pct":99.47,"as_of":"2026-10-05"},{"provider":"DeepInfra","endpoint":"deepinfra/fp4","quantization":"fp4","context_tokens":163840,"input_usd_per_m":0.32,"output_usd_per_m":0.89,"cache_read_usd_per_m":null,"uptime_1d_pct":97.46,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"StreamLake","endpoint":"streamlake","quantization":null,"context_tokens":128000,"input_usd_per_m":0.2574,"output_usd_per_m":1.0287,"cache_read_usd_per_m":null,"uptime_1d_pct":99.47,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1358.44,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"ARC AI2","benchmark_slug":"arc-ai2","score":93.73,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-ai2"},{"benchmark":"HellaSwag","benchmark_slug":"hellaswag","score":85.2,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/hellaswag"},{"benchmark":"BBH","benchmark_slug":"bbh","score":83.33,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/bbh"},{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":83.2,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":83.1,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"MMLU","benchmark_slug":"mmlu","score":82.93,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mmlu"},{"benchmark":"TriviaQA","benchmark_slug":"triviaqa","score":82.9,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/triviaqa"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":77,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":72.3,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"Winogrande","benchmark_slug":"winogrande","score":70.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/winogrande"},{"benchmark":"PIQA","benchmark_slug":"piqa","score":69.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/piqa"},{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":64.85,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":53.8,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"Fiction.LiveBench","benchmark_slug":"fiction-livebench","score":50,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/fiction-livebench"},{"benchmark":"Aider polyglot","benchmark_slug":"aider-polyglot","score":48.4,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/aider-polyglot"},{"benchmark":"Metr Time Horizons","benchmark_slug":"metr-time-horizons","score":47.36,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/metr-time-horizons"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":42.05,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":41.33,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":40.3,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":36.08,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"Lmca","benchmark_slug":"lmca","score":18.2,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lmca"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":15.75,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":3.02,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"SimpleBench","benchmark_slug":"simplebench","score":2.68,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simplebench"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/deepseek-chat","json":"https://benchgecko.ai/api/v1/models/deepseek-chat","price_history":"https://benchgecko.ai/api/v1/price-history/deepseek-chat","provider_prices":"https://benchgecko.ai/pricing/arbitrage/deepseek-chat"}}