{"source":"BenchGecko","url":"https://benchgecko.ai/model/mistral-large","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/mistral-large","cite":"Mistral Large · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/mistral-large","slug":"mistral-large","name":"Mistral Large","provider":{"name":"Mistral AI","slug":"mistral"},"model_type":"text","release_date":"2024-02-26","is_open_source":true,"status":"active","context_window":128000,"description":"This is Mistral AI's flagship model, Mistral Large 2 (version `mistral-large-2407`). It's a proprietary weights-available model and excels at reasoning, code, JSON, chat, and more. Read the launch announcement here....","benchgecko_score":{"value":31.3,"rank":240,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":31.3,"list_price":{"input_usd_per_m":2,"output_usd_per_m":6,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":2,"output":6},"providers":[{"provider":"Mistral","endpoint":"mistral","quantization":null,"context_tokens":128000,"input_usd_per_m":2,"output_usd_per_m":6,"cache_read_usd_per_m":0.2,"uptime_1d_pct":99.94,"as_of":"2026-10-05"},{"provider":"Mistral","endpoint":"mistral/eu","quantization":null,"context_tokens":128000,"input_usd_per_m":2.2,"output_usd_per_m":6.6,"cache_read_usd_per_m":0.22,"uptime_1d_pct":null,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"Mistral","endpoint":"mistral","quantization":null,"context_tokens":128000,"input_usd_per_m":2,"output_usd_per_m":6,"cache_read_usd_per_m":0.2,"uptime_1d_pct":99.94,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":87.6,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":80.1,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":69,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"Aider — Code Editing","benchmark_slug":"aider-edit","score":65.4,"unit":"%","category":"coding","source":"Aider leaderboards","source_url":"https://aider.chat/docs/leaderboards/","benchmark_url":"https://benchgecko.ai/benchmark/aider-edit"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":59.9,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"MMLU","benchmark_slug":"mmlu","score":58.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mmlu"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":43.5,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":28.1,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":26.5,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":24.46,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":18.35,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"SimpleBench","benchmark_slug":"simplebench","score":7,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simplebench"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":1.85,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":0.6,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/mistral-large","json":"https://benchgecko.ai/api/v1/models/mistral-large","price_history":"https://benchgecko.ai/api/v1/price-history/mistral-large","provider_prices":"https://benchgecko.ai/pricing/arbitrage/mistral-large"}}