{"source":"BenchGecko","url":"https://benchgecko.ai/model/gpt-4o-mini-2024-07-18","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/gpt-4o-mini-2024-07-18","cite":"GPT-4o-mini (2024-07-18) · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/gpt-4o-mini-2024-07-18","slug":"gpt-4o-mini-2024-07-18","name":"GPT-4o-mini (2024-07-18)","provider":{"name":"OpenAI","slug":"openai"},"model_type":"multimodal","release_date":"2024-07-18","is_open_source":false,"status":"active","context_window":128000,"description":"GPT-4o mini is OpenAI's newest model after GPT-4 Omni, supporting both text and image inputs with text outputs. As their most advanced small model, it is many multiples more affordable...","benchgecko_score":{"value":33.4,"rank":231,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":33.4,"list_price":{"input_usd_per_m":0.15,"output_usd_per_m":0.6,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":0.15,"output":0.6},"providers":[{"provider":"OpenAI","endpoint":"openai","quantization":null,"context_tokens":128000,"input_usd_per_m":0.15,"output_usd_per_m":0.6,"cache_read_usd_per_m":0.075,"uptime_1d_pct":100,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"OpenAI","endpoint":"openai","quantization":null,"context_tokens":128000,"input_usd_per_m":0.15,"output_usd_per_m":0.6,"cache_read_usd_per_m":0.075,"uptime_1d_pct":100,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"Chatbot Arena Elo — Overall","benchmark_slug":"arena-elo-overall","score":1317.65,"unit":"elo","category":"arena","source":"LMArena","source_url":"https://lmarena.ai/leaderboard","benchmark_url":"https://benchgecko.ai/benchmark/arena-elo-overall"},{"benchmark":"GSM8K","benchmark_slug":"gsm8k","score":91.3,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gsm8k"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":79.1,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":78.2,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"PIQA","benchmark_slug":"piqa","score":77.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/piqa"},{"benchmark":"MMLU","benchmark_slug":"mmlu","score":75.73,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mmlu"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":67.2,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"GeoBench","benchmark_slug":"geobench","score":64,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/geobench"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":60.3,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"VideoMME","benchmark_slug":"videomme","score":53.07,"unit":"%","category":"multimodal","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/videomme"},{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":52.63,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":36.8,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":28,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"Balrog","benchmark_slug":"balrog","score":17.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/balrog"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":16.96,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":11.76,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":6.85,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"Aider polyglot","benchmark_slug":"aider-polyglot","score":3.6,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/aider-polyglot"},{"benchmark":"VPCT","benchmark_slug":"vpct","score":1,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/vpct"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":0.1,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/gpt-4o-mini-2024-07-18","json":"https://benchgecko.ai/api/v1/models/gpt-4o-mini-2024-07-18","price_history":"https://benchgecko.ai/api/v1/price-history/gpt-4o-mini-2024-07-18","provider_prices":null}}