{"source":"BenchGecko","url":"https://benchgecko.ai/model/gpt-5-mini","as_of":"2026-10-05","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/gpt-5-mini","cite":"GPT-5 Mini · benchmarks, pricing and providers. BenchGecko, data as of 2026-10-05. https://benchgecko.ai/model/gpt-5-mini","slug":"gpt-5-mini","name":"GPT-5 Mini","provider":{"name":"OpenAI","slug":"openai"},"model_type":"multimodal","release_date":"2025-08-07","is_open_source":false,"status":"active","context_window":400000,"description":"GPT-5 Mini is a compact version of GPT-5, designed to handle lighter-weight reasoning tasks. It provides the same instruction-following and safety-tuning benefits as GPT-5, but with reduced latency and cost....","benchgecko_score":{"value":56,"rank":111,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":56,"list_price":{"input_usd_per_m":0.25,"output_usd_per_m":2,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":0.25,"output":2},"providers":[{"provider":"OpenAI","endpoint":"openai/flex","quantization":null,"context_tokens":400000,"input_usd_per_m":0.125,"output_usd_per_m":1,"cache_read_usd_per_m":0.0125,"uptime_1d_pct":100,"as_of":"2026-10-05"},{"provider":"Azure","endpoint":"azure","quantization":null,"context_tokens":400000,"input_usd_per_m":0.25,"output_usd_per_m":2,"cache_read_usd_per_m":0.03,"uptime_1d_pct":99.91,"as_of":"2026-10-05"},{"provider":"OpenAI","endpoint":"openai","quantization":null,"context_tokens":400000,"input_usd_per_m":0.25,"output_usd_per_m":2,"cache_read_usd_per_m":0.025,"uptime_1d_pct":99.98,"as_of":"2026-10-05"},{"provider":"Azure","endpoint":"azure/swedencentral","quantization":null,"context_tokens":400000,"input_usd_per_m":0.275,"output_usd_per_m":2.2,"cache_read_usd_per_m":0.033,"uptime_1d_pct":100,"as_of":"2026-10-05"}],"cheapest_provider":{"provider":"OpenAI","endpoint":"openai/flex","quantization":null,"context_tokens":400000,"input_usd_per_m":0.125,"output_usd_per_m":1,"cache_read_usd_per_m":0.0125,"uptime_1d_pct":100,"as_of":"2026-10-05"},"providers_source":"OpenRouter endpoints API, one row per provider, refreshed daily","scores":[{"benchmark":"MATH level 5","benchmark_slug":"math-level-5","score":97.85,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/math-level-5"},{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":92.7,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"OTIS Mock AIME 2024-2025","benchmark_slug":"otis-mock-aime-2024-2025","score":86.65,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/otis-mock-aime-2024-2025"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":85.5,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":83.5,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"Lech Mazur Writing","benchmark_slug":"lech-mazur-writing","score":83.1,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lech-mazur-writing"},{"benchmark":"LiveBench — Coding","benchmark_slug":"livebench-coding","score":76.07,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-coding"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":75.6,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"LiveBench — Mathematics","benchmark_slug":"livebench-mathematics","score":74.38,"unit":"%","category":"math","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-mathematics"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":72.2,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"},{"benchmark":"Fiction.LiveBench","benchmark_slug":"fiction-livebench","score":69.4,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/fiction-livebench"},{"benchmark":"LiveBench — Language","benchmark_slug":"livebench-language","score":69.15,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-language"},{"benchmark":"Dtbench","benchmark_slug":"dtbench","score":67.55,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/dtbench"},{"benchmark":"GPQA diamond","benchmark_slug":"gpqa-diamond","score":66.67,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/gpqa-diamond"},{"benchmark":"SWE-Bench verified","benchmark_slug":"swe-bench-verified","score":64.67,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/swe-bench-verified"},{"benchmark":"LiveBench — If","benchmark_slug":"livebench-if","score":64.22,"unit":"%","category":"language","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-if"},{"benchmark":"LiveBench — Overall","benchmark_slug":"livebench-overall","score":61.01,"unit":"%","category":"knowledge","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-overall"},{"benchmark":"SWE-Bench Verified (Bash Only)","benchmark_slug":"swe-bench-verified-bash-only","score":59.8,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/swe-bench-verified-bash-only"},{"benchmark":"LiveBench — Reasoning","benchmark_slug":"livebench-reasoning","score":58.65,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-reasoning"},{"benchmark":"ARC-AGI","benchmark_slug":"arc-agi","score":54.33,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi"},{"benchmark":"WeirdML","benchmark_slug":"weirdml","score":52.67,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/weirdml"},{"benchmark":"LiveBench — Data Analysis","benchmark_slug":"livebench-data-analysis","score":49.61,"unit":"%","category":"reasoning","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-data-analysis"},{"benchmark":"FrontierMath-Tiers-1-3-v2-Private","benchmark_slug":"frontiermath-tiers-1-3-v2-private","score":46.67,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tiers-1-3-v2-private"},{"benchmark":"Lmca","benchmark_slug":"lmca","score":40.24,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/lmca"},{"benchmark":"LiveBench — Agentic Coding","benchmark_slug":"livebench-agentic-coding","score":35,"unit":"%","category":"coding","source":"LiveBench","source_url":"https://livebench.ai","benchmark_url":"https://benchgecko.ai/benchmark/livebench-agentic-coding"},{"benchmark":"Terminal Bench","benchmark_slug":"terminal-bench","score":34.83,"unit":"%","category":"coding","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/terminal-bench"},{"benchmark":"FrontierMath-2025-02-28-Private","benchmark_slug":"frontiermath-2025-02-28-private","score":27.24,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-2025-02-28-private"},{"benchmark":"Chess Puzzles","benchmark_slug":"chess-puzzles","score":26.35,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/chess-puzzles"},{"benchmark":"SimpleQA Verified","benchmark_slug":"simpleqa-verified","score":21,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/simpleqa-verified"},{"benchmark":"HLE","benchmark_slug":"hle","score":15.38,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/hle"},{"benchmark":"FrontierMath-Tier-4-v2-Private","benchmark_slug":"frontiermath-tier-4-v2-private","score":12.2,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-v2-private"},{"benchmark":"VPCT","benchmark_slug":"vpct","score":10.3,"unit":"%","category":"knowledge","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/vpct"},{"benchmark":"Proofbench","benchmark_slug":"proofbench","score":9,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/proofbench"},{"benchmark":"FrontierMath-Tier-4-2025-07-01-Private","benchmark_slug":"frontiermath-tier-4-2025-07-01-private","score":6.25,"unit":"%","category":"math","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/frontiermath-tier-4-2025-07-01-private"},{"benchmark":"ARC-AGI-2","benchmark_slug":"arc-agi-2","score":4.44,"unit":"%","category":"reasoning","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/arc-agi-2"},{"benchmark":"Mystery Game Puzzles","benchmark_slug":"mystery-game-puzzles","score":0.86,"unit":"%","category":"general","source":"Epoch AI Benchmarking Hub","source_url":"https://epoch.ai/data/ai-benchmarking-dashboard","benchmark_url":"https://benchgecko.ai/benchmark/mystery-game-puzzles"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/gpt-5-mini","json":"https://benchgecko.ai/api/v1/models/gpt-5-mini","price_history":"https://benchgecko.ai/api/v1/price-history/gpt-5-mini","provider_prices":"https://benchgecko.ai/pricing/arbitrage/gpt-5-mini"}}