{"source":"BenchGecko","url":"https://benchgecko.ai/model/grok-3-beta","as_of":"2026-05-03","license":"BenchGecko collected data (prices, provider offers, Gecko Tests) is CC BY 4.0 · https://creativecommons.org/licenses/by/4.0/. Benchmark scores belong to their original publishers (see sources) and are aggregated with attribution. Attribution required: \"Source: BenchGecko\" with a link.","attribution":"Source: BenchGecko · https://benchgecko.ai/model/grok-3-beta","cite":"Grok 3 Beta · benchmarks, pricing and providers. BenchGecko, data as of 2026-05-03. https://benchgecko.ai/model/grok-3-beta","slug":"grok-3-beta","name":"Grok 3 Beta","provider":{"name":"xAI","slug":"xai"},"model_type":"text","release_date":"2025-04-09","is_open_source":false,"status":"preview","context_window":131072,"description":"Grok 3 is the latest model from xAI. It's their flagship model that excels at enterprise use cases like data extraction, coding, and text summarization. Possesses deep domain knowledge in...","benchgecko_score":{"value":67.9,"rank":47,"of":312,"method":"Normalized average of public benchmark scores","method_url":"https://benchgecko.ai/methodology"},"avg_score":67.9,"list_price":{"input_usd_per_m":3,"output_usd_per_m":15,"source":"OpenRouter models API (checked daily)","as_of":"2026-10-05"},"pricing":{"input":3,"output":15},"providers":[],"cheapest_provider":null,"providers_source":null,"scores":[{"benchmark":"HELM — IFEval","benchmark_slug":"helm-ifeval","score":88.4,"unit":"%","category":"language","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-ifeval"},{"benchmark":"HELM — WildBench","benchmark_slug":"helm-wildbench","score":84.9,"unit":"%","category":"reasoning","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-wildbench"},{"benchmark":"HELM — MMLU-Pro","benchmark_slug":"helm-mmlu-pro","score":78.8,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-mmlu-pro"},{"benchmark":"HELM — GPQA","benchmark_slug":"helm-gpqa","score":65,"unit":"%","category":"knowledge","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-gpqa"},{"benchmark":"Aider polyglot","benchmark_slug":"aider-polyglot","score":53.3,"unit":"%","category":"coding","source":"Aider leaderboards","source_url":"https://aider.chat/docs/leaderboards/","benchmark_url":"https://benchgecko.ai/benchmark/aider-polyglot"},{"benchmark":"HELM — Omni-MATH","benchmark_slug":"helm-omni-math","score":46.4,"unit":"%","category":"math","source":"Stanford HELM","source_url":"https://crfm.stanford.edu/helm/","benchmark_url":"https://benchgecko.ai/benchmark/helm-omni-math"}],"gecko_tests":null,"links":{"page":"https://benchgecko.ai/model/grok-3-beta","json":"https://benchgecko.ai/api/v1/models/grok-3-beta","price_history":"https://benchgecko.ai/api/v1/price-history/grok-3-beta","provider_prices":null}}