{ "schemaVersion": "1", "observedAt": "2026-08-23T12:18:43.630Z", "currency": "USD", "unitTokens": 1000000, "providers": [ { "id": "openai", "label": "OpenAI" }, { "id": "anthropic", "label": "Anthropic" }, { "id": "google", "label": "Google" }, { "id": "xai", "label": "xAI" }, { "id": "deepseek", "label": "DeepSeek" }, { "id": "kimi", "label": "Kimi / Moonshot AI" }, { "id": "qwen", "label": "Qwen / Alibaba Cloud" }, { "id": "mistral", "label": "Mistral AI" }, { "id": "cohere", "label": "Cohere" } ], "models": [ { "id": "openai:gpt-5.6-sol:standard:short", "modelId": "gpt-5.6-sol", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Sol", "tierLabel": "Standard · up to 272K input", "inputPerMillionUsd": 5, "cachedInputPerMillionUsd": 0.5, "cacheWritePerMillionUsd": 6.25, "outputPerMillionUsd": 30, "contextWindowTokens": 1050000, "maxInputTokensInclusive": 272000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "openai:gpt-5.6-sol:standard:long", "modelId": "gpt-5.6-sol", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Sol", "tierLabel": "Standard · over 272K input", "inputPerMillionUsd": 10, "cachedInputPerMillionUsd": 1, "cacheWritePerMillionUsd": 12.5, "outputPerMillionUsd": 45, "contextWindowTokens": 1050000, "minInputTokensExclusive": 272000, "maxInputTokensInclusive": 1050000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "openai:gpt-5.6-terra:standard:short", "modelId": "gpt-5.6-terra", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Terra", "tierLabel": "Standard · up to 272K input", "inputPerMillionUsd": 2, "cachedInputPerMillionUsd": 0.2, "cacheWritePerMillionUsd": 2.5, "outputPerMillionUsd": 12, "contextWindowTokens": 1050000, "maxInputTokensInclusive": 272000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "openai:gpt-5.6-terra:standard:long", "modelId": "gpt-5.6-terra", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Terra", "tierLabel": "Standard · over 272K input", "inputPerMillionUsd": 4, "cachedInputPerMillionUsd": 0.4, "cacheWritePerMillionUsd": 5, "outputPerMillionUsd": 18, "contextWindowTokens": 1050000, "minInputTokensExclusive": 272000, "maxInputTokensInclusive": 1050000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "openai:gpt-5.6-luna:standard:short", "modelId": "gpt-5.6-luna", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Luna", "tierLabel": "Standard · up to 272K input", "inputPerMillionUsd": 0.2, "cachedInputPerMillionUsd": 0.02, "cacheWritePerMillionUsd": 0.25, "outputPerMillionUsd": 1.2, "contextWindowTokens": 1050000, "maxInputTokensInclusive": 272000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "openai:gpt-5.6-luna:standard:long", "modelId": "gpt-5.6-luna", "provider": "openai", "providerLabel": "OpenAI", "label": "GPT-5.6 Luna", "tierLabel": "Standard · over 272K input", "inputPerMillionUsd": 0.4, "cachedInputPerMillionUsd": 0.04, "cacheWritePerMillionUsd": 0.5, "outputPerMillionUsd": 1.8, "contextWindowTokens": 1050000, "minInputTokensExclusive": 272000, "maxInputTokensInclusive": 1050000, "region": "Global", "sourceUrl": "https://developers.openai.com/api/docs/pricing", "sourceLabel": "OpenAI API pricing" }, { "id": "anthropic:claude-fable-5:standard", "modelId": "claude-fable-5", "provider": "anthropic", "providerLabel": "Anthropic", "label": "Claude Fable 5", "tierLabel": "Standard", "inputPerMillionUsd": 10, "cachedInputPerMillionUsd": 1, "cacheWritePerMillionUsd": 12.5, "cacheWriteOneHourPerMillionUsd": 20, "outputPerMillionUsd": 50, "contextWindowTokens": 1000000, "region": "Global routing", "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing", "sourceLabel": "Claude API pricing" }, { "id": "anthropic:claude-opus-5:standard", "modelId": "claude-opus-5", "provider": "anthropic", "providerLabel": "Anthropic", "label": "Claude Opus 5", "tierLabel": "Standard", "inputPerMillionUsd": 5, "cachedInputPerMillionUsd": 0.5, "cacheWritePerMillionUsd": 6.25, "cacheWriteOneHourPerMillionUsd": 10, "outputPerMillionUsd": 25, "contextWindowTokens": 1000000, "region": "Global routing", "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing", "sourceLabel": "Claude API pricing" }, { "id": "anthropic:claude-sonnet-5:standard", "modelId": "claude-sonnet-5", "provider": "anthropic", "providerLabel": "Anthropic", "label": "Claude Sonnet 5", "tierLabel": "Standard", "inputPerMillionUsd": 2, "cachedInputPerMillionUsd": 0.2, "cacheWritePerMillionUsd": 2.5, "cacheWriteOneHourPerMillionUsd": 4, "outputPerMillionUsd": 10, "contextWindowTokens": 1000000, "region": "Global routing", "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing", "sourceLabel": "Claude API pricing" }, { "id": "anthropic:claude-sonnet-4.6:standard", "modelId": "claude-sonnet-4-6", "provider": "anthropic", "providerLabel": "Anthropic", "label": "Claude Sonnet 4.6", "tierLabel": "Standard", "inputPerMillionUsd": 3, "cachedInputPerMillionUsd": 0.3, "cacheWritePerMillionUsd": 3.75, "cacheWriteOneHourPerMillionUsd": 6, "outputPerMillionUsd": 15, "contextWindowTokens": 1000000, "region": "Global routing", "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing", "sourceLabel": "Claude API pricing" }, { "id": "anthropic:claude-haiku-4.5:standard", "modelId": "claude-haiku-4-5", "provider": "anthropic", "providerLabel": "Anthropic", "label": "Claude Haiku 4.5", "tierLabel": "Standard", "inputPerMillionUsd": 1, "cachedInputPerMillionUsd": 0.1, "cacheWritePerMillionUsd": 1.25, "cacheWriteOneHourPerMillionUsd": 2, "outputPerMillionUsd": 5, "contextWindowTokens": 200000, "region": "Global routing", "sourceUrl": "https://platform.claude.com/docs/en/about-claude/pricing", "sourceLabel": "Claude API pricing" }, { "id": "google:gemini-3.7-flash:standard:intro", "modelId": "gemini-3.7-flash", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.7 Flash", "tierLabel": "Developer API · Standard · introductory", "inputPerMillionUsd": 0.75, "cachedInputPerMillionUsd": 0.075, "explicitCacheStoragePerMillionTokenHourUsd": 0.5, "outputPerMillionUsd": 3.75, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.7-flash", "sourceLabel": "Gemini API pricing", "effectiveUntil": "2027-01-01T00:00:00Z" }, { "id": "google:gemini-3.6-flash:standard:intro", "modelId": "gemini-3.6-flash", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.6 Flash", "tierLabel": "Developer API · Standard · introductory", "inputPerMillionUsd": 0.75, "cachedInputPerMillionUsd": 0.075, "explicitCacheStoragePerMillionTokenHourUsd": 0.5, "outputPerMillionUsd": 3.75, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.6-flash", "sourceLabel": "Gemini API pricing", "effectiveUntil": "2027-01-01T00:00:00Z" }, { "id": "google:gemini-3.5-flash:standard", "modelId": "gemini-3.5-flash", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.5 Flash", "tierLabel": "Developer API · Standard", "inputPerMillionUsd": 1.5, "cachedInputPerMillionUsd": 0.15, "explicitCacheStoragePerMillionTokenHourUsd": 1, "outputPerMillionUsd": 9, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.5-flash", "sourceLabel": "Gemini API pricing" }, { "id": "google:gemini-3.5-flash-lite:standard", "modelId": "gemini-3.5-flash-lite", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.5 Flash-Lite", "tierLabel": "Developer API · Standard", "inputPerMillionUsd": 0.3, "cachedInputPerMillionUsd": 0.03, "explicitCacheStoragePerMillionTokenHourUsd": 1, "outputPerMillionUsd": 2.5, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.5-flash-lite", "sourceLabel": "Gemini API pricing" }, { "id": "google:gemini-3.1-flash-lite:standard", "modelId": "gemini-3.1-flash-lite", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.1 Flash-Lite", "tierLabel": "Developer API · Standard · text/image/video", "inputPerMillionUsd": 0.25, "cachedInputPerMillionUsd": 0.025, "explicitCacheStoragePerMillionTokenHourUsd": 1, "outputPerMillionUsd": 1.5, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.1-flash-lite", "sourceLabel": "Gemini API pricing" }, { "id": "google:gemini-3.1-pro-preview:standard:short", "modelId": "gemini-3.1-pro-preview", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.1 Pro Preview", "tierLabel": "Developer API · Standard · up to 200K prompt", "inputPerMillionUsd": 2, "cachedInputPerMillionUsd": 0.2, "explicitCacheStoragePerMillionTokenHourUsd": 4.5, "outputPerMillionUsd": 12, "maxInputTokensInclusive": 200000, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.1-pro-preview", "sourceLabel": "Gemini API pricing" }, { "id": "google:gemini-3.1-pro-preview:standard:long", "modelId": "gemini-3.1-pro-preview", "provider": "google", "providerLabel": "Google", "label": "Gemini 3.1 Pro Preview", "tierLabel": "Developer API · Standard · over 200K prompt", "inputPerMillionUsd": 4, "cachedInputPerMillionUsd": 0.4, "explicitCacheStoragePerMillionTokenHourUsd": 4.5, "outputPerMillionUsd": 18, "minInputTokensExclusive": 200000, "region": "Gemini Developer API", "sourceUrl": "https://ai.google.dev/gemini-api/docs/pricing#gemini-3.1-pro-preview", "sourceLabel": "Gemini API pricing" }, { "id": "xai:grok-4.6:standard:short", "modelId": "grok-4.6", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.6", "tierLabel": "Standard · below 200K prompt", "inputPerMillionUsd": 2, "cachedInputPerMillionUsd": 0.5, "outputPerMillionUsd": 6, "contextWindowTokens": 500000, "maxInputTokensExclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-4.6:standard:long", "modelId": "grok-4.6", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.6", "tierLabel": "Standard · 200K+ prompt", "inputPerMillionUsd": 4, "cachedInputPerMillionUsd": 1, "outputPerMillionUsd": 12, "contextWindowTokens": 500000, "minInputTokensInclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-4.5:standard:short", "modelId": "grok-4.5", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.5", "tierLabel": "Standard · below 200K prompt", "inputPerMillionUsd": 2, "cachedInputPerMillionUsd": 0.3, "outputPerMillionUsd": 6, "contextWindowTokens": 500000, "maxInputTokensExclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-4.3:standard:short", "modelId": "grok-4.3", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.3", "tierLabel": "Standard · below 200K prompt", "inputPerMillionUsd": 1.25, "cachedInputPerMillionUsd": 0.2, "outputPerMillionUsd": 2.5, "contextWindowTokens": 1000000, "maxInputTokensExclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-4.3:standard:long", "modelId": "grok-4.3", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.3", "tierLabel": "Standard · 200K+ prompt", "inputPerMillionUsd": 2.5, "cachedInputPerMillionUsd": 0.4, "outputPerMillionUsd": 5, "contextWindowTokens": 1000000, "minInputTokensInclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-build-0.1:standard:short", "modelId": "grok-build-0.1", "provider": "xai", "providerLabel": "xAI", "label": "Grok Build 0.1", "tierLabel": "Standard · below 200K prompt", "inputPerMillionUsd": 1, "cachedInputPerMillionUsd": 0.2, "outputPerMillionUsd": 2, "contextWindowTokens": 256000, "maxInputTokensExclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "xai:grok-4.20-0309-reasoning:standard:short", "modelId": "grok-4.20-0309-reasoning", "provider": "xai", "providerLabel": "xAI", "label": "Grok 4.20 Reasoning", "tierLabel": "Standard · below 200K prompt", "inputPerMillionUsd": 1.25, "cachedInputPerMillionUsd": 0.2, "outputPerMillionUsd": 2.5, "contextWindowTokens": 1000000, "maxInputTokensExclusive": 200000, "region": "xAI direct API", "sourceUrl": "https://docs.x.ai/developers/models", "sourceLabel": "xAI model pricing" }, { "id": "deepseek:deepseek-v4-flash:standard", "modelId": "deepseek-v4-flash", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash", "tierLabel": "Standard · through 2026-08-16 16:00 UTC", "inputPerMillionUsd": 0.14, "cachedInputPerMillionUsd": 0.0028, "outputPerMillionUsd": 0.28, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveUntil": "2026-08-16T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash:off-peak", "modelId": "deepseek-v4-flash", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash", "tierLabel": "Off-peak · scheduled UTC windows", "inputPerMillionUsd": 0.22, "cachedInputPerMillionUsd": 0.007, "outputPerMillionUsd": 0.66, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-16T16:00:00Z", "effectiveUntil": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash:peak", "modelId": "deepseek-v4-flash", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash", "tierLabel": "Peak · 01:00–04:00 and 06:00–10:00 UTC", "inputPerMillionUsd": 0.44, "cachedInputPerMillionUsd": 0.014, "outputPerMillionUsd": 1.32, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-16T16:00:00Z", "effectiveUntil": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-pro:standard", "modelId": "deepseek-v4-pro", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Pro", "tierLabel": "Standard · through 2026-08-16 16:00 UTC", "inputPerMillionUsd": 0.435, "cachedInputPerMillionUsd": 0.003625, "outputPerMillionUsd": 0.87, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveUntil": "2026-08-16T16:00:00Z" }, { "id": "deepseek:deepseek-v4-pro:off-peak", "modelId": "deepseek-v4-pro", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Pro", "tierLabel": "Off-peak · scheduled UTC windows", "inputPerMillionUsd": 0.66, "cachedInputPerMillionUsd": 0.022, "outputPerMillionUsd": 1.98, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-16T16:00:00Z", "effectiveUntil": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-pro:peak", "modelId": "deepseek-v4-pro", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Pro", "tierLabel": "Peak · 01:00–04:00 and 06:00–10:00 UTC", "inputPerMillionUsd": 1.32, "cachedInputPerMillionUsd": 0.044, "outputPerMillionUsd": 3.96, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-16T16:00:00Z", "effectiveUntil": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash:off-peak-weekend", "modelId": "deepseek-v4-flash", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash", "tierLabel": "Off-peak · weekdays outside peak hours; all weekend (Beijing time)", "inputPerMillionUsd": 0.22, "cachedInputPerMillionUsd": 0.007, "outputPerMillionUsd": 0.66, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash:peak-weekday", "modelId": "deepseek-v4-flash", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash", "tierLabel": "Peak · weekdays 01:00–04:00 and 06:00–10:00 UTC", "inputPerMillionUsd": 0.44, "cachedInputPerMillionUsd": 0.014, "outputPerMillionUsd": 1.32, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-pro:off-peak-weekend", "modelId": "deepseek-v4-pro", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Pro", "tierLabel": "Off-peak · weekdays outside peak hours; all weekend (Beijing time)", "inputPerMillionUsd": 0.66, "cachedInputPerMillionUsd": 0.022, "outputPerMillionUsd": 1.98, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-pro:peak-weekday", "modelId": "deepseek-v4-pro", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Pro", "tierLabel": "Peak · weekdays 01:00–04:00 and 06:00–10:00 UTC", "inputPerMillionUsd": 1.32, "cachedInputPerMillionUsd": 0.044, "outputPerMillionUsd": 3.96, "contextWindowTokens": 1000000, "region": "Global", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash-vision-exp:off-peak-weekend", "modelId": "deepseek-v4-flash-vision-exp", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash Vision Experimental", "tierLabel": "Off-peak · weekdays outside peak hours; all weekend (Beijing time)", "inputPerMillionUsd": 0.22, "cachedInputPerMillionUsd": 0.007, "outputPerMillionUsd": 0.66, "contextWindowTokens": 1000000, "region": "Global · images converted to input tokens", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "deepseek:deepseek-v4-flash-vision-exp:peak-weekday", "modelId": "deepseek-v4-flash-vision-exp", "provider": "deepseek", "providerLabel": "DeepSeek", "label": "DeepSeek V4 Flash Vision Experimental", "tierLabel": "Peak · weekdays 01:00–04:00 and 06:00–10:00 UTC", "inputPerMillionUsd": 0.44, "cachedInputPerMillionUsd": 0.014, "outputPerMillionUsd": 1.32, "contextWindowTokens": 1000000, "region": "Global · images converted to input tokens", "sourceUrl": "https://api-docs.deepseek.com/quick_start/pricing/", "sourceLabel": "DeepSeek API pricing", "effectiveFrom": "2026-08-22T16:00:00Z" }, { "id": "kimi:kimi-k3:realtime", "modelId": "kimi-k3", "provider": "kimi", "providerLabel": "Kimi / Moonshot AI", "label": "Kimi K3", "tierLabel": "Realtime", "inputPerMillionUsd": 3, "cachedInputPerMillionUsd": 0.3, "outputPerMillionUsd": 15, "contextWindowTokens": 1048576, "region": "Global", "sourceUrl": "https://platform.kimi.ai/docs/pricing/chat-k3", "sourceLabel": "Kimi K3 pricing" }, { "id": "kimi:kimi-k2.7-code:realtime", "modelId": "kimi-k2.7-code", "provider": "kimi", "providerLabel": "Kimi / Moonshot AI", "label": "Kimi K2.7 Code", "tierLabel": "Realtime", "inputPerMillionUsd": 0.95, "cachedInputPerMillionUsd": 0.19, "outputPerMillionUsd": 4, "contextWindowTokens": 262144, "region": "Global", "sourceUrl": "https://platform.kimi.ai/docs/pricing/chat-k27-code", "sourceLabel": "Kimi K2.7 Code pricing" }, { "id": "kimi:kimi-k2.7-code-highspeed:realtime", "modelId": "kimi-k2.7-code-highspeed", "provider": "kimi", "providerLabel": "Kimi / Moonshot AI", "label": "Kimi K2.7 Code Highspeed", "tierLabel": "Realtime · high-speed tier", "inputPerMillionUsd": 1.9, "cachedInputPerMillionUsd": 0.38, "outputPerMillionUsd": 8, "contextWindowTokens": 262144, "region": "Global", "sourceUrl": "https://platform.kimi.ai/docs/pricing/chat-k27-code", "sourceLabel": "Kimi K2.7 Code pricing" }, { "id": "kimi:kimi-k2.6:realtime", "modelId": "kimi-k2.6", "provider": "kimi", "providerLabel": "Kimi / Moonshot AI", "label": "Kimi K2.6", "tierLabel": "Realtime", "inputPerMillionUsd": 0.95, "cachedInputPerMillionUsd": 0.16, "outputPerMillionUsd": 4, "contextWindowTokens": 262144, "region": "Global", "sourceUrl": "https://platform.kimi.ai/docs/pricing/chat-k26", "sourceLabel": "Kimi K2.6 pricing" }, { "id": "qwen:qwen3.7-max:global", "modelId": "qwen3.7-max", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Max", "tierLabel": "Global scope · up to 1M input", "inputPerMillionUsd": 1.65, "cachedInputPerMillionUsd": 0.33, "cacheWritePerMillionUsd": 2.063, "explicitCacheReadPerMillionUsd": 0.165, "outputPerMillionUsd": 4.951, "contextWindowTokens": 1000000, "region": "US Virginia · Global deployment", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-max", "sourceLabel": "Alibaba Cloud Qwen 3.7 Max pricing" }, { "id": "qwen:qwen3.7-plus:global:short", "modelId": "qwen3.7-plus", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Plus", "tierLabel": "Global scope · up to 256K input", "inputPerMillionUsd": 0.276, "cachedInputPerMillionUsd": 0.056, "cacheWritePerMillionUsd": 0.344, "explicitCacheReadPerMillionUsd": 0.028, "outputPerMillionUsd": 1.101, "contextWindowTokens": 1000000, "maxInputTokensInclusive": 256000, "region": "US Virginia · Global deployment", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-plus", "sourceLabel": "Alibaba Cloud Qwen 3.7 Plus pricing" }, { "id": "qwen:qwen3.7-plus:global:long", "modelId": "qwen3.7-plus", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Plus", "tierLabel": "Global scope · 256K–1M input", "inputPerMillionUsd": 0.826, "cachedInputPerMillionUsd": 0.166, "cacheWritePerMillionUsd": 1.032, "explicitCacheReadPerMillionUsd": 0.083, "outputPerMillionUsd": 3.301, "contextWindowTokens": 1000000, "minInputTokensExclusive": 256000, "region": "US Virginia · Global deployment", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-plus", "sourceLabel": "Alibaba Cloud Qwen 3.7 Plus pricing" }, { "id": "qwen:qwen3.7-flash:beijing:short", "modelId": "qwen3.7-flash", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Flash", "tierLabel": "Beijing · up to 32K input", "inputPerMillionUsd": 0.028, "cachedInputPerMillionUsd": 0.006, "cacheWritePerMillionUsd": 0.034, "explicitCacheReadPerMillionUsd": 0.003, "outputPerMillionUsd": 0.11, "contextWindowTokens": 1000000, "maxInputTokensInclusive": 32000, "region": "China Beijing", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-flash", "sourceLabel": "Alibaba Cloud Qwen 3.7 Flash pricing" }, { "id": "qwen:qwen3.7-flash:beijing:medium", "modelId": "qwen3.7-flash", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Flash", "tierLabel": "Beijing · 32K–256K input", "inputPerMillionUsd": 0.083, "cachedInputPerMillionUsd": 0.017, "cacheWritePerMillionUsd": 0.103, "explicitCacheReadPerMillionUsd": 0.008, "outputPerMillionUsd": 0.33, "contextWindowTokens": 1000000, "minInputTokensExclusive": 32000, "maxInputTokensInclusive": 256000, "region": "China Beijing", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-flash", "sourceLabel": "Alibaba Cloud Qwen 3.7 Flash pricing" }, { "id": "qwen:qwen3.7-flash:beijing:long", "modelId": "qwen3.7-flash", "provider": "qwen", "providerLabel": "Qwen / Alibaba Cloud", "label": "Qwen 3.7 Flash", "tierLabel": "Beijing · 256K–1M input", "inputPerMillionUsd": 0.165, "cachedInputPerMillionUsd": 0.033, "cacheWritePerMillionUsd": 0.206, "explicitCacheReadPerMillionUsd": 0.017, "outputPerMillionUsd": 0.66, "contextWindowTokens": 1000000, "minInputTokensExclusive": 256000, "region": "China Beijing", "sourceUrl": "https://www.alibabacloud.com/help/en/model-studio/qwen3-7-flash", "sourceLabel": "Alibaba Cloud Qwen 3.7 Flash pricing" }, { "id": "mistral:mistral-medium-latest:standard", "modelId": "mistral-medium-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Mistral Medium 3.5", "tierLabel": "Standard", "inputPerMillionUsd": 1.5, "cachedInputPerMillionUsd": 0.15, "outputPerMillionUsd": 7.5, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:mistral-small-latest:standard", "modelId": "mistral-small-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Mistral Small 4", "tierLabel": "Standard", "inputPerMillionUsd": 0.15, "cachedInputPerMillionUsd": 0.015, "outputPerMillionUsd": 0.6, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:mistral-large-latest:standard", "modelId": "mistral-large-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Mistral Large 3", "tierLabel": "Standard", "inputPerMillionUsd": 0.5, "cachedInputPerMillionUsd": 0.05, "outputPerMillionUsd": 1.5, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:codestral-latest:standard", "modelId": "codestral-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Codestral", "tierLabel": "Standard", "inputPerMillionUsd": 0.3, "cachedInputPerMillionUsd": 0.03, "outputPerMillionUsd": 0.9, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:ministral-3b-latest:standard", "modelId": "ministral-3b-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Ministral 3 3B", "tierLabel": "Standard", "inputPerMillionUsd": 0.1, "cachedInputPerMillionUsd": 0.01, "outputPerMillionUsd": 0.1, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:ministral-8b-latest:standard", "modelId": "ministral-8b-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Ministral 3 8B", "tierLabel": "Standard", "inputPerMillionUsd": 0.15, "cachedInputPerMillionUsd": 0.015, "outputPerMillionUsd": 0.15, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "mistral:ministral-14b-latest:standard", "modelId": "ministral-14b-latest", "provider": "mistral", "providerLabel": "Mistral AI", "label": "Ministral 3 14B", "tierLabel": "Standard", "inputPerMillionUsd": 0.2, "cachedInputPerMillionUsd": 0.02, "outputPerMillionUsd": 0.2, "region": "Global", "sourceUrl": "https://mistral.ai/pricing/api/", "sourceLabel": "Mistral API pricing" }, { "id": "cohere:command-a-03-2025:standard", "modelId": "command-a-03-2025", "provider": "cohere", "providerLabel": "Cohere", "label": "Command A", "tierLabel": "Production pay-as-you-go", "inputPerMillionUsd": 2.5, "cachedInputPerMillionUsd": null, "outputPerMillionUsd": 10, "contextWindowTokens": 256000, "region": "Global", "sourceUrl": "https://docs.cohere.com/docs/command-a", "sourceLabel": "Cohere Command A pricing", "provenanceUrls": [ "https://docs.cohere.com/docs/command-a", "https://docs.cohere.com/v1/docs/models" ], "reviewStatus": "manual-review", "reviewNote": "The dedicated Command A page printed a conflicting Command A+ model ID on 2026-08-15; the Command A ID is cross-checked against Cohere's model catalogue." }, { "id": "cohere:command-r7b-12-2024:standard", "modelId": "command-r7b-12-2024", "provider": "cohere", "providerLabel": "Cohere", "label": "Command R7B", "tierLabel": "Production pay-as-you-go", "inputPerMillionUsd": 0.0375, "cachedInputPerMillionUsd": null, "outputPerMillionUsd": 0.15, "region": "Global", "sourceUrl": "https://docs.cohere.com/v2/docs/command-r7b", "sourceLabel": "Cohere Command R7B pricing" }, { "id": "cohere:command-r-08-2024:standard", "modelId": "command-r-08-2024", "provider": "cohere", "providerLabel": "Cohere", "label": "Command R", "tierLabel": "Production pay-as-you-go", "inputPerMillionUsd": 0.15, "cachedInputPerMillionUsd": null, "outputPerMillionUsd": 0.6, "region": "Global", "sourceUrl": "https://docs.cohere.com/docs/command-r", "sourceLabel": "Cohere Command R pricing" }, { "id": "cohere:command-r-plus-08-2024:standard", "modelId": "command-r-plus-08-2024", "provider": "cohere", "providerLabel": "Cohere", "label": "Command R+", "tierLabel": "Production pay-as-you-go", "inputPerMillionUsd": 2.5, "cachedInputPerMillionUsd": null, "outputPerMillionUsd": 10, "region": "Global", "sourceUrl": "https://docs.cohere.com/docs/command-r-plus", "sourceLabel": "Cohere Command R+ pricing" } ] }