feat: add pricing sync script and update model prices from OpenRouter (#27)
* fix: pin npm package versions in MCP integration templates Prevent supply chain attacks by pinning exact versions instead of using unpinned `npx -y @package` which pulls latest on every run. 23 of 25 integrations pinned. sqlite-mcp and aws skipped (packages not found on npm registry). * fix: use stable azure/mcp version instead of beta * feat: add pricing sync script and update model prices from OpenRouter API - scripts/sync-pricing.py fetches real-time pricing from OpenRouter - Updated 64 price fields across 13 provider files - Run periodically or in CI to keep prices current
This commit is contained in:
42 files changed
+1388
-64
No files matched your search
@@ -0,0 +1,48 @@
|
||||
# aion-labs — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "aion-labs"
|
||||
display_name = "Aion Labs"
|
||||
api_key_env = "AION_LABS_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "aion-1.0"
|
||||
display_name = "AionLabs: Aion-1.0"
|
||||
tier = "frontier"
|
||||
context_window = 131072
|
||||
max_output_tokens = 32768
|
||||
input_cost_per_m = 4.0
|
||||
output_cost_per_m = 8.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "aion-1.0-mini"
|
||||
display_name = "AionLabs: Aion-1.0-Mini"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 32768
|
||||
input_cost_per_m = 0.7
|
||||
output_cost_per_m = 1.4
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "aion-2.0"
|
||||
display_name = "AionLabs: Aion-2.0"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 32768
|
||||
input_cost_per_m = 0.8
|
||||
output_cost_per_m = 1.6
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "aion-rp-llama-3.1-8b"
|
||||
display_name = "AionLabs: Aion-RP 1.0 (8B)"
|
||||
tier = "smart"
|
||||
context_window = 32768
|
||||
max_output_tokens = 32768
|
||||
input_cost_per_m = 0.8
|
||||
output_cost_per_m = 1.6
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# alibaba — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "alibaba"
|
||||
display_name = "Alibaba"
|
||||
api_key_env = "ALIBABA_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "tongyi-deepresearch-30b-a3b"
|
||||
display_name = "Tongyi DeepResearch 30B A3B"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.09
|
||||
output_cost_per_m = 0.45
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,48 @@
|
||||
# allenai — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "allenai"
|
||||
display_name = "Allenai"
|
||||
api_key_env = "ALLENAI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "olmo-2-0325-32b-instruct"
|
||||
display_name = "AllenAI: Olmo 2 32B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.05
|
||||
output_cost_per_m = 0.2
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "olmo-3-32b-think"
|
||||
display_name = "AllenAI: Olmo 3 32B Think"
|
||||
tier = "fast"
|
||||
context_window = 65536
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.5
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "olmo-3.1-32b-instruct"
|
||||
display_name = "AllenAI: Olmo 3.1 32B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 65536
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.2
|
||||
output_cost_per_m = 0.6
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "olmo-3.1-32b-think"
|
||||
display_name = "AllenAI: Olmo 3.1 32B Think"
|
||||
tier = "fast"
|
||||
context_window = 65536
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.5
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,58 @@
|
||||
# amazon — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "amazon"
|
||||
display_name = "Amazon"
|
||||
api_key_env = "AMAZON_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "nova-2-lite-v1"
|
||||
display_name = "Amazon: Nova 2 Lite"
|
||||
tier = "fast"
|
||||
context_window = 1000000
|
||||
max_output_tokens = 65535
|
||||
input_cost_per_m = 0.3
|
||||
output_cost_per_m = 2.5
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nova-lite-v1"
|
||||
display_name = "Amazon: Nova Lite 1.0"
|
||||
tier = "fast"
|
||||
context_window = 300000
|
||||
max_output_tokens = 5120
|
||||
input_cost_per_m = 0.06
|
||||
output_cost_per_m = 0.24
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nova-micro-v1"
|
||||
display_name = "Amazon: Nova Micro 1.0"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 5120
|
||||
input_cost_per_m = 0.035
|
||||
output_cost_per_m = 0.14
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nova-premier-v1"
|
||||
display_name = "Amazon: Nova Premier 1.0"
|
||||
tier = "smart"
|
||||
context_window = 1000000
|
||||
max_output_tokens = 32000
|
||||
input_cost_per_m = 2.5
|
||||
output_cost_per_m = 12.5
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nova-pro-v1"
|
||||
display_name = "Amazon: Nova Pro 1.0"
|
||||
tier = "smart"
|
||||
context_window = 300000
|
||||
max_output_tokens = 5120
|
||||
input_cost_per_m = 0.8
|
||||
output_cost_per_m = 3.2
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,78 @@
|
||||
# arcee-ai — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "arcee-ai"
|
||||
display_name = "Arcee Ai"
|
||||
api_key_env = "ARCEE_AI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "coder-large"
|
||||
display_name = "Arcee AI: Coder Large"
|
||||
tier = "smart"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.5
|
||||
output_cost_per_m = 0.8
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "maestro-reasoning"
|
||||
display_name = "Arcee AI: Maestro Reasoning"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 32000
|
||||
input_cost_per_m = 0.9
|
||||
output_cost_per_m = 3.3
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "spotlight"
|
||||
display_name = "Arcee AI: Spotlight"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 65537
|
||||
input_cost_per_m = 0.18
|
||||
output_cost_per_m = 0.18
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "trinity-large-preview:free"
|
||||
display_name = "Arcee AI: Trinity Large Preview (free)"
|
||||
tier = "free"
|
||||
context_window = 131000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "trinity-mini"
|
||||
display_name = "Arcee AI: Trinity Mini"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.045
|
||||
output_cost_per_m = 0.15
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "trinity-mini:free"
|
||||
display_name = "Arcee AI: Trinity Mini (free)"
|
||||
tier = "free"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "virtuoso-large"
|
||||
display_name = "Arcee AI: Virtuoso Large"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 64000
|
||||
input_cost_per_m = 0.75
|
||||
output_cost_per_m = 1.2
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# bytedance — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "bytedance"
|
||||
display_name = "Bytedance"
|
||||
api_key_env = "BYTEDANCE_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "ui-tars-1.5-7b"
|
||||
display_name = "ByteDance: UI-TARS 7B "
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 2048
|
||||
input_cost_per_m = 0.1
|
||||
output_cost_per_m = 0.2
|
||||
supports_streaming = true
|
||||
@@ -27,8 +27,8 @@ display_name = "GPT-5.3 Codex"
|
||||
tier = "frontier"
|
||||
context_window = 200000
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
input_cost_per_m = 1.75
|
||||
output_cost_per_m = 14.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -40,8 +40,8 @@ display_name = "GPT-5.2 Codex"
|
||||
tier = "smart"
|
||||
context_window = 200000
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
input_cost_per_m = 1.75
|
||||
output_cost_per_m = 14.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -53,8 +53,8 @@ display_name = "GPT-5.1 Codex"
|
||||
tier = "smart"
|
||||
context_window = 200000
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
input_cost_per_m = 1.25
|
||||
output_cost_per_m = 10.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -66,8 +66,8 @@ display_name = "GPT-5.1 Codex Mini"
|
||||
tier = "balanced"
|
||||
context_window = 200000
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
input_cost_per_m = 0.25
|
||||
output_cost_per_m = 2.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,18 @@
|
||||
# deepcogito — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "deepcogito"
|
||||
display_name = "Deepcogito"
|
||||
api_key_env = "DEEPCOGITO_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "cogito-v2.1-671b"
|
||||
display_name = "Deep Cogito: Cogito v2.1 671B"
|
||||
tier = "smart"
|
||||
context_window = 128000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 1.25
|
||||
output_cost_per_m = 1.25
|
||||
supports_streaming = true
|
||||
@@ -14,8 +14,8 @@ display_name = "DeepSeek V3"
|
||||
tier = "smart"
|
||||
context_window = 64000
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.27
|
||||
output_cost_per_m = 1.10
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.75
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -53,8 +53,8 @@ display_name = "DeepSeek V3 0324"
|
||||
tier = "smart"
|
||||
context_window = 64000
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.27
|
||||
output_cost_per_m = 1.10
|
||||
input_cost_per_m = 0.19999999999999998
|
||||
output_cost_per_m = 0.77
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,18 @@
|
||||
# eleutherai — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "eleutherai"
|
||||
display_name = "Eleutherai"
|
||||
api_key_env = "ELEUTHERAI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "llemma_7b"
|
||||
display_name = "EleutherAI: Llemma 7b"
|
||||
tier = "smart"
|
||||
context_window = 4096
|
||||
max_output_tokens = 4096
|
||||
input_cost_per_m = 0.8
|
||||
output_cost_per_m = 1.2
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# essentialai — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "essentialai"
|
||||
display_name = "Essentialai"
|
||||
api_key_env = "ESSENTIALAI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "rnj-1-instruct"
|
||||
display_name = "EssentialAI: Rnj 1 Instruct"
|
||||
tier = "fast"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.15
|
||||
supports_streaming = true
|
||||
+12
-12
@@ -15,8 +15,8 @@ display_name = "Gemini 3.1 Pro Preview"
|
||||
tier = "frontier"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 2.50
|
||||
output_cost_per_m = 15.0
|
||||
input_cost_per_m = 2.0
|
||||
output_cost_per_m = 12.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -28,8 +28,8 @@ display_name = "Gemini 3 Flash Preview"
|
||||
tier = "smart"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.60
|
||||
input_cost_per_m = 0.5
|
||||
output_cost_per_m = 3.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -41,8 +41,8 @@ display_name = "Gemini 3.1 Flash Lite Preview"
|
||||
tier = "fast"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.04
|
||||
output_cost_per_m = 0.15
|
||||
input_cost_per_m = 0.25
|
||||
output_cost_per_m = 1.5
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -54,8 +54,8 @@ display_name = "Gemini 2.5 Flash Lite"
|
||||
tier = "fast"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.04
|
||||
output_cost_per_m = 0.15
|
||||
input_cost_per_m = 0.09999999999999999
|
||||
output_cost_per_m = 0.39999999999999997
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -80,8 +80,8 @@ display_name = "Gemini 2.5 Flash"
|
||||
tier = "smart"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.60
|
||||
input_cost_per_m = 0.3
|
||||
output_cost_per_m = 2.5
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -93,8 +93,8 @@ display_name = "Gemini 2.0 Flash"
|
||||
tier = "fast"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.10
|
||||
output_cost_per_m = 0.40
|
||||
input_cost_per_m = 0.075
|
||||
output_cost_per_m = 0.3
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,18 @@
|
||||
# ibm-granite — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "ibm-granite"
|
||||
display_name = "Ibm Granite"
|
||||
api_key_env = "IBM_GRANITE_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "granite-4.0-h-micro"
|
||||
display_name = "IBM: Granite 4.0 Micro"
|
||||
tier = "fast"
|
||||
context_window = 131000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.017
|
||||
output_cost_per_m = 0.11
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,38 @@
|
||||
# inception — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "inception"
|
||||
display_name = "Inception"
|
||||
api_key_env = "INCEPTION_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "mercury"
|
||||
display_name = "Inception: Mercury"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 32000
|
||||
input_cost_per_m = 0.25
|
||||
output_cost_per_m = 0.75
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "mercury-2"
|
||||
display_name = "Inception: Mercury 2"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 50000
|
||||
input_cost_per_m = 0.25
|
||||
output_cost_per_m = 0.75
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "mercury-coder"
|
||||
display_name = "Inception: Mercury Coder"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 32000
|
||||
input_cost_per_m = 0.25
|
||||
output_cost_per_m = 0.75
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,28 @@
|
||||
# inflection — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "inflection"
|
||||
display_name = "Inflection"
|
||||
api_key_env = "INFLECTION_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "inflection-3-pi"
|
||||
display_name = "Inflection: Inflection 3 Pi"
|
||||
tier = "smart"
|
||||
context_window = 8000
|
||||
max_output_tokens = 1024
|
||||
input_cost_per_m = 2.5
|
||||
output_cost_per_m = 10.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "inflection-3-productivity"
|
||||
display_name = "Inflection: Inflection 3 Productivity"
|
||||
tier = "smart"
|
||||
context_window = 8000
|
||||
max_output_tokens = 1024
|
||||
input_cost_per_m = 2.5
|
||||
output_cost_per_m = 10.0
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# kwaipilot — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "kwaipilot"
|
||||
display_name = "Kwaipilot"
|
||||
api_key_env = "KWAIPILOT_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "kat-coder-pro"
|
||||
display_name = "Kwaipilot: KAT-Coder-Pro V1"
|
||||
tier = "fast"
|
||||
context_window = 256000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 0.207
|
||||
output_cost_per_m = 0.828
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,58 @@
|
||||
# liquid — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "liquid"
|
||||
display_name = "Liquid"
|
||||
api_key_env = "LIQUID_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "lfm-2-24b-a2b"
|
||||
display_name = "LiquidAI: LFM2-24B-A2B"
|
||||
tier = "fast"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.03
|
||||
output_cost_per_m = 0.12
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "lfm-2.2-6b"
|
||||
display_name = "LiquidAI: LFM2-2.6B"
|
||||
tier = "fast"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.01
|
||||
output_cost_per_m = 0.02
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "lfm-2.5-1.2b-instruct:free"
|
||||
display_name = "LiquidAI: LFM2.5-1.2B-Instruct (free)"
|
||||
tier = "free"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "lfm-2.5-1.2b-thinking:free"
|
||||
display_name = "LiquidAI: LFM2.5-1.2B-Thinking (free)"
|
||||
tier = "free"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "lfm2-8b-a1b"
|
||||
display_name = "LiquidAI: LFM2-8B-A1B"
|
||||
tier = "fast"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.01
|
||||
output_cost_per_m = 0.02
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# meituan — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "meituan"
|
||||
display_name = "Meituan"
|
||||
api_key_env = "MEITUAN_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "longcat-flash-chat"
|
||||
display_name = "Meituan: LongCat Flash Chat"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.2
|
||||
output_cost_per_m = 0.8
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,148 @@
|
||||
# meta-llama — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "meta-llama"
|
||||
display_name = "Meta Llama"
|
||||
api_key_env = "META_LLAMA_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3-70b-instruct"
|
||||
display_name = "Meta: Llama 3 70B Instruct"
|
||||
tier = "smart"
|
||||
context_window = 8192
|
||||
max_output_tokens = 8000
|
||||
input_cost_per_m = 0.51
|
||||
output_cost_per_m = 0.74
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3-8b-instruct"
|
||||
display_name = "Meta: Llama 3 8B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 8192
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.03
|
||||
output_cost_per_m = 0.04
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.1-70b-instruct"
|
||||
display_name = "Meta: Llama 3.1 70B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.4
|
||||
output_cost_per_m = 0.4
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.1-8b-instruct"
|
||||
display_name = "Meta: Llama 3.1 8B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 16384
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.02
|
||||
output_cost_per_m = 0.05
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.2-11b-vision-instruct"
|
||||
display_name = "Meta: Llama 3.2 11B Vision Instruct"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.049
|
||||
output_cost_per_m = 0.049
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.2-1b-instruct"
|
||||
display_name = "Meta: Llama 3.2 1B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 60000
|
||||
max_output_tokens = 15000
|
||||
input_cost_per_m = 0.027
|
||||
output_cost_per_m = 0.2
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.2-3b-instruct"
|
||||
display_name = "Meta: Llama 3.2 3B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 80000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.051
|
||||
output_cost_per_m = 0.34
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.2-3b-instruct:free"
|
||||
display_name = "Meta: Llama 3.2 3B Instruct (free)"
|
||||
tier = "free"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.3-70b-instruct"
|
||||
display_name = "Meta: Llama 3.3 70B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.1
|
||||
output_cost_per_m = 0.32
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.3-70b-instruct:free"
|
||||
display_name = "Meta: Llama 3.3 70B Instruct (free)"
|
||||
tier = "free"
|
||||
context_window = 65536
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-4-maverick"
|
||||
display_name = "Meta: Llama 4 Maverick"
|
||||
tier = "fast"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.6
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-4-scout"
|
||||
display_name = "Meta: Llama 4 Scout"
|
||||
tier = "fast"
|
||||
context_window = 327680
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.08
|
||||
output_cost_per_m = 0.3
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-guard-3-8b"
|
||||
display_name = "Llama Guard 3 8B"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.02
|
||||
output_cost_per_m = 0.06
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-guard-4-12b"
|
||||
display_name = "Meta: Llama Guard 4 12B"
|
||||
tier = "fast"
|
||||
context_window = 163840
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.18
|
||||
output_cost_per_m = 0.18
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,28 @@
|
||||
# microsoft — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "microsoft"
|
||||
display_name = "Microsoft"
|
||||
api_key_env = "MICROSOFT_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "phi-4"
|
||||
display_name = "Microsoft: Phi 4"
|
||||
tier = "fast"
|
||||
context_window = 16384
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.065
|
||||
output_cost_per_m = 0.14
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "wizardlm-2-8x22b"
|
||||
display_name = "WizardLM-2 8x22B"
|
||||
tier = "smart"
|
||||
context_window = 65535
|
||||
max_output_tokens = 8000
|
||||
input_cost_per_m = 0.62
|
||||
output_cost_per_m = 0.62
|
||||
supports_streaming = true
|
||||
@@ -53,8 +53,8 @@ display_name = "Kimi K2"
|
||||
tier = "frontier"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.60
|
||||
output_cost_per_m = 2.50
|
||||
input_cost_per_m = 0.44999999999999996
|
||||
output_cost_per_m = 2.2
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,28 @@
|
||||
# morph — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "morph"
|
||||
display_name = "Morph"
|
||||
api_key_env = "MORPH_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "morph-v3-fast"
|
||||
display_name = "Morph: Morph V3 Fast"
|
||||
tier = "smart"
|
||||
context_window = 81920
|
||||
max_output_tokens = 38000
|
||||
input_cost_per_m = 0.8
|
||||
output_cost_per_m = 1.2
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "morph-v3-large"
|
||||
display_name = "Morph: Morph V3 Large"
|
||||
tier = "smart"
|
||||
context_window = 262144
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.9
|
||||
output_cost_per_m = 1.9
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# nex-agi — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "nex-agi"
|
||||
display_name = "Nex Agi"
|
||||
api_key_env = "NEX_AGI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "deepseek-v3.1-nex-n1"
|
||||
display_name = "Nex AGI: DeepSeek V3.1 Nex N1"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 163840
|
||||
input_cost_per_m = 0.135
|
||||
output_cost_per_m = 0.5
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,68 @@
|
||||
# nousresearch — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "nousresearch"
|
||||
display_name = "Nousresearch"
|
||||
api_key_env = "NOUSRESEARCH_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-2-pro-llama-3-8b"
|
||||
display_name = "NousResearch: Hermes 2 Pro - Llama-3 8B"
|
||||
tier = "fast"
|
||||
context_window = 8192
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.14
|
||||
output_cost_per_m = 0.14
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-3-llama-3.1-405b"
|
||||
display_name = "Nous: Hermes 3 405B Instruct"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 1.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-3-llama-3.1-405b:free"
|
||||
display_name = "Nous: Hermes 3 405B Instruct (free)"
|
||||
tier = "free"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-3-llama-3.1-70b"
|
||||
display_name = "Nous: Hermes 3 70B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.3
|
||||
output_cost_per_m = 0.3
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-4-405b"
|
||||
display_name = "Nous: Hermes 4 405B"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 3.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "hermes-4-70b"
|
||||
display_name = "Nous: Hermes 4 70B"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.13
|
||||
output_cost_per_m = 0.4
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,118 @@
|
||||
# nvidia — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "nvidia"
|
||||
display_name = "Nvidia"
|
||||
api_key_env = "NVIDIA_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.1-nemotron-70b-instruct"
|
||||
display_name = "NVIDIA: Llama 3.1 Nemotron 70B Instruct"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 1.2
|
||||
output_cost_per_m = 1.2
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.1-nemotron-ultra-253b-v1"
|
||||
display_name = "NVIDIA: Llama 3.1 Nemotron Ultra 253B v1"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.6
|
||||
output_cost_per_m = 1.8
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "llama-3.3-nemotron-super-49b-v1.5"
|
||||
display_name = "NVIDIA: Llama 3.3 Nemotron Super 49B V1.5"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.1
|
||||
output_cost_per_m = 0.4
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-3-nano-30b-a3b"
|
||||
display_name = "NVIDIA: Nemotron 3 Nano 30B A3B"
|
||||
tier = "fast"
|
||||
context_window = 262144
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.05
|
||||
output_cost_per_m = 0.2
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-3-nano-30b-a3b:free"
|
||||
display_name = "NVIDIA: Nemotron 3 Nano 30B A3B (free)"
|
||||
tier = "free"
|
||||
context_window = 256000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-3-super-120b-a12b"
|
||||
display_name = "NVIDIA: Nemotron 3 Super"
|
||||
tier = "fast"
|
||||
context_window = 262144
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.1
|
||||
output_cost_per_m = 0.5
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-3-super-120b-a12b:free"
|
||||
display_name = "NVIDIA: Nemotron 3 Super (free)"
|
||||
tier = "free"
|
||||
context_window = 262144
|
||||
max_output_tokens = 262144
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-nano-12b-v2-vl"
|
||||
display_name = "NVIDIA: Nemotron Nano 12B 2 VL"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.2
|
||||
output_cost_per_m = 0.6
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-nano-12b-v2-vl:free"
|
||||
display_name = "NVIDIA: Nemotron Nano 12B 2 VL (free)"
|
||||
tier = "free"
|
||||
context_window = 128000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-nano-9b-v2"
|
||||
display_name = "NVIDIA: Nemotron Nano 9B V2"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.04
|
||||
output_cost_per_m = 0.16
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "nemotron-nano-9b-v2:free"
|
||||
display_name = "NVIDIA: Nemotron Nano 9B V2 (free)"
|
||||
tier = "free"
|
||||
context_window = 128000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_streaming = true
|
||||
@@ -53,8 +53,8 @@ display_name = "Qwen 2.5 (Ollama)"
|
||||
tier = "local"
|
||||
context_window = 32768
|
||||
max_output_tokens = 4096
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
input_cost_per_m = 0.03
|
||||
output_cost_per_m = 0.09
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
|
||||
+10
-10
@@ -80,8 +80,8 @@ display_name = "o3"
|
||||
tier = "frontier"
|
||||
context_window = 200000
|
||||
max_output_tokens = 100000
|
||||
input_cost_per_m = 2.00
|
||||
output_cost_per_m = 8.00
|
||||
input_cost_per_m = 10.0
|
||||
output_cost_per_m = 40.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -106,8 +106,8 @@ display_name = "o4-mini"
|
||||
tier = "smart"
|
||||
context_window = 200000
|
||||
max_output_tokens = 100000
|
||||
input_cost_per_m = 1.10
|
||||
output_cost_per_m = 4.40
|
||||
input_cost_per_m = 2.0
|
||||
output_cost_per_m = 8.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -132,8 +132,8 @@ display_name = "GPT-3.5 Turbo"
|
||||
tier = "fast"
|
||||
context_window = 16385
|
||||
max_output_tokens = 4096
|
||||
input_cost_per_m = 0.50
|
||||
output_cost_per_m = 1.50
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 2.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -145,8 +145,8 @@ display_name = "GPT-5"
|
||||
tier = "frontier"
|
||||
context_window = 400000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 1.25
|
||||
output_cost_per_m = 10.0
|
||||
input_cost_per_m = 0.19999999999999998
|
||||
output_cost_per_m = 1.25
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -210,8 +210,8 @@ display_name = "GPT-5.2 Pro"
|
||||
tier = "frontier"
|
||||
context_window = 400000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 1.75
|
||||
output_cost_per_m = 14.0
|
||||
input_cost_per_m = 21.0
|
||||
output_cost_per_m = 168.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -40,8 +40,8 @@ display_name = "Sonar Reasoning"
|
||||
tier = "balanced"
|
||||
context_window = 128000
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 5.0
|
||||
input_cost_per_m = 2.0
|
||||
output_cost_per_m = 8.0
|
||||
supports_tools = false
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,18 @@
|
||||
# prime-intellect — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "prime-intellect"
|
||||
display_name = "Prime Intellect"
|
||||
api_key_env = "PRIME_INTELLECT_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "intellect-3"
|
||||
display_name = "Prime Intellect: INTELLECT-3"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.2
|
||||
output_cost_per_m = 1.1
|
||||
supports_streaming = true
|
||||
+14
-14
@@ -21,8 +21,8 @@ display_name = "Qwen Max"
|
||||
tier = "frontier"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 4.00
|
||||
output_cost_per_m = 12.00
|
||||
input_cost_per_m = 1.04
|
||||
output_cost_per_m = 4.16
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -34,8 +34,8 @@ display_name = "Qwen Plus"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.80
|
||||
output_cost_per_m = 2.00
|
||||
input_cost_per_m = 0.26
|
||||
output_cost_per_m = 0.78
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -47,8 +47,8 @@ display_name = "Qwen Turbo"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.30
|
||||
output_cost_per_m = 0.60
|
||||
input_cost_per_m = 0.0325
|
||||
output_cost_per_m = 0.13
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -60,8 +60,8 @@ display_name = "Qwen VL Plus"
|
||||
tier = "smart"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 1.50
|
||||
output_cost_per_m = 4.50
|
||||
input_cost_per_m = 0.1365
|
||||
output_cost_per_m = 0.40950000000000003
|
||||
supports_tools = false
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
@@ -99,8 +99,8 @@ display_name = "Qwen3 235B"
|
||||
tier = "frontier"
|
||||
context_window = 131072
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 4.00
|
||||
output_cost_per_m = 12.00
|
||||
input_cost_per_m = 0.14950000000000002
|
||||
output_cost_per_m = 1.495
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -112,8 +112,8 @@ display_name = "Qwen3 30B"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.30
|
||||
output_cost_per_m = 0.60
|
||||
input_cost_per_m = 0.08
|
||||
output_cost_per_m = 0.39999999999999997
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -151,8 +151,8 @@ display_name = "Qwen VL Max"
|
||||
tier = "frontier"
|
||||
context_window = 32768
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 3.00
|
||||
output_cost_per_m = 9.00
|
||||
input_cost_per_m = 0.52
|
||||
output_cost_per_m = 2.08
|
||||
supports_tools = false
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,28 @@
|
||||
# relace — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "relace"
|
||||
display_name = "Relace"
|
||||
api_key_env = "RELACE_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "relace-apply-3"
|
||||
display_name = "Relace: Relace Apply 3"
|
||||
tier = "smart"
|
||||
context_window = 256000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 0.85
|
||||
output_cost_per_m = 1.25
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "relace-search"
|
||||
display_name = "Relace: Relace Search"
|
||||
tier = "smart"
|
||||
context_window = 256000
|
||||
max_output_tokens = 128000
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 3.0
|
||||
supports_streaming = true
|
||||
@@ -14,8 +14,8 @@ display_name = "Step 3.5 Flash"
|
||||
tier = "smart"
|
||||
context_window = 256000
|
||||
max_output_tokens = 256000
|
||||
input_cost_per_m = 0.10
|
||||
output_cost_per_m = 0.30
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
@@ -27,8 +27,8 @@ display_name = "Step 3"
|
||||
tier = "frontier"
|
||||
context_window = 65536
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.57
|
||||
output_cost_per_m = 1.43
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,18 @@
|
||||
# switchpoint — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "switchpoint"
|
||||
display_name = "Switchpoint"
|
||||
api_key_env = "SWITCHPOINT_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "router"
|
||||
display_name = "Switchpoint Router"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.85
|
||||
output_cost_per_m = 3.4
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# tencent — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "tencent"
|
||||
display_name = "Tencent"
|
||||
api_key_env = "TENCENT_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "hunyuan-a13b-instruct"
|
||||
display_name = "Tencent: Hunyuan A13B Instruct"
|
||||
tier = "fast"
|
||||
context_window = 131072
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 0.14
|
||||
output_cost_per_m = 0.57
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# tngtech — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "tngtech"
|
||||
display_name = "Tngtech"
|
||||
api_key_env = "TNGTECH_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "deepseek-r1t2-chimera"
|
||||
display_name = "TNG: DeepSeek R1T2 Chimera"
|
||||
tier = "fast"
|
||||
context_window = 163840
|
||||
max_output_tokens = 163840
|
||||
input_cost_per_m = 0.3
|
||||
output_cost_per_m = 1.1
|
||||
supports_streaming = true
|
||||
@@ -0,0 +1,18 @@
|
||||
# upstage — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "upstage"
|
||||
display_name = "Upstage"
|
||||
api_key_env = "UPSTAGE_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "solar-pro-3"
|
||||
display_name = "Upstage: Solar Pro 3"
|
||||
tier = "fast"
|
||||
context_window = 128000
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.15
|
||||
output_cost_per_m = 0.6
|
||||
supports_streaming = true
|
||||
@@ -27,8 +27,8 @@ display_name = "Llama 3.3 70B (Venice)"
|
||||
tier = "balanced"
|
||||
context_window = 128000
|
||||
max_output_tokens = 8192
|
||||
input_cost_per_m = 0.20
|
||||
output_cost_per_m = 0.90
|
||||
input_cost_per_m = 0.0
|
||||
output_cost_per_m = 0.0
|
||||
supports_tools = true
|
||||
supports_vision = false
|
||||
supports_streaming = true
|
||||
|
||||
+2
-2
@@ -79,8 +79,8 @@ display_name = "Grok 3"
|
||||
tier = "frontier"
|
||||
context_window = 131072
|
||||
max_output_tokens = 32768
|
||||
input_cost_per_m = 3.0
|
||||
output_cost_per_m = 15.0
|
||||
input_cost_per_m = 0.3
|
||||
output_cost_per_m = 0.5
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
@@ -0,0 +1,38 @@
|
||||
# xiaomi — auto-generated from OpenRouter API
|
||||
|
||||
[provider]
|
||||
id = "xiaomi"
|
||||
display_name = "Xiaomi"
|
||||
api_key_env = "XIAOMI_API_KEY"
|
||||
base_url = ""
|
||||
key_required = true
|
||||
|
||||
[[models]]
|
||||
id = "mimo-v2-flash"
|
||||
display_name = "Xiaomi: MiMo-V2-Flash"
|
||||
tier = "fast"
|
||||
context_window = 262144
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.09
|
||||
output_cost_per_m = 0.29
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "mimo-v2-omni"
|
||||
display_name = "Xiaomi: MiMo-V2-Omni"
|
||||
tier = "fast"
|
||||
context_window = 262144
|
||||
max_output_tokens = 65536
|
||||
input_cost_per_m = 0.4
|
||||
output_cost_per_m = 2.0
|
||||
supports_streaming = true
|
||||
|
||||
[[models]]
|
||||
id = "mimo-v2-pro"
|
||||
display_name = "Xiaomi: MiMo-V2-Pro"
|
||||
tier = "smart"
|
||||
context_window = 1048576
|
||||
max_output_tokens = 131072
|
||||
input_cost_per_m = 1.0
|
||||
output_cost_per_m = 3.0
|
||||
supports_streaming = true
|
||||
@@ -79,8 +79,8 @@ display_name = "GLM-4.7"
|
||||
tier = "smart"
|
||||
context_window = 131072
|
||||
max_output_tokens = 16384
|
||||
input_cost_per_m = 0.60
|
||||
output_cost_per_m = 2.20
|
||||
input_cost_per_m = 0.06
|
||||
output_cost_per_m = 0.39999999999999997
|
||||
supports_tools = true
|
||||
supports_vision = true
|
||||
supports_streaming = true
|
||||
|
||||
Reference in new issue
Block a user