Files
librefang-registry/providers/nvidia.toml
T
Evan fc37ce253b fix: correct invalid tier "free" and teams-mcp id with version (#29)
- Replace tier "free" with "fast" (valid tiers: frontier/smart/balanced/fast/local)
- Remove version suffix from teams-mcp integration id field
- Update sync-pricing.py to not generate invalid tier values
2026-03-25 23:49:39 +09:00

119 lines
2.7 KiB
TOML

# nvidia — auto-generated from OpenRouter API
[provider]
id = "nvidia"
display_name = "Nvidia"
api_key_env = "NVIDIA_API_KEY"
base_url = ""
key_required = true
[[models]]
id = "llama-3.1-nemotron-70b-instruct"
display_name = "NVIDIA: Llama 3.1 Nemotron 70B Instruct"
tier = "smart"
context_window = 131072
max_output_tokens = 16384
input_cost_per_m = 1.2
output_cost_per_m = 1.2
supports_streaming = true
[[models]]
id = "llama-3.1-nemotron-ultra-253b-v1"
display_name = "NVIDIA: Llama 3.1 Nemotron Ultra 253B v1"
tier = "smart"
context_window = 131072
max_output_tokens = 16384
input_cost_per_m = 0.6
output_cost_per_m = 1.8
supports_streaming = true
[[models]]
id = "llama-3.3-nemotron-super-49b-v1.5"
display_name = "NVIDIA: Llama 3.3 Nemotron Super 49B V1.5"
tier = "fast"
context_window = 131072
max_output_tokens = 16384
input_cost_per_m = 0.1
output_cost_per_m = 0.4
supports_streaming = true
[[models]]
id = "nemotron-3-nano-30b-a3b"
display_name = "NVIDIA: Nemotron 3 Nano 30B A3B"
tier = "fast"
context_window = 262144
max_output_tokens = 16384
input_cost_per_m = 0.05
output_cost_per_m = 0.2
supports_streaming = true
[[models]]
id = "nemotron-3-nano-30b-a3b:free"
display_name = "NVIDIA: Nemotron 3 Nano 30B A3B (free)"
tier = "fast"
context_window = 256000
max_output_tokens = 16384
input_cost_per_m = 0.0
output_cost_per_m = 0.0
supports_streaming = true
[[models]]
id = "nemotron-3-super-120b-a12b"
display_name = "NVIDIA: Nemotron 3 Super"
tier = "fast"
context_window = 262144
max_output_tokens = 16384
input_cost_per_m = 0.1
output_cost_per_m = 0.5
supports_streaming = true
[[models]]
id = "nemotron-3-super-120b-a12b:free"
display_name = "NVIDIA: Nemotron 3 Super (free)"
tier = "fast"
context_window = 262144
max_output_tokens = 262144
input_cost_per_m = 0.0
output_cost_per_m = 0.0
supports_streaming = true
[[models]]
id = "nemotron-nano-12b-v2-vl"
display_name = "NVIDIA: Nemotron Nano 12B 2 VL"
tier = "fast"
context_window = 131072
max_output_tokens = 16384
input_cost_per_m = 0.2
output_cost_per_m = 0.6
supports_streaming = true
[[models]]
id = "nemotron-nano-12b-v2-vl:free"
display_name = "NVIDIA: Nemotron Nano 12B 2 VL (free)"
tier = "fast"
context_window = 128000
max_output_tokens = 128000
input_cost_per_m = 0.0
output_cost_per_m = 0.0
supports_streaming = true
[[models]]
id = "nemotron-nano-9b-v2"
display_name = "NVIDIA: Nemotron Nano 9B V2"
tier = "fast"
context_window = 131072
max_output_tokens = 16384
input_cost_per_m = 0.04
output_cost_per_m = 0.16
supports_streaming = true
[[models]]
id = "nemotron-nano-9b-v2:free"
display_name = "NVIDIA: Nemotron Nano 9B V2 (free)"
tier = "fast"
context_window = 128000
max_output_tokens = 16384
input_cost_per_m = 0.0
output_cost_per_m = 0.0
supports_streaming = true