{"data":[{"id":"nvidia-nemotron-nano-3-30b-a3b-reasoning","name":"nvidia-nemotron-nano-3-30b-a3b-reasoning","display_name":"Nemotron Nano 3 30B A3B Reasoning","description":"A reasoning-specialized 30B-parameter hybrid MoE Nemotron Nano 3 model with 3B active parameters, designed for complex multi-step inference in agentic systems.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"","size_in_bn":null,"modalities":{"input":[],"output":[]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-nano-30b-a3b-reasoning","nvidia-nemotron-nano-3-30b-a3b-reasoning"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-04-20 16:51:01"},{"id":"nvidia-nemotron-lightning-3-5","name":"nemotron-lightning-3-5","display_name":"Nemotron Lightning 3.5","description":"NVIDIA's fast open model targeting always-on agents and high-volume specialized tasks with strong agentic capability.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":31.578,"modalities":{"input":["text"],"output":["text"]},"context_window":1048576,"max_output_tokens":235929,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":"Other","capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":true,"prompt_caching":true,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":true,"adaptive_reasoning":false},"release_date":"2026-08-11","earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":4,"ids":["deepinfra/nvidia/NVIDIA-Nemotron-3.5-Lightning","nebius/nvidia/Nemotron-3_5-Lightning","nemotron-3-5-lightning","nvidia-nemotron-lightning-3-5","nvidia/nemotron-3.5-lightning","nvidia/nemotron-3.5-lightning:free","nvidia/NVIDIA-Nemotron-3.5-Lightning","openrouter/nvidia/nemotron-3.5-lightning"],"hf_likes":120,"hf_downloads":15740,"hf_downloads_all_time":15740,"hf_trending_score":120,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-cascade-2-30b-a3b","name":"nemotron-cascade-2-30b-a3b","display_name":"Nemotron Cascade 2 30B A3B","description":"A 30B-parameter cascaded Nemotron model from NVIDIA with 3B active parameters, designed for efficient multi-stage reasoning and inference pipelines.","creator":"nvidia","family":"nemotron","tier":"","version":"2","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nemotron-cascade-2-30b-a3b","nvidia-nemotron-cascade-2-30b-a3b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-3-30b-a3b-omni","name":"nvidia-nemotron-nano-3-30b-a3b-omni","display_name":"Nemotron Nano 3 30B A3B Omni","description":"An open multimodal Mixture-of-Experts model from NVIDIA with 30B total and 3B active parameters, capable of reasoning across text, images, video, and audio inputs.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["accounts/fireworks/models/nvidia-nemotron-3-nano-omni-30b-a3b","nemotron-3-nano-omni-30b-a3b","nvidia-nemotron-nano-3-30b-a3b-omni"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-2-12b-vl-reasoning","name":"nvidia-nemotron-nano-2-12b-vl-reasoning","display_name":"Nemotron Nano 2 12B VL Reasoning","description":"A reasoning-specialized 12B-parameter vision-language variant of NVIDIA's Nemotron Nano v2, designed for complex visual and document reasoning tasks.","creator":"nvidia","family":"nemotron","tier":"","version":"2","type":"","size_in_bn":null,"modalities":{"input":[],"output":[]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-nano-12b-v2-vl-reasoning","nvidia-nemotron-nano-2-12b-vl-reasoning"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-04-20 16:51:01"},{"id":"nvidia-nemotron-nano-2-9b-reasoning","name":"nvidia-nemotron-nano-2-9b-reasoning","display_name":"Nemotron Nano 2 9B Reasoning","description":"A reasoning-focused 9B-parameter variant of NVIDIA's Nemotron Nano v2, optimized for extended chain-of-thought inference and complex problem solving.","creator":"nvidia","family":"nemotron","tier":"","version":"2","type":"","size_in_bn":null,"modalities":{"input":[],"output":[]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-nano-2-9b-reasoning","nvidia-nemotron-nano-9b-v2-reasoning"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-04-20 16:51:01"},{"id":"nvidia-nemotron-nano-3-4b","name":"nvidia-nemotron-nano-3-4b","display_name":"Nemotron Nano 3 4B","description":"A compact 4B-parameter third-generation Nemotron Nano model from NVIDIA, designed for efficient on-device reasoning and instruction-following tasks.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":4,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-nano-4b","nvidia-nemotron-nano-3-4b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-120b-a12b","name":"nemotron-120b-a12b","display_name":"Nemotron 120B A12B","description":"A 120B-parameter hybrid Mixture-of-Experts LLM from NVIDIA with 12B active parameters, designed for compute-efficient reasoning and agentic workloads.","creator":"nvidia","family":"nemotron","tier":"","version":null,"type":"language","size_in_bn":120,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["baseten/nvidia/Nemotron-120B-A12B","nvidia-nemotron-120b-a12b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3","name":"nemotron-3","display_name":"Nemotron 3","description":"The third-generation Nemotron LLM from NVIDIA, designed for multilingual text generation, coding, and instruction-following tasks.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3","publishers/nvidia/models/nemotron-v3"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-550b-a55b-ultra-free","name":"nemotron-3-550b-a55b-ultra-free","display_name":"Nemotron 3 550B A55B Ultra Free","description":"A free, very large mixture-of-experts LLM in NVIDIA's Nemotron series designed for high-capacity general-purpose reasoning.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":550,"modalities":{"input":["text"],"output":["text"]},"context_window":1000000,"max_output_tokens":65536,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-550b-a55b-ultra-free","openrouter/nvidia/nemotron-3-ultra-550b-a55b:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-nano-30b","name":"nemotron-3-nano-30b","display_name":"Nemotron 3 Nano 30B","description":"A compact large language model in NVIDIA's Nemotron series, sized at 30B parameters for efficient reasoning and generation tasks.","creator":"nvidia","family":"nemotron","tier":"nano","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":262144,"max_output_tokens":8192,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":true,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["nvidia-nemotron-3-nano-30b","us-gov.nvidia.nemotron-nano-3-30b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-nano-30b-a3b","name":"nemotron-3-nano-30b-a3b","display_name":"Nemotron 3 Nano 30B A3B","description":"A 30B-parameter mixture-of-experts Nemotron model with 3B active parameters, optimized for efficient reasoning tasks.","creator":"nvidia","family":"nemotron","tier":"nano","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["huggingface-reasoning-nvidia-nemotron-3-nano-30b-a3b-bf16","nvidia-nemotron-3-nano-30b-a3b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-nano-30b-a3b-omni-reasoning-free","name":"nemotron-3-nano-30b-a3b-omni-reasoning-free","display_name":"Nemotron 3 Nano 30B A3B Omni Reasoning Free","description":"A free compact mixture-of-experts multimodal LLM in NVIDIA's Nemotron series tuned for omni-modal reasoning tasks.","creator":"nvidia","family":"nemotron","tier":"nano","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["audio","image"],"output":["text"]},"context_window":256000,"max_output_tokens":65536,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-nano-30b-a3b-omni-reasoning-free","openrouter/nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-super","name":"nemotron-3-super","display_name":"Nemotron 3 Super","description":"A Super-tier variant of NVIDIA's third-generation Nemotron, optimized for high compute efficiency in multi-agent and specialized agentic applications.","creator":"nvidia","family":"nemotron","tier":"super","version":"3","type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-super","publishers/nvidia/models/nemotron-3-super"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-super-120b","name":"nemotron-3-super-120b","display_name":"Nemotron 3 Super 120B","description":"A 120B-parameter Super-tier Nemotron 3 model from NVIDIA, engineered for maximum compute efficiency and accuracy in agentic and multi-agent systems.","creator":"nvidia","family":"nemotron","tier":"super","version":"3","type":"language","size_in_bn":120,"modalities":{"input":["text"],"output":["text"]},"context_window":256000,"max_output_tokens":32768,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["nvidia-nemotron-3-super-120b","nvidia-nemotron-3-super-120b-fp8-sg","publishers/nvidia/models/nemotron-v3-super-120b","us-gov.nvidia.nemotron-super-3-120b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-super-120b-a12b","name":"nemotron-3-super-120b-a12b","display_name":"Nemotron 3 Super 120B A12B","description":"A 120B-parameter hybrid MoE Nemotron 3 Super model with 12B active parameters, optimized by NVIDIA for reasoning-intensive agentic and production workloads.","creator":"nvidia","family":"nemotron","tier":"super","version":"3","type":"language","size_in_bn":120,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["accounts/fireworks/models/nemotron-3-super-120b-a12b-bf16","huggingface-reasoning-nvidia-nemotron-3-super-120b-a12b-fp8","nvidia-nemotron-3-super-120b-a12b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-super-120b-a12b-free","name":"nemotron-3-super-120b-a12b-free","display_name":"Nemotron 3 Super 120B A12B Free","description":"A free large-scale mixture-of-experts LLM in NVIDIA's Nemotron series built for high-throughput general-purpose reasoning.","creator":"nvidia","family":"nemotron","tier":"super","version":"3","type":"language","size_in_bn":120,"modalities":{"input":["text"],"output":["text"]},"context_window":262144,"max_output_tokens":235929,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":true,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-super-120b-a12b-free","openrouter/nvidia/nemotron-3-super-120b-a12b:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-ultra-nvfp4","name":"nemotron-3-ultra-nvfp4","display_name":"Nemotron 3 Ultra NVfp4","description":"A large-scale Nemotron LLM optimized for NVIDIA's NVfp4 quantization format, targeting high-throughput inference on NVIDIA hardware.","creator":"nvidia","family":"nemotron","tier":"ultra","version":"3","type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":262144,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":true,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["azure_ai/FW-Nemotron-3-Ultra-NVFP4","nvidia-nemotron-3-ultra-nvfp4"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-0-6b-asr-streaming","name":"nemotron-3-5-0-6b-asr-streaming","display_name":"Nemotron 3.5 0.6B ASR Streaming","description":"A compact 0.6B-parameter streaming automatic speech recognition model from NVIDIA's Nemotron family designed for low-latency real-time transcription.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"speech-to-text","size_in_bn":0.6,"modalities":{"input":["audio"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["huggingface-asr-nvidia-nemotron-3-5-asr-streaming-0-6b","nvidia-nemotron-3-5-0-6b-asr-streaming"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-0-6b-asr-streaming-multilingual","name":"nemotron-3-5-0-6b-asr-streaming-multilingual","display_name":"Nemotron 3.5 0.6B ASR Streaming Multilingual","description":"A compact 0.6B-parameter FastConformer-RNNT streaming ASR model optimized for low-latency multilingual transcription across 40+ languages.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"speech-to-text","size_in_bn":0.6,"modalities":{"input":["audio"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-5-0-6b-asr-streaming-multilingual","nvidia/Nemotron-3.5-ASR-Streaming-Multilingual-0.6b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-30b-a3b-lightning","name":"nvidia-nemotron-3-5-30b-a3b-lightning","display_name":"Nemotron 3.5 30B A3B Lightning","description":"A mixture-of-experts large language model in the Nemotron series optimized for fast, low-latency inference at scale.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":262144,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":true,"reasoning":true,"web_search":true,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":2,"ids":["nvidia-nemotron-3-5-30b-a3b-lightning","perplexity/perplexity/nemotron-3.5-lightning-30b-a3b","wandb/nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-content-safety","name":"nemotron-3-5-content-safety","display_name":"Nemotron 3.5 Content Safety","description":"A multimodal safety classifier from NVIDIA that handles text and images with support for custom policies, outputting safe/unsafe classifications for content moderation.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":null,"modalities":{"input":["image","text"],"output":["text"]},"context_window":131072,"max_output_tokens":117964,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":"Other","capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":true,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":"2026-06-04","earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":2,"ids":["deepinfra/nvidia/Nemotron-Content-Safety-3.5","nvidia-nemotron-3-5-content-safety","nvidia/nemotron-3.5-content-safety","nvidia/nemotron-3.5-content-safety:free","nvidia/Nemotron-Content-Safety-3.5","openrouter/nvidia/nemotron-3.5-content-safety"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-content-safety-free","name":"nemotron-3-5-content-safety-free","display_name":"Nemotron 3.5 Content Safety Free","description":"A free content-safety moderation model in NVIDIA's Nemotron series, designed to detect unsafe or policy-violating text content.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":null,"modalities":{"input":["image"],"output":["text"]},"context_window":128000,"max_output_tokens":8192,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-5-content-safety-free","openrouter/nvidia/nemotron-3.5-content-safety:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-lightning","name":"nemotron-3-5-lightning","display_name":"Nemotron 3.5 Lightning 30B","description":"A hybrid Mixture-of-Experts LLM with interleaved Mamba-2 and MoE layers, built for reasoning and tool-use workloads with implicit caching support.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":1000000,"max_output_tokens":32768,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":true,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":"2026-08-11","earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-5-lightning","nvidia/nemotron-3.5-lightning-free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-5-lightning-free","name":"nemotron-3-5-lightning-free","display_name":"Nemotron 3.5 Lightning Free","description":"A free, low-latency variant of NVIDIA's Nemotron LLM series optimized for fast inference on general-purpose tasks.","creator":"nvidia","family":"nemotron","tier":"","version":"3-5","type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":1000000,"max_output_tokens":65536,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-3-5-lightning-free","openrouter/nvidia/nemotron-3.5-lightning:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-4-15b","name":"nvidia-nemotron-4-15b","display_name":"Nemotron 4 15B","description":"A 15B-parameter multilingual LLM from NVIDIA's Nemotron-4 generation, trained on 8 trillion tokens for text generation, coding, and multilingual tasks.","creator":"nvidia","family":"nemotron","tier":"","version":"4","type":"language","size_in_bn":15,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-4-15b","nvidia-nemotron-4-15b-nim"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-3-30b","name":"nemotron-nano-3-30b","display_name":"Nemotron Nano 3 30B","description":"A 30B-parameter hybrid MoE Nemotron Nano model from NVIDIA's third generation, optimized for fast and cost-efficient reasoning in agentic production workloads.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["text"],"output":["text"]},"context_window":262144,"max_output_tokens":8192,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":true,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":true,"adaptive_reasoning":false},"release_date":"2025-12-23","earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["bedrock/us-gov-east-1/nvidia.nemotron-nano-3-30b","bedrock/us-gov-west-1/nvidia.nemotron-nano-3-30b","nvidia-nemotron-nano-3-30b","nvidia-nim-nemotron-3-nano-30b","nvidia.nemotron-nano-3-30b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-3-30b-a3b-omni-reasoning","name":"nvidia-nemotron-nano-3-30b-a3b-omni-reasoning","display_name":"Nemotron Nano 3 30B A3B Omni Reasoning","description":"A reasoning-optimized variant of NVIDIA's Nemotron Nano Omni multimodal MoE model, combining 30B total and 3B active parameters with extended chain-of-thought capabilities across text, image, video, and audio.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":30,"modalities":{"input":["audio","image","text","video"],"output":["text"]},"context_window":256000,"max_output_tokens":65536,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":"Other","capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":"2026-04-28","earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["huggingface-vlm-nemotron-3-nano-omni-30b-a3b-reasoning-bf16","huggingface-vlm-nvidia-nemotron3-nano-omni-30ba3b-reasoning-fp8","nvidia-nemotron-nano-3-30b-a3b-omni-reasoning","nvidia/nemotron-3-nano-omni-30b-a3b-reasoning","nvidia/Nemotron-3-Nano-Omni-30B-A3B-Reasoning","nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-3-omni","name":"nemotron-nano-3-omni","display_name":"Nemotron Nano 3 Omni","description":"A compact omni-modal Nemotron model from NVIDIA's Nano tier, combining text, vision, and audio understanding in a small, efficient footprint.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":null,"modalities":{"input":["image","text"],"output":["text"]},"context_window":262144,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nebius/nvidia/Nemotron-3-Nano-Omni","nvidia-nemotron-nano-3-omni"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-nano-8b","name":"nvidia-nemotron-nano-8b","display_name":"Nemotron Nano 8B","description":"An 8B-parameter Nemotron Nano LLM derived from Llama 3.1, adding reasoning capabilities to a compact model suited for efficient deployment.","creator":"nvidia","family":"nemotron","tier":"","version":null,"type":"language","size_in_bn":8,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-nano-8b","nvidia-nemotron-nano-8b-nim"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-parse","name":"nvidia-nemotron-parse","display_name":"Nemotron Parse","description":"A document-parsing model from NVIDIA's Nemotron family that extracts formatted text and bounding boxes from document images for structured information retrieval.","creator":"nvidia","family":"nemotron","tier":"","version":null,"type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-parse"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-super-1-5-49b","name":"nvidia-nemotron-super-1-5-49b","display_name":"Nemotron Super 1.5 49B","description":"A 49B-parameter Nemotron Super v1.5 LLM from NVIDIA, derived from Llama 3.3 70B and optimized for reasoning, RAG, and tool-calling in agentic workflows.","creator":"nvidia","family":"nemotron","tier":"","version":"1-5","type":"language","size_in_bn":49,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-super-1-5-49b","nvidia-nemotron-super-49b-nim-1-5"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-super-3-120b","name":"nemotron-super-3-120b","display_name":"Nemotron Super 3 120B","description":"A 120B-parameter third-generation Nemotron Super model from NVIDIA, engineered for highest compute efficiency and accuracy in multi-agent applications.","creator":"nvidia","family":"nemotron","tier":"","version":"3","type":"language","size_in_bn":120,"modalities":{"input":["text"],"output":["text"]},"context_window":256000,"max_output_tokens":32768,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":true,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":"2026-03-18","earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["bedrock/us-gov-east-1/nvidia.nemotron-super-3-120b","bedrock/us-gov-west-1/nvidia.nemotron-super-3-120b","nvidia-nemotron-super-3-120b","nvidia.nemotron-super-3-120b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-3-120b-a12b","name":"nemotron-3-120b-a12b","display_name":"Nemotron Super 3 120B A12B","description":"Hybrid mixture-of-experts model with 120B total and 12B active parameters, designed for high-accuracy multi-agent and specialized agentic AI applications.","creator":"nvidia","family":"nemotron","tier":"","version":null,"type":"language","size_in_bn":null,"modalities":{"input":["text"],"output":["text"]},"context_window":256000,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":["default"],"tokenizer":null,"capabilities":{"function_calling":true,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":true,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":true,"provider_count":1,"ids":["@cf/nvidia/nemotron-3-120b-a12b","cloudflare/@cf/nvidia/nemotron-3-120b-a12b","nvidia-nemotron-3-120b-a12b"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"},{"id":"nvidia-nemotron-super-49b","name":"nvidia-nemotron-super-49b","display_name":"Nemotron Super 49B","description":"A large language model derived from Meta Llama 3.3 70B Instruct, optimized for reasoning tasks with a 49B parameter footprint.","creator":"nvidia","family":"nemotron","tier":"","version":null,"type":"language","size_in_bn":49,"modalities":{"input":["text"],"output":["text"]},"context_window":null,"max_output_tokens":null,"tool_use_system_prompt_tokens":0,"output_vector_sizes":[],"knowledge_cutoff":null,"training_data_cutoff":null,"supported_reasoning_efforts":[],"tokenizer":null,"capabilities":{"function_calling":false,"parallel_function_calling":false,"structured_outputs":false,"prompt_caching":false,"reasoning":false,"web_search":false,"computer_use":false,"code_execution":false,"file_search":false,"url_context":false,"assistant_prefill":false,"native_structured_output":false,"adaptive_reasoning":false},"release_date":null,"earliest_deprecation_date":null,"deprecated":false,"has_pricing":false,"provider_count":0,"ids":["nvidia-nemotron-super-49b","nvidia-nemotron-super-49b-nim"],"hf_likes":null,"hf_downloads":null,"hf_downloads_all_time":null,"hf_trending_score":null,"updated_at":"2026-09-22 08:02:55"}],"pagination":{"page_size":50,"has_next":false,"next_token":null,"total_count":35},"meta":{"updated_at":"2026-09-22","request_id":"42866c34-5ef2-4ae7-b193-e93657a59d80","execution_ms":8}}