@agentjido/llmdb 2026.9.6 → 2026.9.7

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (85) hide show
  1. package/dist/generated/compact-strings.js +1 -1
  2. package/dist/generated/manifest.js +1 -1
  3. package/dist/providers/302ai.js +1 -1
  4. package/dist/providers/abacus.js +1 -1
  5. package/dist/providers/ai_router.js +1 -1
  6. package/dist/providers/aihubmix.js +1 -1
  7. package/dist/providers/amazon_bedrock.js +1 -1
  8. package/dist/providers/anthropic.js +1 -1
  9. package/dist/providers/azure.js +1 -1
  10. package/dist/providers/azure_cognitive_services.js +1 -1
  11. package/dist/providers/cline_pass.js +1 -1
  12. package/dist/providers/cloudflare_ai_gateway.js +1 -1
  13. package/dist/providers/cortecs.js +1 -1
  14. package/dist/providers/crof.js +1 -1
  15. package/dist/providers/crossmodel.js +1 -1
  16. package/dist/providers/daoxe.js +1 -1
  17. package/dist/providers/databricks.js +1 -1
  18. package/dist/providers/deepinfra.js +1 -1
  19. package/dist/providers/digitalocean.js +1 -1
  20. package/dist/providers/edenai.js +1 -1
  21. package/dist/providers/empiriolabs.js +1 -1
  22. package/dist/providers/fastrouter.js +1 -1
  23. package/dist/providers/fireworks_ai.js +1 -1
  24. package/dist/providers/freemodel.js +1 -1
  25. package/dist/providers/github_copilot.js +1 -1
  26. package/dist/providers/gitlab.js +1 -1
  27. package/dist/providers/gmicloud.js +1 -1
  28. package/dist/providers/google.js +1 -1
  29. package/dist/providers/google_vertex.js +1 -1
  30. package/dist/providers/google_vertex_anthropic.js +1 -1
  31. package/dist/providers/hpc_ai.js +1 -1
  32. package/dist/providers/huggingface.js +1 -1
  33. package/dist/providers/hyper.js +1 -1
  34. package/dist/providers/impossibl.js +1 -1
  35. package/dist/providers/inco.js +1 -1
  36. package/dist/providers/iteracompute.js +1 -1
  37. package/dist/providers/jalapeno.js +1 -1
  38. package/dist/providers/kenari.js +1 -1
  39. package/dist/providers/kilo.js +1 -1
  40. package/dist/providers/lilac.js +1 -1
  41. package/dist/providers/llmgateway.js +1 -1
  42. package/dist/providers/llmgateway_providers.js +1 -1
  43. package/dist/providers/llmtr.js +1 -1
  44. package/dist/providers/merge_gateway.js +1 -1
  45. package/dist/providers/minimax.js +1 -1
  46. package/dist/providers/model_oracle_ai.js +1 -1
  47. package/dist/providers/morph.js +1 -1
  48. package/dist/providers/nano_gpt.js +1 -1
  49. package/dist/providers/nearai.js +1 -1
  50. package/dist/providers/nebius.js +1 -1
  51. package/dist/providers/neon.js +1 -1
  52. package/dist/providers/nvidia.js +1 -1
  53. package/dist/providers/ofox.js +1 -1
  54. package/dist/providers/openai.js +1 -1
  55. package/dist/providers/opencode.js +1 -1
  56. package/dist/providers/opencode_go.js +1 -1
  57. package/dist/providers/openrouter.js +1 -1
  58. package/dist/providers/opper.js +1 -1
  59. package/dist/providers/orcarouter.js +1 -1
  60. package/dist/providers/pioneer.js +1 -1
  61. package/dist/providers/poe.js +1 -1
  62. package/dist/providers/qiniu_ai.js +1 -1
  63. package/dist/providers/requesty.js +1 -1
  64. package/dist/providers/sap_ai_core.js +1 -1
  65. package/dist/providers/scnet_token_plan.js +1 -1
  66. package/dist/providers/siliconflow.js +1 -1
  67. package/dist/providers/siliconflow_cn.js +1 -1
  68. package/dist/providers/snowflake_cortex.js +1 -1
  69. package/dist/providers/synthetic.js +1 -1
  70. package/dist/providers/tempr.js +1 -1
  71. package/dist/providers/tencent_coding_plan.js +1 -1
  72. package/dist/providers/tencent_token_plan.js +1 -1
  73. package/dist/providers/tencent_tokenhub.js +1 -1
  74. package/dist/providers/tensorx.js +1 -1
  75. package/dist/providers/unorouter.js +1 -1
  76. package/dist/providers/upstage.js +1 -1
  77. package/dist/providers/vancine.js +1 -1
  78. package/dist/providers/venice.js +1 -1
  79. package/dist/providers/vercel.js +1 -1
  80. package/dist/providers/vivgrid.js +1 -1
  81. package/dist/providers/volcengine_coding_plan.js +1 -1
  82. package/dist/providers/wafer_ai.js +1 -1
  83. package/dist/providers/xpersona.js +1 -1
  84. package/dist/providers/zenmux.js +1 -1
  85. package/package.json +1 -1
@@ -1,2 +1,2 @@
1
1
  // Generated by scripts/generate.mjs. Do not edit.
2
- export const compactStrings = ["openai_chat_compatible", "/chat/completions", "Qwen vision-language model for visual reasoning, documents, and agent tasks", "openai_chat", "reasoning_content", "Compact GPT model for low-latency assistance and high-volume workloads", "Qwen instruction model for multilingual chat, reasoning, and tool use", "Image model for prompt-driven generation, editing, and visual design workflows", "Flagship Claude model for deep reasoning, coding, and long-horizon agents", "General purpose text generation", "Fast Gemini model balancing multimodal reasoning, tool use, and cost", "Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work", "Balanced Claude model for coding, analysis, agent workflows, and cost control", "Open flagship GLM for long-horizon coding agents and million-token context work", "Open Llama instruction model for multilingual chat, reasoning, and coding", "Coding-optimized GPT model for repository edits, reviews, and agentic software work", "GPT model for general reasoning, writing, coding, and tool-assisted tasks", "Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding", "Prompt length selects the rate for all tokens in the request; the long-context threshold is inclusive. See https://docs.x.ai/developers/pricing.", "Multimodal reasoning model for visual analysis, planning, and tool use", "llmgateway_providers", "Flagship GLM model for hybrid reasoning, coding, and agentic engineering", "DeepSeek chat model for instruction following, coding, and analysis", "Open GPT reasoning model for self-hosted agents and controllable deployments", "budget_tokens", "Flagship GLM model for long-horizon coding, agents, and complex project delivery", "Open Gemma instruction model for efficient chat and self-hosted deployments", "Efficient model for low-latency assistance, extraction, and routine automation", "Native multimodal GLM model for efficient coding and long-horizon agent tasks", "Multimodal, vision and text", "Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking", "Embedding model for semantic search, retrieval, clustering, and ranking pipelines", "Fast DeepSeek V4 lane for economical reasoning, coding, and long-context work", "Open MoE flagship with million-token context for coding and long agent runs", "Qwen coding model for software agents, repository edits, and code reasoning", "text-generation", "DeepSeek V4 Pro snapshot with million-token context and support for thinking and non-thinking modes", "Speech generation model for controllable voice, narration, and audio delivery", "Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents", "MiniMax model for chat, coding, office work, and agentic tasks", "Multimodal Kimi workhorse for agent loops, coding tasks, and visual context", "Open-weight GPT model for self-hosted reasoning and instruction-following workloads", "2.4-trillion-parameter MoE flagship for coding, professional work, multimodal understanding, and long-horizon agentic workflows", "Grok model for agentic tool use, reasoning, coding, and live assistance", "Strong GLM coding model for agentic engineering, terminals, and repository generation", "Fast Claude model for responsive assistance, classification, and lightweight agents", "Low-latency Gemini model for high-volume multimodal and agent workloads", "effort", "Kimi multimodal agent model for visual understanding, coding, and planning", "Legacy model retained for compatibility with older integrations", "Claude workhorse for coding agents, careful analysis, and production cost control", "xAI's frontier model for long-running agents, coding, knowledge work, and visual projects", "Frontier GPT-5.6 model for complex professional work, coding, and agentic workflows", "MiniMax multimodal model for long-context coding, perception, and agent planning", "Everyday Claude agent model for coding, planning, browsing, and general work", "openai_responses_compatible", "Claude model for creative writing, analysis, and controlled agent workflows", "Mistral model for multilingual chat, reasoning, and tool-assisted workflows", "DeepSeek V4.1 Flash model for reasoning and agentic coding", "Reasoning model for deliberate analysis, multi-step problem solving, and tool use", "Advanced Gemini model for complex reasoning, coding, and multimodal analysis", "O-series reasoning model for hard analysis, math, coding, and planning", "image-text-to-text", "Flagship model for demanding analysis, coding, and production agent workflows", "Tencent Hy reasoning model for coding, instruction following, and agent tasks", "Default frontier GPT for coding, computer use, research, and knowledge work", "google_generate_content", "openrouter", "openai_responses", "Video model for prompt-guided generation, editing, and motion workflows", "/models/{provider_model_id}:generateContent", "medium", "Stronger MiMo Pro tier for multimodal reasoning and coding-agent execution", "Open-weight sparse MoE (2.4T total, 95B active), the open-weight twin of Qwen3.8 Max for coding, research, complex reasoning, and agentic workflows", "GLM vision model for visual reasoning, documents, and multimodal agents", "Google's most intelligent Flash model, engineered for long-horizon software engineering, autonomous agents, and complex enterprise workflows", "Stronger Opus tier for advanced software work and high-stakes reasoning", "Open-weight instruction model for adaptable chat and self-hosted production workloads", "Largest Gemma 4 instruction model for open, self-hosted chat and reasoning", "Muse Glimmer is a 30-billion-parameter open-weight multimodal model from Meta Superintelligence Labs, distilled from Muse Spark for always-on local agents, tool use, coding, and image understanding.", "High-end Claude for difficult coding, planning, and slower expert reasoning", "GPT-6 Astra is OpenAI's most capable model for complex reasoning, coding, computer use, research, and document creation.", "Frontier GPT model for professional reasoning, coding, and multimodal work", "Open MiniMax flagship for coding agents, office automation, and complex environments", "Fast Grok model for responsive chat, reasoning, and tool-assisted work", "DeepSeek reasoning model for multi-step analysis, math, coding, and tools", "Strongest Claude Opus model for coding, agents, and professional work", "Muse Spark 1.2 is a coding-focused update to Muse Spark 1.1 with improvements in code generation, complex debugging, codebase understanding, and end-to-end developer workflows.", "Qwen frontier model tuned for agent frameworks, coding assistants, and long tasks", "Balanced GPT-5.6 model for capable, cost-efficient everyday work", "Dense 27B vision-language model for coding, agent tasks, and image and video understanding", "General GLM flagship for coding, analysis, and tool-heavy engineering workflows", "Agent-ready GPT for coding and computer-use workflows at a lower cost", "Cost-efficient GPT-5.6 model for fast, high-volume workloads", "Efficient Mistral model for fast chat, extraction, and production assistants", "Muse Spark 1.3 is a multimodal reasoning model from Meta for long-running agentic, multi-agent, and coding workflows. It improves long-horizon agent collaboration, instruction following, and coding efficiency relative to Muse Spark 1.2.", "Open multimodal Qwen MoE for local agents that need vision, audio, and code", "Chat-tuned GPT model for conversational assistance, writing, and tool workflows", "Compact Mistral model for edge, latency-sensitive, and cost-efficient workloads", "Safety model for policy screening, moderation, and risk-aware routing workflows", "High-efficiency Gemini model for agentic workflows, coding, and multimodal reasoning", "Flagship Qwen model for complex reasoning, coding, and agentic workflows", "mradermacher", "Efficient GLM model for fast reasoning, coding, and agent workflows", "Strong small GPT for coding subagents, quick tool use, and high-volume work", "Multimodal Qwen workhorse for long-context agents, visual inputs, and coding", "deepseek-thinking", "Qwen reasoning model for deliberate problem solving, math, and coding", "claude-opus", "deepseek-flash", "Reasoning-first Gemini preview for agentic coding and complex problem solving", "Earlier Kimi frontier model for long-context agents, coding, and multimodal work", "Open MiMo model for multimodal coding agents and long-context automation", "https://${GOOGLE_VERTEX_ENDPOINT}/v1/projects/${GOOGLE_VERTEX_PROJECT}/locations/${GOOGLE_VERTEX_LOCATION}/endpoints/openapi", "provider_docs", "tool.web_search", "amazon_bedrock", "nano_gpt", "Multimodal MoE reasoning model (975B total, 41B active) for text, image, and audio", "General-purpose chat model for instruction following, writing, and analysis", "Claude model for long-running agentic coding and knowledge work", "Mature GLM model for dependable coding, reasoning, and structured agent tasks", "Fast Claude lane for lightweight agents, office tasks, and responsive chat", "2025-08-05", "Claude model for demanding reasoning and long-horizon agentic work", "Popular open Llama workhorse for multilingual chat, coding, and self-hosting", "anthropic_messages", "Long-lived GPT workhorse for coding, instruction following, and production apps", "Cheapest GPT-5.4 lane for simple routing, extraction, and bulk automation", "2026-07-09", "Kimi model for long-context chat, coding, and agentic reasoning", "Speech transcription model for accurate audio-to-text and captioning workflows", "Nemotron middle tier for collaborative agents and high-volume reasoning workloads", "xAI's default Grok for chat, coding, agentic tools, and lower hallucination risk", "claude-sonnet", "Large open Qwen multimodal MoE for visual agents and long technical tasks", "gemini-flash", "Reliable GPT generation for broad coding, writing, and tool-assisted product work", "2026-04-24", "2025-08-31", "Prior MiniMax coding model for agent workflows, office edits, and automation", "deprecated", "New Gemini flash lane bringing frontier-style multimodal reasoning to cheaper runs", "Original GPT-5 workhorse for reasoning, coding, writing, and tool workflows", "xAI's Grok model for chat, coding, agentic tools, and lower hallucination risk", "merge_gateway", "Small GPT-5 for responsive agents, coding help, and everyday automation", "2026-09-22", "Largest Nemotron 3 model for maximum open-weight reasoning and agent accuracy", "Affordable GPT-4.1 lane for fast coding help and structured extraction", "Flagship Mistral model for advanced reasoning, coding, and multilingual work", "Credit consumption for the current token-based plan; select a confirmed pricing period. Legacy plan quotas and USD API charges are separate.", "azure_cognitive_services", "Coding model for repository understanding, refactors, and agentic engineering tasks", "Newer StepFun flash model for faster agents, coding, and multimodal prompts", "Automatic model router for matching prompts to suitable backends and budgets", "StepFun flash model for efficient multimodal reasoning, coding, and tool use", "Small omni GPT for cheap multimodal assistance and production-scale traffic", "Tiny GPT-4.1 option for classification, routing, and very high-volume tasks", "Instruction following, chat", "Sharper GPT-5 generation for coding, product work, and tool-assisted tasks", "OpenAI's most efficient model for focused, high-volume tasks", "Multimodal model for analyzing text, images, documents, and rich media", "merge_by_id", "codecategories", "Earlier Qwen multimodal workhorse for million-token agent and document tasks", "toggle", "2026-04-22", "Tool-capable chat model for instruction following and agentic application workflows", "full_request", "minimal", "gemini-flash-lite", "Tiny GPT-5 lane for routing, extraction, classification, and bulk jobs", "2026-06-13", "openai/gpt-oss-120b", "Fast Gemini workhorse for multimodal apps where latency and price matter", "@ai-sdk/anthropic", "Omni-era GPT for multimodal chat, practical coding, and general assistants", "Highest-accuracy GPT-5.5 tier for slower, precision-heavy reasoning and coding", "High-speed MiniMax model for low-latency coding and agent workflows", "Budget GLM lane for fast coding help, routing, and everyday automation", "Hybrid-reasoning DeepSeek model with thinking and non-thinking modes, sparse attention, and tool-use", "Google's proven reasoning model for coding, math, and multimodal analysis", "Nemotron model for efficient reasoning, coding, and specialized AI agents", "Quality-first multi-agent model for hard research, analysis, and competitions", "2026-08-14", "Experimental multimodal DeepSeek V4 Flash model for image understanding, coding, and agentic work", "Deliberate o-series reasoner for hard math, coding, and multi-step analysis", "Qwen omni model for text, vision, audio, and multimodal agent tasks", "2026-08-26", "Muse Spark is a natively multimodal reasoning model with support for tool-use, visual chain of thought, and multi-agent orchestration.", "2026-04-02", "OpenAI model for complex coding and agentic workflows", "2025-01", "models", "@ai-sdk/openai-compatible", "2026-02-16", "Late GLM-4 workhorse for coding agents, reasoning, and structured tasks", "2025-04", "Balanced Mistral model for enterprise assistants, multilingual work, and tools", "Efficient Qwen model for fast chat, extraction, and high-volume workloads", "unsloth", "openai_images", "More exact GPT-5.4 tier for demanding professional reasoning and agent tasks", "Flagship Qwen3 model for coding agents, complex reasoning, and tool use", "@ai-sdk/openai", "Fast o-series model for compact reasoning, coding, and tool use", "Flagship DeepSeek model for coding, reasoning, and agentic work", "Kimi reasoning model for long-horizon research, planning, and tool use", "https://${AZURE_COGNITIVE_SERVICES_RESOURCE_NAME}.services.ai.azure.com/anthropic/v1", "llmgateway", "MiMo Flash model for multimodal coding agents and long-context automation", "Smaller o-series reasoner for economical coding, math, and planning tasks", "Fast DeepSeek model for efficient chat, coding help, and agent loops", "Fast Grok coding model tuned for agentic engineering and iterative edits", "Open multimodal Llama model for strong reasoning and fast responses", "Lean Gemini 2.5 lane for cheap multimodal traffic and quick agents", "2025-08-07", "Open Llama multimodal model for image understanding and text reasoning", "A next-generation productivity model with significantly enhanced Agent and complex task execution capabilities.", "openai_embeddings", "/images/generations", "Small Nemotron 3 MoE for efficient coding, math, and long-context agents", "2026-07-16", "StepFun's next-generation flagship base model for coding and professional knowledge work, with native text, image, and video input and a 1M-token context window", "2026-04-21", "2026-06-12", "Classic open reasoning model for transparent math, coding, and deliberate problem solving", "Fast GLM vision model for screenshots, documents, and multimodal agent tasks", "2025-04-14", "web_search", "2025-12-01", "2026-02-12", "2026-03-18", "uicomponent", "2026-04-16", "Fast Mistral production model for chat, extraction, and cost-sensitive agents", "Hybrid-reasoning GLM release that made the 4.5 line broadly useful", "Lower-latency Kimi Code variant for interactive edits and coding-agent loops", "2026-08-12", "deepseek-ai/DeepSeek-V4-Flash-0731", "Qwen/Qwen3-235B-A22B-Instruct-2507", "2026-05-28", "Mistral's largest general model for enterprise agents, coding, and multilingual reasoning", "Nano Banana Pro for higher-fidelity image generation and design-heavy edits", "Lighter GLM-4.5 variant for fast coding assistance and cheaper agents", "2026-09-10", "Fast NVIDIA Nemotron MoE for reliable agentic tasks across enterprise workloads", "openai_completion", "2026-01-31", "cloudflare_ai_gateway", "deepseek-ai/DeepSeek-V4-Pro", "Earlier MiniMax agent model for practical coding and productivity tasks", "http_request", "Mistral's coding-agent model for repository work, terminal tasks, and software fixes", "Reranking model for improving retrieval quality in search and recommendation systems"];
2
+ export const compactStrings = ["openai_chat_compatible", "/chat/completions", "Qwen vision-language model for visual reasoning, documents, and agent tasks", "openai_chat", "reasoning_content", "Compact GPT model for low-latency assistance and high-volume workloads", "Qwen instruction model for multilingual chat, reasoning, and tool use", "Image model for prompt-driven generation, editing, and visual design workflows", "Flagship Claude model for deep reasoning, coding, and long-horizon agents", "General purpose text generation", "Fast Gemini model balancing multimodal reasoning, tool use, and cost", "Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work", "Balanced Claude model for coding, analysis, agent workflows, and cost control", "Open flagship GLM for long-horizon coding agents and million-token context work", "Open Llama instruction model for multilingual chat, reasoning, and coding", "Coding-optimized GPT model for repository edits, reviews, and agentic software work", "GPT model for general reasoning, writing, coding, and tool-assisted tasks", "Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding", "Prompt length selects the rate for all tokens in the request; the long-context threshold is inclusive. See https://docs.x.ai/developers/pricing.", "Multimodal reasoning model for visual analysis, planning, and tool use", "llmgateway_providers", "Flagship GLM model for hybrid reasoning, coding, and agentic engineering", "DeepSeek chat model for instruction following, coding, and analysis", "Open GPT reasoning model for self-hosted agents and controllable deployments", "budget_tokens", "Flagship GLM model for long-horizon coding, agents, and complex project delivery", "Open Gemma instruction model for efficient chat and self-hosted deployments", "Efficient model for low-latency assistance, extraction, and routine automation", "Native multimodal GLM model for efficient coding and long-horizon agent tasks", "Multimodal, vision and text", "Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking", "Embedding model for semantic search, retrieval, clustering, and ranking pipelines", "Fast DeepSeek V4 lane for economical reasoning, coding, and long-context work", "Open MoE flagship with million-token context for coding and long agent runs", "Qwen coding model for software agents, repository edits, and code reasoning", "text-generation", "DeepSeek V4 Pro snapshot with million-token context and support for thinking and non-thinking modes", "Speech generation model for controllable voice, narration, and audio delivery", "Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents", "MiniMax model for chat, coding, office work, and agentic tasks", "Multimodal Kimi workhorse for agent loops, coding tasks, and visual context", "Open-weight GPT model for self-hosted reasoning and instruction-following workloads", "2.4-trillion-parameter MoE flagship for coding, professional work, multimodal understanding, and long-horizon agentic workflows", "Grok model for agentic tool use, reasoning, coding, and live assistance", "Strong GLM coding model for agentic engineering, terminals, and repository generation", "Fast Claude model for responsive assistance, classification, and lightweight agents", "Low-latency Gemini model for high-volume multimodal and agent workloads", "effort", "Kimi multimodal agent model for visual understanding, coding, and planning", "Legacy model retained for compatibility with older integrations", "Claude workhorse for coding agents, careful analysis, and production cost control", "xAI's frontier model for long-running agents, coding, knowledge work, and visual projects", "Frontier GPT-5.6 model for complex professional work, coding, and agentic workflows", "openai_responses_compatible", "MiniMax multimodal model for long-context coding, perception, and agent planning", "Everyday Claude agent model for coding, planning, browsing, and general work", "Claude model for creative writing, analysis, and controlled agent workflows", "Mistral model for multilingual chat, reasoning, and tool-assisted workflows", "DeepSeek V4.1 Flash model for reasoning and agentic coding", "Reasoning model for deliberate analysis, multi-step problem solving, and tool use", "Advanced Gemini model for complex reasoning, coding, and multimodal analysis", "O-series reasoning model for hard analysis, math, coding, and planning", "image-text-to-text", "openai_responses", "Flagship model for demanding analysis, coding, and production agent workflows", "Tencent Hy reasoning model for coding, instruction following, and agent tasks", "Default frontier GPT for coding, computer use, research, and knowledge work", "google_generate_content", "openrouter", "Video model for prompt-guided generation, editing, and motion workflows", "/models/{provider_model_id}:generateContent", "medium", "Stronger MiMo Pro tier for multimodal reasoning and coding-agent execution", "Open-weight sparse MoE (2.4T total, 95B active), the open-weight twin of Qwen3.8 Max for coding, research, complex reasoning, and agentic workflows", "GLM vision model for visual reasoning, documents, and multimodal agents", "Google's most intelligent Flash model, engineered for long-horizon software engineering, autonomous agents, and complex enterprise workflows", "Stronger Opus tier for advanced software work and high-stakes reasoning", "Open-weight instruction model for adaptable chat and self-hosted production workloads", "Largest Gemma 4 instruction model for open, self-hosted chat and reasoning", "Muse Glimmer is a 30-billion-parameter open-weight multimodal model from Meta Superintelligence Labs, distilled from Muse Spark for always-on local agents, tool use, coding, and image understanding.", "High-end Claude for difficult coding, planning, and slower expert reasoning", "GPT-6 Astra is OpenAI's most capable model for complex reasoning, coding, computer use, research, and document creation.", "Frontier GPT model for professional reasoning, coding, and multimodal work", "Open MiniMax flagship for coding agents, office automation, and complex environments", "Fast Grok model for responsive chat, reasoning, and tool-assisted work", "DeepSeek reasoning model for multi-step analysis, math, coding, and tools", "Strongest Claude Opus model for coding, agents, and professional work", "Muse Spark 1.2 is a coding-focused update to Muse Spark 1.1 with improvements in code generation, complex debugging, codebase understanding, and end-to-end developer workflows.", "Qwen frontier model tuned for agent frameworks, coding assistants, and long tasks", "Balanced GPT-5.6 model for capable, cost-efficient everyday work", "Dense 27B vision-language model for coding, agent tasks, and image and video understanding", "General GLM flagship for coding, analysis, and tool-heavy engineering workflows", "Agent-ready GPT for coding and computer-use workflows at a lower cost", "Cost-efficient GPT-5.6 model for fast, high-volume workloads", "Efficient Mistral model for fast chat, extraction, and production assistants", "Muse Spark 1.3 is a multimodal reasoning model from Meta for long-running agentic, multi-agent, and coding workflows. It improves long-horizon agent collaboration, instruction following, and coding efficiency relative to Muse Spark 1.2.", "Open multimodal Qwen MoE for local agents that need vision, audio, and code", "Chat-tuned GPT model for conversational assistance, writing, and tool workflows", "Compact Mistral model for edge, latency-sensitive, and cost-efficient workloads", "Safety model for policy screening, moderation, and risk-aware routing workflows", "High-efficiency Gemini model for agentic workflows, coding, and multimodal reasoning", "Flagship Qwen model for complex reasoning, coding, and agentic workflows", "mradermacher", "Efficient GLM model for fast reasoning, coding, and agent workflows", "Strong small GPT for coding subagents, quick tool use, and high-volume work", "Multimodal Qwen workhorse for long-context agents, visual inputs, and coding", "deepseek-thinking", "Qwen reasoning model for deliberate problem solving, math, and coding", "claude-opus", "deepseek-flash", "Reasoning-first Gemini preview for agentic coding and complex problem solving", "Earlier Kimi frontier model for long-context agents, coding, and multimodal work", "Open MiMo model for multimodal coding agents and long-context automation", "https://${GOOGLE_VERTEX_ENDPOINT}/v1/projects/${GOOGLE_VERTEX_PROJECT}/locations/${GOOGLE_VERTEX_LOCATION}/endpoints/openapi", "provider_docs", "tool.web_search", "amazon_bedrock", "nano_gpt", "Multimodal MoE reasoning model (975B total, 41B active) for text, image, and audio", "General-purpose chat model for instruction following, writing, and analysis", "Claude model for long-running agentic coding and knowledge work", "Mature GLM model for dependable coding, reasoning, and structured agent tasks", "Fast Claude lane for lightweight agents, office tasks, and responsive chat", "2025-08-05", "Claude model for demanding reasoning and long-horizon agentic work", "Popular open Llama workhorse for multilingual chat, coding, and self-hosting", "anthropic_messages", "Long-lived GPT workhorse for coding, instruction following, and production apps", "Cheapest GPT-5.4 lane for simple routing, extraction, and bulk automation", "2026-07-09", "Kimi model for long-context chat, coding, and agentic reasoning", "Speech transcription model for accurate audio-to-text and captioning workflows", "Nemotron middle tier for collaborative agents and high-volume reasoning workloads", "xAI's default Grok for chat, coding, agentic tools, and lower hallucination risk", "claude-sonnet", "Large open Qwen multimodal MoE for visual agents and long technical tasks", "gemini-flash", "Reliable GPT generation for broad coding, writing, and tool-assisted product work", "2026-04-24", "2025-08-31", "Prior MiniMax coding model for agent workflows, office edits, and automation", "deprecated", "New Gemini flash lane bringing frontier-style multimodal reasoning to cheaper runs", "Original GPT-5 workhorse for reasoning, coding, writing, and tool workflows", "xAI's Grok model for chat, coding, agentic tools, and lower hallucination risk", "merge_gateway", "Small GPT-5 for responsive agents, coding help, and everyday automation", "2026-09-22", "Largest Nemotron 3 model for maximum open-weight reasoning and agent accuracy", "Affordable GPT-4.1 lane for fast coding help and structured extraction", "Flagship Mistral model for advanced reasoning, coding, and multilingual work", "Credit consumption for the current token-based plan; select a confirmed pricing period. Legacy plan quotas and USD API charges are separate.", "azure_cognitive_services", "Coding model for repository understanding, refactors, and agentic engineering tasks", "Newer StepFun flash model for faster agents, coding, and multimodal prompts", "Automatic model router for matching prompts to suitable backends and budgets", "StepFun flash model for efficient multimodal reasoning, coding, and tool use", "Small omni GPT for cheap multimodal assistance and production-scale traffic", "Tiny GPT-4.1 option for classification, routing, and very high-volume tasks", "Instruction following, chat", "Sharper GPT-5 generation for coding, product work, and tool-assisted tasks", "OpenAI's most efficient model for focused, high-volume tasks", "Multimodal model for analyzing text, images, documents, and rich media", "merge_by_id", "codecategories", "Earlier Qwen multimodal workhorse for million-token agent and document tasks", "toggle", "2026-04-22", "Tool-capable chat model for instruction following and agentic application workflows", "full_request", "minimal", "gemini-flash-lite", "Tiny GPT-5 lane for routing, extraction, classification, and bulk jobs", "2026-06-13", "openai/gpt-oss-120b", "Fast Gemini workhorse for multimodal apps where latency and price matter", "@ai-sdk/anthropic", "Omni-era GPT for multimodal chat, practical coding, and general assistants", "Highest-accuracy GPT-5.5 tier for slower, precision-heavy reasoning and coding", "High-speed MiniMax model for low-latency coding and agent workflows", "Budget GLM lane for fast coding help, routing, and everyday automation", "Hybrid-reasoning DeepSeek model with thinking and non-thinking modes, sparse attention, and tool-use", "Google's proven reasoning model for coding, math, and multimodal analysis", "Nemotron model for efficient reasoning, coding, and specialized AI agents", "Quality-first multi-agent model for hard research, analysis, and competitions", "2026-08-14", "Experimental multimodal DeepSeek V4 Flash model for image understanding, coding, and agentic work", "Deliberate o-series reasoner for hard math, coding, and multi-step analysis", "Qwen omni model for text, vision, audio, and multimodal agent tasks", "2026-08-26", "Muse Spark is a natively multimodal reasoning model with support for tool-use, visual chain of thought, and multi-agent orchestration.", "2026-04-02", "OpenAI model for complex coding and agentic workflows", "2025-01", "models", "@ai-sdk/openai-compatible", "2026-02-16", "Late GLM-4 workhorse for coding agents, reasoning, and structured tasks", "2025-04", "Balanced Mistral model for enterprise assistants, multilingual work, and tools", "Efficient Qwen model for fast chat, extraction, and high-volume workloads", "unsloth", "openai_images", "More exact GPT-5.4 tier for demanding professional reasoning and agent tasks", "Flagship Qwen3 model for coding agents, complex reasoning, and tool use", "@ai-sdk/openai", "Fast o-series model for compact reasoning, coding, and tool use", "Flagship DeepSeek model for coding, reasoning, and agentic work", "Kimi reasoning model for long-horizon research, planning, and tool use", "https://${AZURE_COGNITIVE_SERVICES_RESOURCE_NAME}.services.ai.azure.com/anthropic/v1", "llmgateway", "MiMo Flash model for multimodal coding agents and long-context automation", "Smaller o-series reasoner for economical coding, math, and planning tasks", "Fast DeepSeek model for efficient chat, coding help, and agent loops", "Fast Grok coding model tuned for agentic engineering and iterative edits", "Open multimodal Llama model for strong reasoning and fast responses", "Lean Gemini 2.5 lane for cheap multimodal traffic and quick agents", "2025-08-07", "Open Llama multimodal model for image understanding and text reasoning", "A next-generation productivity model with significantly enhanced Agent and complex task execution capabilities.", "openai_embeddings", "/images/generations", "Small Nemotron 3 MoE for efficient coding, math, and long-context agents", "2026-07-16", "StepFun's next-generation flagship base model for coding and professional knowledge work, with native text, image, and video input and a 1M-token context window", "2026-04-21", "2026-06-12", "Classic open reasoning model for transparent math, coding, and deliberate problem solving", "Fast GLM vision model for screenshots, documents, and multimodal agent tasks", "2025-04-14", "web_search", "2025-12-01", "2026-02-12", "2026-03-18", "uicomponent", "2026-04-16", "Fast Mistral production model for chat, extraction, and cost-sensitive agents", "Hybrid-reasoning GLM release that made the 4.5 line broadly useful", "Lower-latency Kimi Code variant for interactive edits and coding-agent loops", "2026-08-12", "deepseek-ai/DeepSeek-V4-Flash-0731", "Qwen/Qwen3-235B-A22B-Instruct-2507", "2026-05-28", "Mistral's largest general model for enterprise agents, coding, and multilingual reasoning", "Nano Banana Pro for higher-fidelity image generation and design-heavy edits", "Lighter GLM-4.5 variant for fast coding assistance and cheaper agents", "2026-09-10", "Fast NVIDIA Nemotron MoE for reliable agentic tasks across enterprise workloads", "openai_completion", "2026-01-31", "cloudflare_ai_gateway", "deepseek-ai/DeepSeek-V4-Pro", "Earlier MiniMax agent model for practical coding and productivity tasks", "http_request", "Mistral's coding-agent model for repository work, terminal tasks, and software fixes", "Reranking model for improving retrieval quality in search and recommendation systems"];
@@ -1 +1 @@
1
- export const manifest = { "catalog_version": 2, "format_version": 1, "generated_at": "2026-09-25T12:55:24.551887Z", "model_count": 8800, "provider_count": 228, "providers": { "302ai": { "alias_of": null, "base_url": "https://api.302.ai/v1", "catalog_only": true, "doc": "https://doc.302.ai", "id": "302ai", "model_count": 117, "name": "302.AI" }, "a2agent": { "alias_of": null, "base_url": "https://api.a2agent.me/v1", "catalog_only": false, "doc": "https://docs.a2agent.me", "id": "a2agent", "model_count": 20, "name": "A2Agent" }, "abacus": { "alias_of": null, "base_url": "https://routellm.abacus.ai/v1", "catalog_only": true, "doc": "https://abacus.ai/help/api", "id": "abacus", "model_count": 108, "name": "Abacus" }, "abliteration_ai": { "alias_of": null, "base_url": "https://api.abliteration.ai/v1", "catalog_only": true, "doc": "https://docs.abliteration.ai/models", "id": "abliteration_ai", "model_count": 3, "name": "abliteration.ai" }, "above": { "alias_of": null, "base_url": "https://api.above.dev/v1", "catalog_only": true, "doc": "https://above.dev/docs", "id": "above", "model_count": 9, "name": "above.dev" }, "agentrouter": { "alias_of": null, "base_url": "https://agentrouter.org/v1", "catalog_only": true, "doc": "https://agentrouter.org/docs/opencode.html", "id": "agentrouter", "model_count": 5, "name": "AgentRouter" }, "agnes": { "alias_of": null, "base_url": "https://apihub.agnes-ai.com/v1", "catalog_only": true, "doc": "https://agnes-ai.com/doc", "id": "agnes", "model_count": 3, "name": "Agnes AI" }, "ai21": { "alias_of": null, "base_url": "https://api.ai21.com/studio/v1", "catalog_only": true, "doc": "https://docs.ai21.com/docs/jamba-foundation-models", "id": "ai21", "model_count": 2, "name": "AI21 Labs" }, "ai_router": { "alias_of": null, "base_url": "https://api.ai-router.dev/v1", "catalog_only": true, "doc": "https://ai-router.dev/openai-compatible-api-gateway/", "id": "ai_router", "model_count": 5, "name": "AI-ROUTER" }, "aiand": { "alias_of": null, "base_url": "https://api.aiand.com/v1", "catalog_only": true, "doc": "https://docs.aiand.com/", "id": "aiand", "model_count": 11, "name": "ai&" }, "aihubmix": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.aihubmix.com", "id": "aihubmix", "model_count": 106, "name": "AIHubMix" }, "ainetcafe": { "alias_of": null, "base_url": "https://microquickjs.com/v1", "catalog_only": true, "doc": "https://ainetcafe.com/k3/guides/", "id": "ainetcafe", "model_count": 1, "name": "ainetcafe" }, "aixy": { "alias_of": null, "base_url": "https://api.aixy-gateway.com/v1", "catalog_only": true, "doc": "https://docs.aixy-gateway.com/integrations/overview", "id": "aixy", "model_count": 1, "name": "Aixy" }, "aki_io": { "alias_of": null, "base_url": "https://aki.io/v1", "catalog_only": true, "doc": "https://aki.io/docs/", "id": "aki_io", "model_count": 7, "name": "AKI.IO" }, "alibaba": { "alias_of": null, "base_url": "https://dashscope-intl.aliyuncs.com/compatible-mode/v1", "catalog_only": false, "doc": "https://www.alibabacloud.com/help/en/model-studio/models", "id": "alibaba", "model_count": 56, "name": "Alibaba" }, "alibaba_cn": { "alias_of": null, "base_url": "https://dashscope.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/models", "id": "alibaba_cn", "model_count": 90, "name": "Alibaba (China)" }, "alibaba_coding_plan": { "alias_of": null, "base_url": "https://coding-intl.dashscope.aliyuncs.com/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/coding-plan", "id": "alibaba_coding_plan", "model_count": 12, "name": "Alibaba Coding Plan" }, "alibaba_coding_plan_cn": { "alias_of": null, "base_url": "https://coding.dashscope.aliyuncs.com/v1", "catalog_only": true, "doc": "https://help.aliyun.com/zh/model-studio/coding-plan", "id": "alibaba_coding_plan_cn", "model_count": 12, "name": "Alibaba Coding Plan (China)" }, "alibaba_token_plan": { "alias_of": null, "base_url": "https://token-plan.ap-southeast-1.maas.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/token-plan-overview", "id": "alibaba_token_plan", "model_count": 28, "name": "Alibaba Token Plan" }, "alibaba_token_plan_cn": { "alias_of": null, "base_url": "https://token-plan.cn-beijing.maas.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/zh/model-studio/token-plan-overview", "id": "alibaba_token_plan_cn", "model_count": 28, "name": "Alibaba Token Plan (China)" }, "amazon_bedrock": { "alias_of": null, "base_url": "https://bedrock-runtime.{region}.amazonaws.com", "catalog_only": true, "doc": "https://docs.aws.amazon.com/bedrock/", "id": "amazon_bedrock", "model_count": 199, "name": "Amazon Bedrock" }, "ambient": { "alias_of": null, "base_url": "https://api.ambient.xyz/v1", "catalog_only": true, "doc": "https://ambient.xyz", "id": "ambient", "model_count": 10, "name": "Ambient" }, "amd": { "alias_of": null, "base_url": "https://developer.amd.com.cn/radeon/api/v1", "catalog_only": true, "doc": "https://developer.amd.com.cn/radeon/tokenfactory", "id": "amd", "model_count": 6, "name": "AMD" }, "anthropic": { "alias_of": null, "base_url": "https://api.anthropic.com", "catalog_only": false, "doc": "https://docs.anthropic.com", "id": "anthropic", "model_count": 18, "name": "Anthropic" }, "anyapi": { "alias_of": null, "base_url": "https://api.anyapi.ai/v1", "catalog_only": true, "doc": "https://docs.anyapi.ai", "id": "anyapi", "model_count": 30, "name": "AnyAPI" }, "arcee": { "alias_of": null, "base_url": "https://api.arcee.ai/api/v1", "catalog_only": true, "doc": "https://docs.arcee.ai", "id": "arcee", "model_count": 7, "name": "Arcee" }, "atomic_chat": { "alias_of": null, "base_url": "http://127.0.0.1:1337/v1", "catalog_only": true, "doc": "https://atomic.chat", "id": "atomic_chat", "model_count": 5, "name": "Atomic Chat" }, "auriko": { "alias_of": null, "base_url": "https://api.auriko.ai/v1", "catalog_only": true, "doc": "https://docs.auriko.ai", "id": "auriko", "model_count": 15, "name": "Auriko" }, "azure": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://learn.microsoft.com/en-us/azure/ai-services/openai/concepts/models", "id": "azure", "model_count": 118, "name": "Azure" }, "azure_cognitive_services": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://learn.microsoft.com/en-us/azure/ai-services/openai/concepts/models", "id": "azure_cognitive_services", "model_count": 78, "name": "Azure Cognitive Services" }, "bailing": { "alias_of": null, "base_url": "https://api.tbox.cn/api/llm/v1/chat/completions", "catalog_only": true, "doc": "https://alipaytbox.yuque.com/sxs0ba/ling/intro", "id": "bailing", "model_count": 2, "name": "Bailing" }, "baseten": { "alias_of": null, "base_url": "https://inference.baseten.co/v1", "catalog_only": true, "doc": "https://docs.baseten.co/inference/model-apis/overview", "id": "baseten", "model_count": 23, "name": "Baseten" }, "berget": { "alias_of": null, "base_url": "https://api.berget.ai/v1", "catalog_only": true, "doc": "https://api.berget.ai", "id": "berget", "model_count": 6, "name": "Berget.AI" }, "blueclaw": { "alias_of": null, "base_url": "https://openai.blueclaw.network/v1", "catalog_only": true, "doc": "https://blueclaw.network", "id": "blueclaw", "model_count": 2, "name": "Blue Claw" }, "bothub": { "alias_of": null, "base_url": "https://openai.bothub.ru/v1", "catalog_only": true, "doc": "https://bothub.ru/models", "id": "bothub", "model_count": 8, "name": "Bothub" }, "cerebras": { "alias_of": null, "base_url": "https://api.cerebras.ai/v1", "catalog_only": false, "doc": "https://cerebras.ai/docs", "id": "cerebras", "model_count": 4, "name": "Cerebras" }, "chutes": { "alias_of": null, "base_url": "https://llm.chutes.ai/v1", "catalog_only": true, "doc": "https://llm.chutes.ai/v1/models", "id": "chutes", "model_count": 14, "name": "Chutes" }, "clarifai": { "alias_of": null, "base_url": "https://api.clarifai.com/v2/ext/openai/v1", "catalog_only": true, "doc": "https://docs.clarifai.com/compute/inference/", "id": "clarifai", "model_count": 12, "name": "Clarifai" }, "claudinio": { "alias_of": null, "base_url": "https://api.claudin.io/v1", "catalog_only": true, "doc": "https://claudin.io", "id": "claudinio", "model_count": 2, "name": "Claudinio" }, "cline_pass": { "alias_of": null, "base_url": "https://api.cline.bot/api/v1", "catalog_only": true, "doc": "https://docs.cline.bot/getting-started/clinepass", "id": "cline_pass", "model_count": 18, "name": "ClinePass" }, "cloudferro_sherlock": { "alias_of": null, "base_url": "https://api-sherlock.cloudferro.com/openai/v1", "catalog_only": true, "doc": "https://docs.sherlock.cloudferro.com/", "id": "cloudferro_sherlock", "model_count": 5, "name": "CloudFerro Sherlock" }, "cloudflare_ai_gateway": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://developers.cloudflare.com/ai-gateway/", "id": "cloudflare_ai_gateway", "model_count": 50, "name": "Cloudflare AI Gateway" }, "cloudflare_workers_ai": { "alias_of": null, "base_url": "https://api.cloudflare.com/client/v4/accounts/${CLOUDFLARE_ACCOUNT_ID}/ai/v1", "catalog_only": false, "doc": "https://developers.cloudflare.com/workers-ai/models/", "id": "cloudflare_workers_ai", "model_count": 28, "name": "Cloudflare Workers AI" }, "cohere": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.cohere.com/docs/models", "id": "cohere", "model_count": 19, "name": "Cohere" }, "coralbricks": { "alias_of": null, "base_url": "https://inference.coralbricks.ai/v1", "catalog_only": true, "doc": "https://www.coralbricks.ai/docs", "id": "coralbricks", "model_count": 4, "name": "CoralBricks" }, "cortecs": { "alias_of": null, "base_url": "https://api.cortecs.ai/v1", "catalog_only": true, "doc": "https://api.cortecs.ai/v1/models", "id": "cortecs", "model_count": 109, "name": "Cortecs" }, "crof": { "alias_of": null, "base_url": "https://crof.ai/v1", "catalog_only": true, "doc": "https://crof.ai/docs", "id": "crof", "model_count": 24, "name": "CrofAI" }, "crossmodel": { "alias_of": null, "base_url": "https://api.crossmodel.ai/v1", "catalog_only": true, "doc": "https://www.crossmodel.ai/docs", "id": "crossmodel", "model_count": 66, "name": "CrossModel" }, "crusoe": { "alias_of": null, "base_url": "https://api.inference.crusoecloud.com/v1", "catalog_only": true, "doc": "https://docs.crusoecloud.com/managed-inference/overview", "id": "crusoe", "model_count": 11, "name": "Crusoe" }, "daoxe": { "alias_of": null, "base_url": "https://daoxe.com/v1", "catalog_only": true, "doc": "https://daoxe.com/pricing", "id": "daoxe", "model_count": 9, "name": "DaoXE" }, "databricks": { "alias_of": null, "base_url": "https://${DATABRICKS_HOST}/ai-gateway/mlflow/v1", "catalog_only": true, "doc": "https://docs.databricks.com/aws/en/machine-learning/foundation-models/", "id": "databricks", "model_count": 30, "name": "Databricks" }, "deepinfra": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://deepinfra.com/models", "id": "deepinfra", "model_count": 70, "name": "Deep Infra" }, "deepseek": { "alias_of": null, "base_url": "https://api.deepseek.com", "catalog_only": false, "doc": "https://api-docs.deepseek.com/quick_start/pricing", "id": "deepseek", "model_count": 4, "name": "DeepSeek" }, "digitalocean": { "alias_of": null, "base_url": "https://inference.do-ai.run/v1", "catalog_only": true, "doc": "https://docs.digitalocean.com/products/gradient-ai-platform/details/models/", "id": "digitalocean", "model_count": 100, "name": "DigitalOcean" }, "dinference": { "alias_of": null, "base_url": "https://api.dinference.com/v1", "catalog_only": true, "doc": "https://dinference.com", "id": "dinference", "model_count": 6, "name": "DInference" }, "drun": { "alias_of": null, "base_url": "https://chat.d.run/v1", "catalog_only": true, "doc": "https://www.d.run", "id": "drun", "model_count": 3, "name": "D.Run (China)" }, "ebcloud": { "alias_of": null, "base_url": "https://maas-api.ebcloud.com/v1", "catalog_only": true, "doc": "https://docs.ebtech.com/ai/model-api.html", "id": "ebcloud", "model_count": 4, "name": "EBCloud" }, "echo": { "alias_of": null, "base_url": "https://echo.tracerml.ai/v1", "catalog_only": true, "doc": "https://echo.tracerml.ai/docs/api", "id": "echo", "model_count": 1, "name": "Echo" }, "edenai": { "alias_of": null, "base_url": "https://api.edenai.run/v3", "catalog_only": true, "doc": "https://docs.edenai.co", "id": "edenai", "model_count": 285, "name": "Eden AI" }, "elevenlabs": { "alias_of": null, "base_url": "https://api.elevenlabs.io", "catalog_only": false, "doc": "https://elevenlabs.io/docs/api-reference/introduction", "id": "elevenlabs", "model_count": 4, "name": "ElevenLabs" }, "empiriolabs": { "alias_of": null, "base_url": "https://api.empiriolabs.ai/v1", "catalog_only": true, "doc": "https://docs.empiriolabs.ai", "id": "empiriolabs", "model_count": 66, "name": "EmpirioLabs AI" }, "evroc": { "alias_of": null, "base_url": "https://models.think.evroc.com/v1", "catalog_only": true, "doc": "https://docs.evroc.com/products/think/overview.html", "id": "evroc", "model_count": 16, "name": "evroc" }, "fastrouter": { "alias_of": null, "base_url": "https://go.fastrouter.ai/api/v1", "catalog_only": true, "doc": "https://fastrouter.ai/models", "id": "fastrouter", "model_count": 47, "name": "FastRouter" }, "fireworks_ai": { "alias_of": null, "base_url": "https://api.fireworks.ai/inference/v1", "catalog_only": false, "doc": "https://fireworks.ai/docs/", "id": "fireworks_ai", "model_count": 34, "name": "Fireworks AI" }, "freemodel": { "alias_of": null, "base_url": "https://cc.freemodel.dev/v1", "catalog_only": true, "doc": "https://freemodel.dev", "id": "freemodel", "model_count": 10, "name": "FreeModel" }, "friendli": { "alias_of": null, "base_url": "https://api.friendli.ai/serverless/v1", "catalog_only": false, "doc": "https://friendli.ai/docs/guides/serverless_endpoints/introduction", "id": "friendli", "model_count": 7, "name": "Friendli" }, "frogbot": { "alias_of": null, "base_url": "https://app.frogbot.ai/api/v1", "catalog_only": true, "doc": "https://docs.frogbot.ai", "id": "frogbot", "model_count": 26, "name": "FrogBot" }, "github_copilot": { "alias_of": null, "base_url": "https://api.githubcopilot.com", "catalog_only": true, "doc": "https://docs.github.com/en/copilot", "id": "github_copilot", "model_count": 32, "name": "GitHub Copilot" }, "github_models": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": null, "id": "github_models", "model_count": 0, "name": null }, "gitlab": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.gitlab.com/user/duo_agent_platform/", "id": "gitlab", "model_count": 28, "name": "GitLab Duo" }, "gmicloud": { "alias_of": null, "base_url": "https://api.gmi-serving.com/v1", "catalog_only": true, "doc": "https://docs.gmicloud.ai/inference-engine/api-reference/llm-api-reference", "id": "gmicloud", "model_count": 15, "name": "GMI Cloud" }, "google": { "alias_of": null, "base_url": "https://generativelanguage.googleapis.com/v1beta", "catalog_only": false, "doc": "https://ai.google.dev/gemini-api/docs", "id": "google", "model_count": 66, "name": "Google" }, "google_vertex": { "alias_of": null, "base_url": "https://{region}-aiplatform.googleapis.com", "catalog_only": true, "doc": "https://cloud.google.com/vertex-ai/docs", "id": "google_vertex", "model_count": 58, "name": "Google Vertex AI" }, "google_vertex_anthropic": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://cloud.google.com/vertex-ai/generative-ai/docs/partner-models/claude", "id": "google_vertex_anthropic", "model_count": 15, "name": "Vertex (Anthropic)" }, "greenpt": { "alias_of": null, "base_url": "https://api.greenpt.ai/v1", "catalog_only": true, "doc": "https://docs.greenpt.ai", "id": "greenpt", "model_count": 40, "name": "GreenPT" }, "groq": { "alias_of": null, "base_url": "https://api.groq.com/openai/v1", "catalog_only": false, "doc": "https://groq.com/docs", "id": "groq", "model_count": 16, "name": "Groq" }, "helicone": { "alias_of": null, "base_url": "https://ai-gateway.helicone.ai/v1", "catalog_only": true, "doc": "https://helicone.ai/models", "id": "helicone", "model_count": 90, "name": "Helicone" }, "hetzner": { "alias_of": null, "base_url": "https://inference.hetzner.com/api/v1", "catalog_only": true, "doc": "https://experiments.hetzner.com/docs/inference", "id": "hetzner", "model_count": 2, "name": "Hetzner" }, "hpc_ai": { "alias_of": null, "base_url": "https://api.hpc-ai.com/inference/v1", "catalog_only": true, "doc": "https://www.hpc-ai.com/doc/docs/quickstart/", "id": "hpc_ai", "model_count": 9, "name": "HPC-AI" }, "huggingface": { "alias_of": null, "base_url": "https://router.huggingface.co/v1", "catalog_only": true, "doc": "https://huggingface.co/docs/inference-providers", "id": "huggingface", "model_count": 78, "name": "Hugging Face" }, "hyper": { "alias_of": null, "base_url": "https://hyper.charm.land/v1", "catalog_only": true, "doc": "https://hyper.charm.land", "id": "hyper", "model_count": 23, "name": "Charm Hyper" }, "iflowcn": { "alias_of": null, "base_url": "https://apis.iflow.cn/v1", "catalog_only": true, "doc": "https://platform.iflow.cn/en/docs", "id": "iflowcn", "model_count": 14, "name": "iFlow" }, "impossibl": { "alias_of": null, "base_url": "https://api.impossibl.com/v1", "catalog_only": true, "doc": "https://impossibl.com/docs/models", "id": "impossibl", "model_count": 76, "name": "Impossibl" }, "inception": { "alias_of": null, "base_url": "https://api.inceptionlabs.ai/v1", "catalog_only": true, "doc": "https://docs.inceptionlabs.ai/get-started/models", "id": "inception", "model_count": 3, "name": "Inception" }, "inceptron": { "alias_of": null, "base_url": "https://api.inceptron.io/v1", "catalog_only": true, "doc": "https://docs.inceptron.io", "id": "inceptron", "model_count": 4, "name": "Inceptron" }, "inco": { "alias_of": null, "base_url": "https://api.inco.ai/v1", "catalog_only": true, "doc": "https://platform.inco.ai/docs", "id": "inco", "model_count": 7, "name": "Inco" }, "infer": { "alias_of": null, "base_url": "https://infer.flow7.org/v1", "catalog_only": true, "doc": "https://infer.flow7.org/opencode", "id": "infer", "model_count": 2, "name": "Infer by Flow7" }, "inference": { "alias_of": null, "base_url": "https://inference.net/v1", "catalog_only": true, "doc": "https://inference.net/models", "id": "inference", "model_count": 9, "name": "Inference" }, "inferx": { "alias_of": null, "base_url": "https://model.inferx.net/endpoints/v1", "catalog_only": true, "doc": "https://model.inferx.net/endpoints", "id": "inferx", "model_count": 12, "name": "InferX" }, "infomaniak": { "alias_of": null, "base_url": "https://api.infomaniak.com/2/ai/${INFOMANIAK_PRODUCT_ID}/openai/v1", "catalog_only": true, "doc": "https://www.infomaniak.com/en/hosting/ai-services/open-source-models", "id": "infomaniak", "model_count": 10, "name": "Infomaniak" }, "io_net": { "alias_of": null, "base_url": "https://api.intelligence.io.solutions/api/v1", "catalog_only": true, "doc": "https://io.net/docs/guides/intelligence/io-intelligence", "id": "io_net", "model_count": 17, "name": "IO.NET" }, "iteracompute": { "alias_of": null, "base_url": "https://api.iteracompute.com/v1", "catalog_only": true, "doc": "https://iteracompute.com/docs.html", "id": "iteracompute", "model_count": 9, "name": "IteraCompute" }, "jalapeno": { "alias_of": null, "base_url": "https://api.jalapeno-cloud.ai/v1", "catalog_only": true, "doc": "https://www.jalapeno-cloud.ai/docs/", "id": "jalapeno", "model_count": 17, "name": "Jalapeno Cloud" }, "jiekou": { "alias_of": null, "base_url": "https://api.jiekou.ai/openai", "catalog_only": true, "doc": "https://docs.jiekou.ai/docs/support/quickstart?utm_source=github_models.dev", "id": "jiekou", "model_count": 61, "name": "Jiekou.AI" }, "kenari": { "alias_of": null, "base_url": "https://kenari.id/v1", "catalog_only": true, "doc": "https://kenari.id/docs", "id": "kenari", "model_count": 60, "name": "Kenari" }, "kilo": { "alias_of": null, "base_url": "https://api.kilo.ai/api/gateway", "catalog_only": true, "doc": "https://kilo.ai", "id": "kilo", "model_count": 394, "name": "Kilo Gateway" }, "kimi_code_plan_cn": { "alias_of": null, "base_url": "https://api.kimi.com/coding/v1", "catalog_only": true, "doc": "https://www.kimi.com/code/docs/en/kimi-code/models.html", "id": "kimi_code_plan_cn", "model_count": 4, "name": "Kimi For Coding (kimi.com)" }, "kimi_code_plan_global": { "alias_of": null, "base_url": "https://api.kimi.ai/coding/v1", "catalog_only": true, "doc": "https://www.kimi.ai/code/docs/en/kimi-code/models.html", "id": "kimi_code_plan_global", "model_count": 4, "name": "Kimi For Coding (kimi.ai)" }, "klokintegration": { "alias_of": null, "base_url": "https://api-gw.klok.ipaas.se/proxy/kloker-key/v1", "catalog_only": true, "doc": "https://klokintegration.se/docs/ai-api", "id": "klokintegration", "model_count": 3, "name": "klokintegration.se" }, "kosmik": { "alias_of": null, "base_url": "https://api.koscompute.com/v1", "catalog_only": true, "doc": "https://api.koscompute.com/docs/", "id": "kosmik", "model_count": 1, "name": "Kosmik Compute" }, "kuae_cloud_coding_plan": { "alias_of": null, "base_url": "https://coding-plan-endpoint.kuaecloud.net/v1", "catalog_only": true, "doc": "https://docs.mthreads.com/kuaecloud/kuaecloud-doc-online/coding_plan/", "id": "kuae_cloud_coding_plan", "model_count": 1, "name": "KUAE Cloud Coding Plan" }, "lilac": { "alias_of": null, "base_url": "https://api.getlilac.com/v1", "catalog_only": true, "doc": "https://docs.getlilac.com/inference/models", "id": "lilac", "model_count": 4, "name": "Lilac" }, "llama": { "alias_of": null, "base_url": "https://api.llama.com/compat/v1", "catalog_only": true, "doc": "https://llama.developer.meta.com/docs/models", "id": "llama", "model_count": 7, "name": "Llama" }, "llmgateway": { "alias_of": null, "base_url": "https://api.llmgateway.io/v1", "catalog_only": true, "doc": "https://llmgateway.io/docs", "id": "llmgateway", "model_count": 205, "name": "DevPass (LLM Gateway)" }, "llmgateway_providers": { "alias_of": null, "base_url": "https://api.llmgateway.io/v1", "catalog_only": true, "doc": "https://llmgateway.io/docs", "id": "llmgateway_providers", "model_count": 428, "name": "LLM Gateway" }, "llmtech": { "alias_of": null, "base_url": "https://api.llmtech.eu/v1", "catalog_only": true, "doc": "https://llmtech.eu/models/qwen3.8-27b", "id": "llmtech", "model_count": 1, "name": "LLM Tech" }, "llmtr": { "alias_of": null, "base_url": "https://llmtr.com/v1", "catalog_only": true, "doc": "https://llmtr.com/docs", "id": "llmtr", "model_count": 32, "name": "LLMTR" }, "lmstudio": { "alias_of": null, "base_url": "http://127.0.0.1:1234/v1", "catalog_only": true, "doc": "https://lmstudio.ai/models", "id": "lmstudio", "model_count": 3, "name": "LMStudio" }, "longcat": { "alias_of": null, "base_url": "https://api.longcat.chat/openai", "catalog_only": true, "doc": "https://longcat.chat/platform/docs/", "id": "longcat", "model_count": 1, "name": "LongCat" }, "lucidquery": { "alias_of": null, "base_url": "https://api.lucidquery.com/v1", "catalog_only": true, "doc": "https://lucidquery.com/docs", "id": "lucidquery", "model_count": 4, "name": "LucidQuery" }, "lynkr": { "alias_of": null, "base_url": "http://127.0.0.1:8081/v1", "catalog_only": true, "doc": "https://github.com/Fast-Editor/Lynkr", "id": "lynkr", "model_count": 1, "name": "Lynkr" }, "meganova": { "alias_of": null, "base_url": "https://api.meganova.ai/v1", "catalog_only": true, "doc": "https://docs.meganova.ai", "id": "meganova", "model_count": 19, "name": "Meganova" }, "melious": { "alias_of": null, "base_url": "https://api.melious.ai/v1", "catalog_only": true, "doc": "https://melious.ai/docs/reference/models", "id": "melious", "model_count": 15, "name": "Melious" }, "merge_gateway": { "alias_of": null, "base_url": "https://api-gateway.merge.dev/v1/ai-sdk", "catalog_only": true, "doc": "https://docs.merge.dev/merge-gateway", "id": "merge_gateway", "model_count": 191, "name": "Merge Gateway" }, "meta": { "alias_of": null, "base_url": "https://api.meta.ai/v1", "catalog_only": true, "doc": "https://dev.meta.ai/docs", "id": "meta", "model_count": 5, "name": "Meta" }, "minimax": { "alias_of": null, "base_url": "https://api.minimax.io/v1", "catalog_only": false, "doc": "https://platform.minimax.io/docs/guides/quickstart", "id": "minimax", "model_count": 8, "name": "MiniMax" }, "minimax_cn": { "alias_of": null, "base_url": "https://api.minimax.cn/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimaxi.com/docs/guides/quickstart", "id": "minimax_cn", "model_count": 8, "name": "MiniMax (minimax.cn)" }, "minimax_cn_coding_plan": { "alias_of": null, "base_url": "https://api.minimax.cn/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimaxi.com/docs/token-plan/intro", "id": "minimax_cn_coding_plan", "model_count": 8, "name": "MiniMax Token Plan (minimax.cn)" }, "minimax_coding_plan": { "alias_of": null, "base_url": "https://api.minimax.io/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimax.io/docs/token-plan/intro", "id": "minimax_coding_plan", "model_count": 8, "name": "MiniMax Token Plan (minimax.io)" }, "mistral": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.mistral.ai/getting-started/models/", "id": "mistral", "model_count": 35, "name": "Mistral" }, "mixlayer": { "alias_of": null, "base_url": "https://models.mixlayer.ai/v1", "catalog_only": true, "doc": "https://docs.mixlayer.com", "id": "mixlayer", "model_count": 5, "name": "Mixlayer" }, "moark": { "alias_of": null, "base_url": "https://moark.com/v1", "catalog_only": true, "doc": "https://moark.com/docs/openapi/v1#tag/%E6%96%87%E6%9C%AC%E7%94%9F%E6%88%90", "id": "moark", "model_count": 2, "name": "Moark" }, "modal": { "alias_of": null, "base_url": "https://inference.us-west.modal.direct/v1", "catalog_only": true, "doc": "https://modal.com/docs/guide/endpoints", "id": "modal", "model_count": 4, "name": "Modal" }, "model_oracle_ai": { "alias_of": null, "base_url": "https://api.modeloracle.com/api/v1", "catalog_only": true, "doc": "https://modeloracle.com/setup/", "id": "model_oracle_ai", "model_count": 15, "name": "Model Oracle AI" }, "modelis": { "alias_of": null, "base_url": "https://modelishub.com/v1", "catalog_only": true, "doc": "https://modelishub.com/pricing", "id": "modelis", "model_count": 9, "name": "Modelis" }, "modelscope": { "alias_of": null, "base_url": "https://api-inference.modelscope.cn/v1", "catalog_only": true, "doc": "https://modelscope.cn/docs/model-service/API-Inference/intro", "id": "modelscope", "model_count": 7, "name": "ModelScope" }, "moonshotai": { "alias_of": null, "base_url": "https://api.moonshot.ai/v1", "catalog_only": false, "doc": "https://platform.kimi.ai/docs/api/chat", "id": "moonshotai", "model_count": 10, "name": "Moonshot AI" }, "moonshotai_cn": { "alias_of": null, "base_url": "https://api.moonshot.cn/v1", "catalog_only": false, "doc": "https://platform.moonshot.cn/docs/api/chat", "id": "moonshotai_cn", "model_count": 10, "name": "Moonshot AI (China)" }, "morph": { "alias_of": null, "base_url": "https://api.morphllm.com/v1", "catalog_only": true, "doc": "https://docs.morphllm.com/api-reference/introduction", "id": "morph", "model_count": 3, "name": "Morph" }, "nan": { "alias_of": null, "base_url": "https://api.nan.builders/v1", "catalog_only": true, "doc": "https://nan.builders/docs/models", "id": "nan", "model_count": 7, "name": "NaN" }, "nano_gpt": { "alias_of": null, "base_url": "https://nano-gpt.com/api/v1", "catalog_only": true, "doc": "https://docs.nano-gpt.com", "id": "nano_gpt", "model_count": 593, "name": "NanoGPT" }, "nearai": { "alias_of": null, "base_url": "https://cloud-api.near.ai/v1", "catalog_only": false, "doc": "https://docs.near.ai/", "id": "nearai", "model_count": 36, "name": "NEAR AI Cloud" }, "nebius": { "alias_of": null, "base_url": "https://api.tokenfactory.nebius.com/v1", "catalog_only": true, "doc": "https://docs.tokenfactory.nebius.com/", "id": "nebius", "model_count": 20, "name": "Nebius Token Factory" }, "neon": { "alias_of": null, "base_url": "${NEON_AI_GATEWAY_BASE_URL}/v1", "catalog_only": true, "doc": "https://neon.com/docs", "id": "neon", "model_count": 46, "name": "Neon" }, "neosmith": { "alias_of": null, "base_url": "https://router.neosmith.ai/v1", "catalog_only": true, "doc": "https://neosmith.ai/docs", "id": "neosmith", "model_count": 4, "name": "NeoSmith" }, "neuralwatt": { "alias_of": null, "base_url": "https://api.neuralwatt.com/v1", "catalog_only": true, "doc": "https://portal.neuralwatt.com/docs", "id": "neuralwatt", "model_count": 29, "name": "Neuralwatt" }, "nova": { "alias_of": null, "base_url": "https://api.nova.amazon.com/v1", "catalog_only": true, "doc": "https://nova.amazon.com/dev/documentation", "id": "nova", "model_count": 2, "name": "Nova" }, "novita_ai": { "alias_of": null, "base_url": "https://api.novita.ai/openai", "catalog_only": true, "doc": "https://novita.ai/docs/guides/introduction", "id": "novita_ai", "model_count": 110, "name": "NovitaAI" }, "nvidia": { "alias_of": null, "base_url": "https://integrate.api.nvidia.com/v1", "catalog_only": true, "doc": "https://docs.api.nvidia.com/nim/", "id": "nvidia", "model_count": 105, "name": "Nvidia" }, "oci": { "alias_of": null, "base_url": "https://inference.generativeai.us-chicago-1.oci.oraclecloud.com/openai/v1", "catalog_only": true, "doc": "https://docs.oracle.com/en-us/iaas/Content/generative-ai/pretrained-models.htm", "id": "oci", "model_count": 9, "name": "OCI Generative AI" }, "ofox": { "alias_of": null, "base_url": "https://api.ofox.ai/v1", "catalog_only": true, "doc": "https://ofox.ai/docs", "id": "ofox", "model_count": 148, "name": "Ofox" }, "ollama_cloud": { "alias_of": null, "base_url": "https://ollama.com/v1", "catalog_only": false, "doc": "https://docs.ollama.com/cloud", "id": "ollama_cloud", "model_count": 24, "name": "Ollama Cloud" }, "openai": { "alias_of": null, "base_url": "https://api.openai.com/v1", "catalog_only": false, "doc": "https://platform.openai.com/docs", "id": "openai", "model_count": 154, "name": "OpenAI" }, "opencode": { "alias_of": null, "base_url": "https://opencode.ai/zen/v1", "catalog_only": false, "doc": "https://opencode.ai/docs/zen/", "id": "opencode", "model_count": 111, "name": "OpenCode Zen" }, "opencode_go": { "alias_of": null, "base_url": "https://opencode.ai/zen/go/v1", "catalog_only": true, "doc": "https://opencode.ai/docs/go", "id": "opencode_go", "model_count": 41, "name": "OpenCode Go" }, "openreason": { "alias_of": null, "base_url": "https://api.openreason.app/v1", "catalog_only": true, "doc": "https://openreason.app/docs", "id": "openreason", "model_count": 3, "name": "OpenReason" }, "openrouter": { "alias_of": null, "base_url": "https://openrouter.ai/api/v1", "catalog_only": false, "doc": "https://openrouter.ai/docs", "id": "openrouter", "model_count": 626, "name": "OpenRouter" }, "opper": { "alias_of": null, "base_url": "https://api.opper.ai/v3/compat", "catalog_only": true, "doc": "https://opper.ai/models", "id": "opper", "model_count": 57, "name": "Opper" }, "orcarouter": { "alias_of": null, "base_url": "https://api.orcarouter.ai/v1", "catalog_only": true, "doc": "https://docs.orcarouter.ai", "id": "orcarouter", "model_count": 117, "name": "OrcaRouter" }, "ovhcloud": { "alias_of": null, "base_url": "https://oai.endpoints.kepler.ai.cloud.ovh.net/v1", "catalog_only": true, "doc": "https://www.ovhcloud.com/en/public-cloud/ai-endpoints/catalog//", "id": "ovhcloud", "model_count": 14, "name": "OVHcloud AI Endpoints" }, "pendra": { "alias_of": null, "base_url": "https://api.pendra.ai/api/v1", "catalog_only": true, "doc": "https://pendra.ai/docs/integrations/opencode", "id": "pendra", "model_count": 6, "name": "Pendra" }, "perplexity": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.perplexity.ai", "id": "perplexity", "model_count": 4, "name": "Perplexity" }, "perplexity_agent": { "alias_of": null, "base_url": "https://api.perplexity.ai/v1", "catalog_only": true, "doc": "https://docs.perplexity.ai/docs/agent-api/models", "id": "perplexity_agent", "model_count": 22, "name": "Perplexity Agent" }, "pioneer": { "alias_of": null, "base_url": "https://api.pioneer.ai/v1", "catalog_only": true, "doc": "https://agent.pioneer.ai/llms.txt", "id": "pioneer", "model_count": 114, "name": "Pioneer" }, "poe": { "alias_of": null, "base_url": "https://api.poe.com/v1", "catalog_only": true, "doc": "https://creator.poe.com/docs/external-applications/openai-compatible-api", "id": "poe", "model_count": 137, "name": "Poe" }, "poolside": { "alias_of": null, "base_url": "https://inference.poolside.ai/v1", "catalog_only": true, "doc": "https://platform.poolside.ai", "id": "poolside", "model_count": 3, "name": "Poolside" }, "privatemode_ai": { "alias_of": null, "base_url": "http://localhost:8080/v1", "catalog_only": true, "doc": "https://docs.privatemode.ai/api/overview", "id": "privatemode_ai", "model_count": 11, "name": "Privatemode AI" }, "qihang_ai": { "alias_of": null, "base_url": "https://api.qhaigc.net/v1", "catalog_only": true, "doc": "https://www.qhaigc.net/docs", "id": "qihang_ai", "model_count": 9, "name": "QiHang" }, "qiniu_ai": { "alias_of": null, "base_url": "https://api.qnaigc.com/v1", "catalog_only": true, "doc": "https://developer.qiniu.com/aitokenapi", "id": "qiniu_ai", "model_count": 91, "name": "Qiniu" }, "qvac": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://www.npmjs.com/package/@qvac/ai-sdk-provider", "id": "qvac", "model_count": 9, "name": "QVAC" }, "regolo_ai": { "alias_of": null, "base_url": "https://api.regolo.ai/v1", "catalog_only": true, "doc": "https://docs.regolo.ai/", "id": "regolo_ai", "model_count": 18, "name": "Regolo AI" }, "requesty": { "alias_of": null, "base_url": "https://router.requesty.ai/v1", "catalog_only": false, "doc": "https://requesty.ai/solution/llm-routing/models", "id": "requesty", "model_count": 159, "name": "Requesty" }, "routing_run": { "alias_of": null, "base_url": "https://api.routing.run/v1", "catalog_only": true, "doc": "https://docs.routing.run/api-reference/models", "id": "routing_run", "model_count": 15, "name": "routing.run" }, "runinfra": { "alias_of": null, "base_url": "https://api.runinfra.ai/v1", "catalog_only": true, "doc": "https://runinfra.ai/docs", "id": "runinfra", "model_count": 7, "name": "RunInfra" }, "sakana": { "alias_of": null, "base_url": "https://api.sakana.ai/v1", "catalog_only": true, "doc": "https://console.sakana.ai/models", "id": "sakana", "model_count": 4, "name": "Sakana AI" }, "salad_cloud": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.salad.com/ai-gateway/explanation/overview", "id": "salad_cloud", "model_count": 1, "name": "SaladCloud AI Gateway" }, "sap_ai_core": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://help.sap.com/docs/sap-ai-core", "id": "sap_ai_core", "model_count": 49, "name": "SAP AI Core" }, "sarvam": { "alias_of": null, "base_url": "https://api.sarvam.ai/v1", "catalog_only": true, "doc": "https://docs.sarvam.ai/api-reference-docs/getting-started/models", "id": "sarvam", "model_count": 2, "name": "Sarvam AI" }, "scaleway": { "alias_of": null, "base_url": "https://api.scaleway.ai/v1", "catalog_only": true, "doc": "https://www.scaleway.com/en/docs/generative-apis/", "id": "scaleway", "model_count": 15, "name": "Scaleway" }, "scnet_token_plan": { "alias_of": null, "base_url": "https://api.scnet.cn/api/llm/v1", "catalog_only": true, "doc": "https://www.scnet.cn/ac/openapi/doc/2.0/moduleapi/plans/token-plan.html", "id": "scnet_token_plan", "model_count": 19, "name": "SCNet Token Plan" }, "scx_ai": { "alias_of": null, "base_url": "https://api.scx.ai/v1", "catalog_only": true, "doc": "https://platform.scx.ai/docs", "id": "scx_ai", "model_count": 4, "name": "SCX.ai" }, "sensenova": { "alias_of": null, "base_url": "https://token.sensenova.cn/v1", "catalog_only": true, "doc": "https://platform.sensenova.cn/docs", "id": "sensenova", "model_count": 5, "name": "SenseNova (China)" }, "siliconflow": { "alias_of": null, "base_url": "https://api.siliconflow.com/v1", "catalog_only": true, "doc": "https://cloud.siliconflow.com/models", "id": "siliconflow", "model_count": 57, "name": "SiliconFlow" }, "siliconflow_cn": { "alias_of": null, "base_url": "https://api.siliconflow.cn/v1", "catalog_only": true, "doc": "https://cloud.siliconflow.com/models", "id": "siliconflow_cn", "model_count": 44, "name": "SiliconFlow (China)" }, "snowflake_cortex": { "alias_of": null, "base_url": "https://${SNOWFLAKE_ACCOUNT}.snowflakecomputing.com/api/v2/cortex/v1", "catalog_only": true, "doc": "https://docs.snowflake.com/en/user-guide/snowflake-cortex/cortex-rest-api", "id": "snowflake_cortex", "model_count": 25, "name": "Snowflake Cortex" }, "stackit": { "alias_of": null, "base_url": "https://api.openai-compat.model-serving.eu01.onstackit.cloud/v1", "catalog_only": true, "doc": "https://docs.stackit.cloud/products/data-and-ai/ai-model-serving/basics/available-shared-models", "id": "stackit", "model_count": 8, "name": "STACKIT" }, "standardcompute": { "alias_of": null, "base_url": "https://api.stdcmpt.com/v1", "catalog_only": true, "doc": "https://standardcompute.com/models", "id": "standardcompute", "model_count": 1, "name": "Standard Compute" }, "stepfun": { "alias_of": null, "base_url": "https://api.stepfun.com/v1", "catalog_only": true, "doc": "https://platform.stepfun.com/docs/zh/overview/concept", "id": "stepfun", "model_count": 9, "name": "StepFun (China)" }, "stepfun_ai": { "alias_of": null, "base_url": "https://api.stepfun.ai/v1", "catalog_only": true, "doc": "https://platform.stepfun.ai/docs/en/overview/concept", "id": "stepfun_ai", "model_count": 9, "name": "StepFun (Global)" }, "stepfun_ai_step_plan": { "alias_of": null, "base_url": "https://api.stepfun.ai/step_plan/v1", "catalog_only": true, "doc": "https://platform.stepfun.ai/docs/en/step-plan/integrations/reasoning-api", "id": "stepfun_ai_step_plan", "model_count": 4, "name": "StepFun Step Plan (Global)" }, "stepfun_step_plan": { "alias_of": null, "base_url": "https://api.stepfun.com/step_plan/v1", "catalog_only": true, "doc": "https://platform.stepfun.com/docs/zh/step-plan/integrations/reasoning-api", "id": "stepfun_step_plan", "model_count": 5, "name": "StepFun Step Plan (China)" }, "subconscious": { "alias_of": null, "base_url": "https://api.subconscious.dev/v1", "catalog_only": true, "doc": "https://docs.subconscious.dev", "id": "subconscious", "model_count": 2, "name": "Subconscious" }, "submodel": { "alias_of": null, "base_url": "https://llm.submodel.ai/v1", "catalog_only": true, "doc": "https://submodel.gitbook.io", "id": "submodel", "model_count": 9, "name": "submodel" }, "synthetic": { "alias_of": null, "base_url": "https://api.synthetic.new/openai/v1", "catalog_only": true, "doc": "https://synthetic.new/pricing", "id": "synthetic", "model_count": 10, "name": "Synthetic" }, "tempr": { "alias_of": null, "base_url": "https://api.temprhq.io/v1", "catalog_only": true, "doc": "https://temprhq.io/docs/gateway-reference.html", "id": "tempr", "model_count": 39, "name": "Tempr" }, "tencent_coding_plan": { "alias_of": null, "base_url": "https://api.lkeap.cloud.tencent.com/coding/v3", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1772/128947", "id": "tencent_coding_plan", "model_count": 8, "name": "Tencent Coding Plan (China)" }, "tencent_token_plan": { "alias_of": null, "base_url": "https://api.lkeap.cloud.tencent.com/plan/v3", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1823/130060", "id": "tencent_token_plan", "model_count": 2, "name": "Tencent Token Plan" }, "tencent_tokenhub": { "alias_of": null, "base_url": "https://tokenhub.tencentmaas.com/v1", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1823/130050", "id": "tencent_tokenhub", "model_count": 3, "name": "Tencent TokenHub" }, "tensorx": { "alias_of": null, "base_url": "https://api.tensorx.ai/v1", "catalog_only": true, "doc": "https://docs.tensorx.ai/", "id": "tensorx", "model_count": 25, "name": "TensorX" }, "the_grid_ai": { "alias_of": null, "base_url": "https://api.thegrid.ai/v1", "catalog_only": true, "doc": "https://thegrid.ai/docs", "id": "the_grid_ai", "model_count": 9, "name": "The Grid AI" }, "thinkingmachines": { "alias_of": null, "base_url": "https://tinker.thinkingmachines.dev/services/tinker-prod/anthropic/api/v1", "catalog_only": true, "doc": "https://tinker-docs.thinkingmachines.ai/tinker/compatible-apis/anthropic/", "id": "thinkingmachines", "model_count": 2, "name": "Thinking Machines" }, "tinfoil": { "alias_of": null, "base_url": "https://inference.tinfoil.sh/v1", "catalog_only": true, "doc": "https://docs.tinfoil.sh", "id": "tinfoil", "model_count": 9, "name": "Tinfoil" }, "togetherai": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.together.ai/docs/serverless-models", "id": "togetherai", "model_count": 39, "name": "Together AI" }, "tokengo": { "alias_of": null, "base_url": "https://api.tokengo.com/v1", "catalog_only": true, "doc": "https://www.tokengo.com/docs", "id": "tokengo", "model_count": 13, "name": "TokenGo" }, "tokenrouter": { "alias_of": null, "base_url": "https://api.tokenrouter.com/v1", "catalog_only": true, "doc": "https://www.tokenrouter.com/docs/tokenrouter-feature-guide/", "id": "tokenrouter", "model_count": 1, "name": "TokenRouter" }, "trustedrouter": { "alias_of": null, "base_url": "https://api.trustedrouter.com/v1", "catalog_only": true, "doc": "https://trustedrouter.com/docs", "id": "trustedrouter", "model_count": 7, "name": "TrustedRouter" }, "typesafe": { "alias_of": null, "base_url": "https://api.typesafe.ai", "catalog_only": false, "doc": "https://docs.typesafe.ai/api", "id": "typesafe", "model_count": 3, "name": "TypeSafe AI" }, "umans_ai": { "alias_of": null, "base_url": "https://api.code.umans.ai/v1", "catalog_only": true, "doc": "https://app.umans.ai/offers/code/docs/orgs", "id": "umans_ai", "model_count": 6, "name": "Umans AI" }, "umans_ai_coding_plan": { "alias_of": null, "base_url": "https://api.code.umans.ai/v1", "catalog_only": true, "doc": "https://app.umans.ai/offers/code/docs", "id": "umans_ai_coding_plan", "model_count": 7, "name": "Umans AI Coding Plan" }, "unorouter": { "alias_of": null, "base_url": "https://api.unorouter.com/v1", "catalog_only": true, "doc": "https://unorouter.com/models", "id": "unorouter", "model_count": 23, "name": "UnoRouter" }, "upstage": { "alias_of": null, "base_url": "https://api.upstage.ai/v1/solar", "catalog_only": true, "doc": "https://developers.upstage.ai/docs/apis/chat", "id": "upstage", "model_count": 4, "name": "Upstage" }, "v0": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://sdk.vercel.ai/providers/ai-sdk-providers/vercel", "id": "v0", "model_count": 3, "name": "v0" }, "vancine": { "alias_of": null, "base_url": "https://vancine.com/v1", "catalog_only": true, "doc": "https://vancine.com/docs", "id": "vancine", "model_count": 8, "name": "Vancine" }, "venice": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.venice.ai", "id": "venice", "model_count": 111, "name": "Venice AI" }, "vercel": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://github.com/vercel/ai/tree/5eb85cc45a259553501f535b8ac79a77d0e79223/packages/gateway", "id": "vercel", "model_count": 389, "name": "Vercel AI Gateway" }, "vispark": { "alias_of": null, "base_url": "https://api.lab.vispark.in/v1", "catalog_only": true, "doc": "https://lab.vispark.in/#vision", "id": "vispark", "model_count": 3, "name": "Vispark" }, "vivgrid": { "alias_of": null, "base_url": "https://api.vivgrid.com/v1", "catalog_only": true, "doc": "https://docs.vivgrid.com/models", "id": "vivgrid", "model_count": 34, "name": "Vivgrid" }, "volcengine": { "alias_of": null, "base_url": "https://ark.cn-beijing.volces.com/api/v3", "catalog_only": true, "doc": "https://www.volcengine.com/docs/82379/1330310", "id": "volcengine", "model_count": 16, "name": "Volcengine Ark" }, "volcengine_coding_plan": { "alias_of": null, "base_url": "https://ark.cn-beijing.volces.com/api/coding/v3", "catalog_only": true, "doc": "https://www.volcengine.com/docs/82379/1928261", "id": "volcengine_coding_plan", "model_count": 10, "name": "Volcengine Ark Coding Plan" }, "vultr": { "alias_of": null, "base_url": "https://api.vultrinference.com/v1", "catalog_only": true, "doc": "https://api.vultrinference.com/", "id": "vultr", "model_count": 10, "name": "Vultr" }, "wafer_ai": { "alias_of": null, "base_url": "https://pass.wafer.ai/v1", "catalog_only": true, "doc": "https://docs.wafer.ai/wafer-pass", "id": "wafer_ai", "model_count": 5, "name": "Wafer" }, "wallaby": { "alias_of": null, "base_url": "https://api.wallabytoken.com/v1", "catalog_only": true, "doc": "https://wallabytoken.com/docs", "id": "wallaby", "model_count": 1, "name": "Wallaby" }, "wandb": { "alias_of": null, "base_url": "https://api.inference.wandb.ai/v1", "catalog_only": true, "doc": "https://docs.wandb.ai/inference", "id": "wandb", "model_count": 29, "name": "CoreWeave" }, "watsonx": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://www.ibm.com/docs/en/watsonx/saas?topic=solutions-supported-foundation-models", "id": "watsonx", "model_count": 5, "name": "watsonx.ai" }, "xai": { "alias_of": null, "base_url": "https://api.x.ai/v1", "catalog_only": false, "doc": "https://docs.x.ai", "id": "xai", "model_count": 27, "name": "xAI" }, "xiaomi": { "alias_of": null, "base_url": "https://api.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi", "model_count": 9, "name": "Xiaomi" }, "xiaomi_token_plan_ams": { "alias_of": null, "base_url": "https://token-plan-ams.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_ams", "model_count": 9, "name": "Xiaomi Token Plan (Europe)" }, "xiaomi_token_plan_cn": { "alias_of": null, "base_url": "https://token-plan-cn.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_cn", "model_count": 9, "name": "Xiaomi Token Plan (China)" }, "xiaomi_token_plan_sgp": { "alias_of": null, "base_url": "https://token-plan-sgp.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_sgp", "model_count": 9, "name": "Xiaomi Token Plan (Singapore)" }, "xpersona": { "alias_of": null, "base_url": "https://www.xpersona.co/v1", "catalog_only": true, "doc": "https://www.xpersona.co/docs", "id": "xpersona", "model_count": 13, "name": "Xpersona" }, "zai": { "alias_of": null, "base_url": "https://api.z.ai/api/paas/v4", "catalog_only": false, "doc": "https://docs.z.ai/guides/overview/pricing", "id": "zai", "model_count": 18, "name": "Z.AI" }, "zai_coder": { "alias_of": null, "base_url": "https://api.zai.chat", "catalog_only": true, "doc": "https://docs.zai.chat", "id": "zai_coder", "model_count": 5, "name": "Z.AI Coder" }, "zai_coding_plan": { "alias_of": null, "base_url": "https://api.z.ai/api/coding/paas/v4", "catalog_only": true, "doc": "https://docs.z.ai/devpack/overview", "id": "zai_coding_plan", "model_count": 7, "name": "Z.AI Coding Plan" }, "zeldoc": { "alias_of": null, "base_url": "https://api.zeldoc.ai/v1", "catalog_only": true, "doc": "https://docs.zeldoc.ai", "id": "zeldoc", "model_count": 1, "name": "Zeldoc" }, "zenifra": { "alias_of": null, "base_url": "https://ai.zenifra.com/v1", "catalog_only": true, "doc": "https://docs.zenifra.com", "id": "zenifra", "model_count": 1, "name": "Zenifra" }, "zenmux": { "alias_of": null, "base_url": "https://zenmux.ai/api/v1", "catalog_only": false, "doc": "https://docs.zenmux.ai", "id": "zenmux", "model_count": 237, "name": "Zenmux" }, "zhipuai": { "alias_of": null, "base_url": "https://open.bigmodel.cn/api/paas/v4", "catalog_only": true, "doc": "https://docs.z.ai/guides/overview/pricing", "id": "zhipuai", "model_count": 17, "name": "Zhipu AI" }, "zhipuai_coding_plan": { "alias_of": null, "base_url": "https://open.bigmodel.cn/api/coding/paas/v4", "catalog_only": true, "doc": "https://docs.bigmodel.cn/cn/coding-plan/overview", "id": "zhipuai_coding_plan", "model_count": 4, "name": "Zhipu AI Coding Plan" } }, "snapshot_id": "52464f773a83f25ee942758f9b4350f8afbb238a723eb4e26db2ddc30ec8d50d", "snapshot_schema_version": 1 };
1
+ export const manifest = { "catalog_version": 2, "format_version": 1, "generated_at": "2026-09-25T18:03:32.815437Z", "model_count": 8800, "provider_count": 228, "providers": { "302ai": { "alias_of": null, "base_url": "https://api.302.ai/v1", "catalog_only": true, "doc": "https://doc.302.ai", "id": "302ai", "model_count": 117, "name": "302.AI" }, "a2agent": { "alias_of": null, "base_url": "https://api.a2agent.me/v1", "catalog_only": false, "doc": "https://docs.a2agent.me", "id": "a2agent", "model_count": 20, "name": "A2Agent" }, "abacus": { "alias_of": null, "base_url": "https://routellm.abacus.ai/v1", "catalog_only": true, "doc": "https://abacus.ai/help/api", "id": "abacus", "model_count": 108, "name": "Abacus" }, "abliteration_ai": { "alias_of": null, "base_url": "https://api.abliteration.ai/v1", "catalog_only": true, "doc": "https://docs.abliteration.ai/models", "id": "abliteration_ai", "model_count": 3, "name": "abliteration.ai" }, "above": { "alias_of": null, "base_url": "https://api.above.dev/v1", "catalog_only": true, "doc": "https://above.dev/docs", "id": "above", "model_count": 9, "name": "above.dev" }, "agentrouter": { "alias_of": null, "base_url": "https://agentrouter.org/v1", "catalog_only": true, "doc": "https://agentrouter.org/docs/opencode.html", "id": "agentrouter", "model_count": 5, "name": "AgentRouter" }, "agnes": { "alias_of": null, "base_url": "https://apihub.agnes-ai.com/v1", "catalog_only": true, "doc": "https://agnes-ai.com/doc", "id": "agnes", "model_count": 3, "name": "Agnes AI" }, "ai21": { "alias_of": null, "base_url": "https://api.ai21.com/studio/v1", "catalog_only": true, "doc": "https://docs.ai21.com/docs/jamba-foundation-models", "id": "ai21", "model_count": 2, "name": "AI21 Labs" }, "ai_router": { "alias_of": null, "base_url": "https://api.ai-router.dev/v1", "catalog_only": true, "doc": "https://ai-router.dev/openai-compatible-api-gateway/", "id": "ai_router", "model_count": 5, "name": "AI-ROUTER" }, "aiand": { "alias_of": null, "base_url": "https://api.aiand.com/v1", "catalog_only": true, "doc": "https://docs.aiand.com/", "id": "aiand", "model_count": 11, "name": "ai&" }, "aihubmix": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.aihubmix.com", "id": "aihubmix", "model_count": 106, "name": "AIHubMix" }, "ainetcafe": { "alias_of": null, "base_url": "https://microquickjs.com/v1", "catalog_only": true, "doc": "https://ainetcafe.com/k3/guides/", "id": "ainetcafe", "model_count": 1, "name": "ainetcafe" }, "aixy": { "alias_of": null, "base_url": "https://api.aixy-gateway.com/v1", "catalog_only": true, "doc": "https://docs.aixy-gateway.com/integrations/overview", "id": "aixy", "model_count": 1, "name": "Aixy" }, "aki_io": { "alias_of": null, "base_url": "https://aki.io/v1", "catalog_only": true, "doc": "https://aki.io/docs/", "id": "aki_io", "model_count": 7, "name": "AKI.IO" }, "alibaba": { "alias_of": null, "base_url": "https://dashscope-intl.aliyuncs.com/compatible-mode/v1", "catalog_only": false, "doc": "https://www.alibabacloud.com/help/en/model-studio/models", "id": "alibaba", "model_count": 56, "name": "Alibaba" }, "alibaba_cn": { "alias_of": null, "base_url": "https://dashscope.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/models", "id": "alibaba_cn", "model_count": 90, "name": "Alibaba (China)" }, "alibaba_coding_plan": { "alias_of": null, "base_url": "https://coding-intl.dashscope.aliyuncs.com/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/coding-plan", "id": "alibaba_coding_plan", "model_count": 12, "name": "Alibaba Coding Plan" }, "alibaba_coding_plan_cn": { "alias_of": null, "base_url": "https://coding.dashscope.aliyuncs.com/v1", "catalog_only": true, "doc": "https://help.aliyun.com/zh/model-studio/coding-plan", "id": "alibaba_coding_plan_cn", "model_count": 12, "name": "Alibaba Coding Plan (China)" }, "alibaba_token_plan": { "alias_of": null, "base_url": "https://token-plan.ap-southeast-1.maas.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/en/model-studio/token-plan-overview", "id": "alibaba_token_plan", "model_count": 28, "name": "Alibaba Token Plan" }, "alibaba_token_plan_cn": { "alias_of": null, "base_url": "https://token-plan.cn-beijing.maas.aliyuncs.com/compatible-mode/v1", "catalog_only": true, "doc": "https://www.alibabacloud.com/help/zh/model-studio/token-plan-overview", "id": "alibaba_token_plan_cn", "model_count": 28, "name": "Alibaba Token Plan (China)" }, "amazon_bedrock": { "alias_of": null, "base_url": "https://bedrock-runtime.{region}.amazonaws.com", "catalog_only": true, "doc": "https://docs.aws.amazon.com/bedrock/", "id": "amazon_bedrock", "model_count": 199, "name": "Amazon Bedrock" }, "ambient": { "alias_of": null, "base_url": "https://api.ambient.xyz/v1", "catalog_only": true, "doc": "https://ambient.xyz", "id": "ambient", "model_count": 10, "name": "Ambient" }, "amd": { "alias_of": null, "base_url": "https://developer.amd.com.cn/radeon/api/v1", "catalog_only": true, "doc": "https://developer.amd.com.cn/radeon/tokenfactory", "id": "amd", "model_count": 6, "name": "AMD" }, "anthropic": { "alias_of": null, "base_url": "https://api.anthropic.com", "catalog_only": false, "doc": "https://docs.anthropic.com", "id": "anthropic", "model_count": 18, "name": "Anthropic" }, "anyapi": { "alias_of": null, "base_url": "https://api.anyapi.ai/v1", "catalog_only": true, "doc": "https://docs.anyapi.ai", "id": "anyapi", "model_count": 30, "name": "AnyAPI" }, "arcee": { "alias_of": null, "base_url": "https://api.arcee.ai/api/v1", "catalog_only": true, "doc": "https://docs.arcee.ai", "id": "arcee", "model_count": 7, "name": "Arcee" }, "atomic_chat": { "alias_of": null, "base_url": "http://127.0.0.1:1337/v1", "catalog_only": true, "doc": "https://atomic.chat", "id": "atomic_chat", "model_count": 5, "name": "Atomic Chat" }, "auriko": { "alias_of": null, "base_url": "https://api.auriko.ai/v1", "catalog_only": true, "doc": "https://docs.auriko.ai", "id": "auriko", "model_count": 15, "name": "Auriko" }, "azure": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://learn.microsoft.com/en-us/azure/ai-services/openai/concepts/models", "id": "azure", "model_count": 118, "name": "Azure" }, "azure_cognitive_services": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://learn.microsoft.com/en-us/azure/ai-services/openai/concepts/models", "id": "azure_cognitive_services", "model_count": 78, "name": "Azure Cognitive Services" }, "bailing": { "alias_of": null, "base_url": "https://api.tbox.cn/api/llm/v1/chat/completions", "catalog_only": true, "doc": "https://alipaytbox.yuque.com/sxs0ba/ling/intro", "id": "bailing", "model_count": 2, "name": "Bailing" }, "baseten": { "alias_of": null, "base_url": "https://inference.baseten.co/v1", "catalog_only": true, "doc": "https://docs.baseten.co/inference/model-apis/overview", "id": "baseten", "model_count": 23, "name": "Baseten" }, "berget": { "alias_of": null, "base_url": "https://api.berget.ai/v1", "catalog_only": true, "doc": "https://api.berget.ai", "id": "berget", "model_count": 6, "name": "Berget.AI" }, "blueclaw": { "alias_of": null, "base_url": "https://openai.blueclaw.network/v1", "catalog_only": true, "doc": "https://blueclaw.network", "id": "blueclaw", "model_count": 2, "name": "Blue Claw" }, "bothub": { "alias_of": null, "base_url": "https://openai.bothub.ru/v1", "catalog_only": true, "doc": "https://bothub.ru/models", "id": "bothub", "model_count": 8, "name": "Bothub" }, "cerebras": { "alias_of": null, "base_url": "https://api.cerebras.ai/v1", "catalog_only": false, "doc": "https://cerebras.ai/docs", "id": "cerebras", "model_count": 4, "name": "Cerebras" }, "chutes": { "alias_of": null, "base_url": "https://llm.chutes.ai/v1", "catalog_only": true, "doc": "https://llm.chutes.ai/v1/models", "id": "chutes", "model_count": 14, "name": "Chutes" }, "clarifai": { "alias_of": null, "base_url": "https://api.clarifai.com/v2/ext/openai/v1", "catalog_only": true, "doc": "https://docs.clarifai.com/compute/inference/", "id": "clarifai", "model_count": 12, "name": "Clarifai" }, "claudinio": { "alias_of": null, "base_url": "https://api.claudin.io/v1", "catalog_only": true, "doc": "https://claudin.io", "id": "claudinio", "model_count": 2, "name": "Claudinio" }, "cline_pass": { "alias_of": null, "base_url": "https://api.cline.bot/api/v1", "catalog_only": true, "doc": "https://docs.cline.bot/getting-started/clinepass", "id": "cline_pass", "model_count": 18, "name": "ClinePass" }, "cloudferro_sherlock": { "alias_of": null, "base_url": "https://api-sherlock.cloudferro.com/openai/v1", "catalog_only": true, "doc": "https://docs.sherlock.cloudferro.com/", "id": "cloudferro_sherlock", "model_count": 5, "name": "CloudFerro Sherlock" }, "cloudflare_ai_gateway": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://developers.cloudflare.com/ai-gateway/", "id": "cloudflare_ai_gateway", "model_count": 50, "name": "Cloudflare AI Gateway" }, "cloudflare_workers_ai": { "alias_of": null, "base_url": "https://api.cloudflare.com/client/v4/accounts/${CLOUDFLARE_ACCOUNT_ID}/ai/v1", "catalog_only": false, "doc": "https://developers.cloudflare.com/workers-ai/models/", "id": "cloudflare_workers_ai", "model_count": 28, "name": "Cloudflare Workers AI" }, "cohere": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.cohere.com/docs/models", "id": "cohere", "model_count": 19, "name": "Cohere" }, "coralbricks": { "alias_of": null, "base_url": "https://inference.coralbricks.ai/v1", "catalog_only": true, "doc": "https://www.coralbricks.ai/docs", "id": "coralbricks", "model_count": 4, "name": "CoralBricks" }, "cortecs": { "alias_of": null, "base_url": "https://api.cortecs.ai/v1", "catalog_only": true, "doc": "https://api.cortecs.ai/v1/models", "id": "cortecs", "model_count": 109, "name": "Cortecs" }, "crof": { "alias_of": null, "base_url": "https://crof.ai/v1", "catalog_only": true, "doc": "https://crof.ai/docs", "id": "crof", "model_count": 24, "name": "CrofAI" }, "crossmodel": { "alias_of": null, "base_url": "https://api.crossmodel.ai/v1", "catalog_only": true, "doc": "https://www.crossmodel.ai/docs", "id": "crossmodel", "model_count": 66, "name": "CrossModel" }, "crusoe": { "alias_of": null, "base_url": "https://api.inference.crusoecloud.com/v1", "catalog_only": true, "doc": "https://docs.crusoecloud.com/managed-inference/overview", "id": "crusoe", "model_count": 11, "name": "Crusoe" }, "daoxe": { "alias_of": null, "base_url": "https://daoxe.com/v1", "catalog_only": true, "doc": "https://daoxe.com/pricing", "id": "daoxe", "model_count": 9, "name": "DaoXE" }, "databricks": { "alias_of": null, "base_url": "https://${DATABRICKS_HOST}/ai-gateway/mlflow/v1", "catalog_only": true, "doc": "https://docs.databricks.com/aws/en/machine-learning/foundation-models/", "id": "databricks", "model_count": 30, "name": "Databricks" }, "deepinfra": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://deepinfra.com/models", "id": "deepinfra", "model_count": 70, "name": "Deep Infra" }, "deepseek": { "alias_of": null, "base_url": "https://api.deepseek.com", "catalog_only": false, "doc": "https://api-docs.deepseek.com/quick_start/pricing", "id": "deepseek", "model_count": 4, "name": "DeepSeek" }, "digitalocean": { "alias_of": null, "base_url": "https://inference.do-ai.run/v1", "catalog_only": true, "doc": "https://docs.digitalocean.com/products/gradient-ai-platform/details/models/", "id": "digitalocean", "model_count": 100, "name": "DigitalOcean" }, "dinference": { "alias_of": null, "base_url": "https://api.dinference.com/v1", "catalog_only": true, "doc": "https://dinference.com", "id": "dinference", "model_count": 6, "name": "DInference" }, "drun": { "alias_of": null, "base_url": "https://chat.d.run/v1", "catalog_only": true, "doc": "https://www.d.run", "id": "drun", "model_count": 3, "name": "D.Run (China)" }, "ebcloud": { "alias_of": null, "base_url": "https://maas-api.ebcloud.com/v1", "catalog_only": true, "doc": "https://docs.ebtech.com/ai/model-api.html", "id": "ebcloud", "model_count": 4, "name": "EBCloud" }, "echo": { "alias_of": null, "base_url": "https://echo.tracerml.ai/v1", "catalog_only": true, "doc": "https://echo.tracerml.ai/docs/api", "id": "echo", "model_count": 1, "name": "Echo" }, "edenai": { "alias_of": null, "base_url": "https://api.edenai.run/v3", "catalog_only": true, "doc": "https://docs.edenai.co", "id": "edenai", "model_count": 285, "name": "Eden AI" }, "elevenlabs": { "alias_of": null, "base_url": "https://api.elevenlabs.io", "catalog_only": false, "doc": "https://elevenlabs.io/docs/api-reference/introduction", "id": "elevenlabs", "model_count": 4, "name": "ElevenLabs" }, "empiriolabs": { "alias_of": null, "base_url": "https://api.empiriolabs.ai/v1", "catalog_only": true, "doc": "https://docs.empiriolabs.ai", "id": "empiriolabs", "model_count": 66, "name": "EmpirioLabs AI" }, "evroc": { "alias_of": null, "base_url": "https://models.think.evroc.com/v1", "catalog_only": true, "doc": "https://docs.evroc.com/products/think/overview.html", "id": "evroc", "model_count": 16, "name": "evroc" }, "fastrouter": { "alias_of": null, "base_url": "https://go.fastrouter.ai/api/v1", "catalog_only": true, "doc": "https://fastrouter.ai/models", "id": "fastrouter", "model_count": 47, "name": "FastRouter" }, "fireworks_ai": { "alias_of": null, "base_url": "https://api.fireworks.ai/inference/v1", "catalog_only": false, "doc": "https://fireworks.ai/docs/", "id": "fireworks_ai", "model_count": 34, "name": "Fireworks AI" }, "freemodel": { "alias_of": null, "base_url": "https://cc.freemodel.dev/v1", "catalog_only": true, "doc": "https://freemodel.dev", "id": "freemodel", "model_count": 10, "name": "FreeModel" }, "friendli": { "alias_of": null, "base_url": "https://api.friendli.ai/serverless/v1", "catalog_only": false, "doc": "https://friendli.ai/docs/guides/serverless_endpoints/introduction", "id": "friendli", "model_count": 7, "name": "Friendli" }, "frogbot": { "alias_of": null, "base_url": "https://app.frogbot.ai/api/v1", "catalog_only": true, "doc": "https://docs.frogbot.ai", "id": "frogbot", "model_count": 26, "name": "FrogBot" }, "github_copilot": { "alias_of": null, "base_url": "https://api.githubcopilot.com", "catalog_only": true, "doc": "https://docs.github.com/en/copilot", "id": "github_copilot", "model_count": 32, "name": "GitHub Copilot" }, "github_models": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": null, "id": "github_models", "model_count": 0, "name": null }, "gitlab": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.gitlab.com/user/duo_agent_platform/", "id": "gitlab", "model_count": 28, "name": "GitLab Duo" }, "gmicloud": { "alias_of": null, "base_url": "https://api.gmi-serving.com/v1", "catalog_only": true, "doc": "https://docs.gmicloud.ai/inference-engine/api-reference/llm-api-reference", "id": "gmicloud", "model_count": 15, "name": "GMI Cloud" }, "google": { "alias_of": null, "base_url": "https://generativelanguage.googleapis.com/v1beta", "catalog_only": false, "doc": "https://ai.google.dev/gemini-api/docs", "id": "google", "model_count": 66, "name": "Google" }, "google_vertex": { "alias_of": null, "base_url": "https://{region}-aiplatform.googleapis.com", "catalog_only": true, "doc": "https://cloud.google.com/vertex-ai/docs", "id": "google_vertex", "model_count": 58, "name": "Google Vertex AI" }, "google_vertex_anthropic": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://cloud.google.com/vertex-ai/generative-ai/docs/partner-models/claude", "id": "google_vertex_anthropic", "model_count": 15, "name": "Vertex (Anthropic)" }, "greenpt": { "alias_of": null, "base_url": "https://api.greenpt.ai/v1", "catalog_only": true, "doc": "https://docs.greenpt.ai", "id": "greenpt", "model_count": 40, "name": "GreenPT" }, "groq": { "alias_of": null, "base_url": "https://api.groq.com/openai/v1", "catalog_only": false, "doc": "https://groq.com/docs", "id": "groq", "model_count": 16, "name": "Groq" }, "helicone": { "alias_of": null, "base_url": "https://ai-gateway.helicone.ai/v1", "catalog_only": true, "doc": "https://helicone.ai/models", "id": "helicone", "model_count": 90, "name": "Helicone" }, "hetzner": { "alias_of": null, "base_url": "https://inference.hetzner.com/api/v1", "catalog_only": true, "doc": "https://experiments.hetzner.com/docs/inference", "id": "hetzner", "model_count": 2, "name": "Hetzner" }, "hpc_ai": { "alias_of": null, "base_url": "https://api.hpc-ai.com/inference/v1", "catalog_only": true, "doc": "https://www.hpc-ai.com/doc/docs/quickstart/", "id": "hpc_ai", "model_count": 9, "name": "HPC-AI" }, "huggingface": { "alias_of": null, "base_url": "https://router.huggingface.co/v1", "catalog_only": true, "doc": "https://huggingface.co/docs/inference-providers", "id": "huggingface", "model_count": 78, "name": "Hugging Face" }, "hyper": { "alias_of": null, "base_url": "https://hyper.charm.land/v1", "catalog_only": true, "doc": "https://hyper.charm.land", "id": "hyper", "model_count": 23, "name": "Charm Hyper" }, "iflowcn": { "alias_of": null, "base_url": "https://apis.iflow.cn/v1", "catalog_only": true, "doc": "https://platform.iflow.cn/en/docs", "id": "iflowcn", "model_count": 14, "name": "iFlow" }, "impossibl": { "alias_of": null, "base_url": "https://api.impossibl.com/v1", "catalog_only": true, "doc": "https://impossibl.com/docs/models", "id": "impossibl", "model_count": 76, "name": "Impossibl" }, "inception": { "alias_of": null, "base_url": "https://api.inceptionlabs.ai/v1", "catalog_only": true, "doc": "https://docs.inceptionlabs.ai/get-started/models", "id": "inception", "model_count": 3, "name": "Inception" }, "inceptron": { "alias_of": null, "base_url": "https://api.inceptron.io/v1", "catalog_only": true, "doc": "https://docs.inceptron.io", "id": "inceptron", "model_count": 4, "name": "Inceptron" }, "inco": { "alias_of": null, "base_url": "https://api.inco.ai/v1", "catalog_only": true, "doc": "https://platform.inco.ai/docs", "id": "inco", "model_count": 7, "name": "Inco" }, "infer": { "alias_of": null, "base_url": "https://infer.flow7.org/v1", "catalog_only": true, "doc": "https://infer.flow7.org/opencode", "id": "infer", "model_count": 2, "name": "Infer by Flow7" }, "inference": { "alias_of": null, "base_url": "https://inference.net/v1", "catalog_only": true, "doc": "https://inference.net/models", "id": "inference", "model_count": 9, "name": "Inference" }, "inferx": { "alias_of": null, "base_url": "https://model.inferx.net/endpoints/v1", "catalog_only": true, "doc": "https://model.inferx.net/endpoints", "id": "inferx", "model_count": 12, "name": "InferX" }, "infomaniak": { "alias_of": null, "base_url": "https://api.infomaniak.com/2/ai/${INFOMANIAK_PRODUCT_ID}/openai/v1", "catalog_only": true, "doc": "https://www.infomaniak.com/en/hosting/ai-services/open-source-models", "id": "infomaniak", "model_count": 10, "name": "Infomaniak" }, "io_net": { "alias_of": null, "base_url": "https://api.intelligence.io.solutions/api/v1", "catalog_only": true, "doc": "https://io.net/docs/guides/intelligence/io-intelligence", "id": "io_net", "model_count": 17, "name": "IO.NET" }, "iteracompute": { "alias_of": null, "base_url": "https://api.iteracompute.com/v1", "catalog_only": true, "doc": "https://iteracompute.com/docs.html", "id": "iteracompute", "model_count": 9, "name": "IteraCompute" }, "jalapeno": { "alias_of": null, "base_url": "https://api.jalapeno-cloud.ai/v1", "catalog_only": true, "doc": "https://www.jalapeno-cloud.ai/docs/", "id": "jalapeno", "model_count": 17, "name": "Jalapeno Cloud" }, "jiekou": { "alias_of": null, "base_url": "https://api.jiekou.ai/openai", "catalog_only": true, "doc": "https://docs.jiekou.ai/docs/support/quickstart?utm_source=github_models.dev", "id": "jiekou", "model_count": 61, "name": "Jiekou.AI" }, "kenari": { "alias_of": null, "base_url": "https://kenari.id/v1", "catalog_only": true, "doc": "https://kenari.id/docs", "id": "kenari", "model_count": 60, "name": "Kenari" }, "kilo": { "alias_of": null, "base_url": "https://api.kilo.ai/api/gateway", "catalog_only": true, "doc": "https://kilo.ai", "id": "kilo", "model_count": 394, "name": "Kilo Gateway" }, "kimi_code_plan_cn": { "alias_of": null, "base_url": "https://api.kimi.com/coding/v1", "catalog_only": true, "doc": "https://www.kimi.com/code/docs/en/kimi-code/models.html", "id": "kimi_code_plan_cn", "model_count": 4, "name": "Kimi For Coding (kimi.com)" }, "kimi_code_plan_global": { "alias_of": null, "base_url": "https://api.kimi.ai/coding/v1", "catalog_only": true, "doc": "https://www.kimi.ai/code/docs/en/kimi-code/models.html", "id": "kimi_code_plan_global", "model_count": 4, "name": "Kimi For Coding (kimi.ai)" }, "klokintegration": { "alias_of": null, "base_url": "https://api-gw.klok.ipaas.se/proxy/kloker-key/v1", "catalog_only": true, "doc": "https://klokintegration.se/docs/ai-api", "id": "klokintegration", "model_count": 3, "name": "klokintegration.se" }, "kosmik": { "alias_of": null, "base_url": "https://api.koscompute.com/v1", "catalog_only": true, "doc": "https://api.koscompute.com/docs/", "id": "kosmik", "model_count": 1, "name": "Kosmik Compute" }, "kuae_cloud_coding_plan": { "alias_of": null, "base_url": "https://coding-plan-endpoint.kuaecloud.net/v1", "catalog_only": true, "doc": "https://docs.mthreads.com/kuaecloud/kuaecloud-doc-online/coding_plan/", "id": "kuae_cloud_coding_plan", "model_count": 1, "name": "KUAE Cloud Coding Plan" }, "lilac": { "alias_of": null, "base_url": "https://api.getlilac.com/v1", "catalog_only": true, "doc": "https://docs.getlilac.com/inference/models", "id": "lilac", "model_count": 4, "name": "Lilac" }, "llama": { "alias_of": null, "base_url": "https://api.llama.com/compat/v1", "catalog_only": true, "doc": "https://llama.developer.meta.com/docs/models", "id": "llama", "model_count": 7, "name": "Llama" }, "llmgateway": { "alias_of": null, "base_url": "https://api.llmgateway.io/v1", "catalog_only": true, "doc": "https://llmgateway.io/docs", "id": "llmgateway", "model_count": 205, "name": "DevPass (LLM Gateway)" }, "llmgateway_providers": { "alias_of": null, "base_url": "https://api.llmgateway.io/v1", "catalog_only": true, "doc": "https://llmgateway.io/docs", "id": "llmgateway_providers", "model_count": 428, "name": "LLM Gateway" }, "llmtech": { "alias_of": null, "base_url": "https://api.llmtech.eu/v1", "catalog_only": true, "doc": "https://llmtech.eu/models/qwen3.8-27b", "id": "llmtech", "model_count": 1, "name": "LLM Tech" }, "llmtr": { "alias_of": null, "base_url": "https://llmtr.com/v1", "catalog_only": true, "doc": "https://llmtr.com/docs", "id": "llmtr", "model_count": 32, "name": "LLMTR" }, "lmstudio": { "alias_of": null, "base_url": "http://127.0.0.1:1234/v1", "catalog_only": true, "doc": "https://lmstudio.ai/models", "id": "lmstudio", "model_count": 3, "name": "LMStudio" }, "longcat": { "alias_of": null, "base_url": "https://api.longcat.chat/openai", "catalog_only": true, "doc": "https://longcat.chat/platform/docs/", "id": "longcat", "model_count": 1, "name": "LongCat" }, "lucidquery": { "alias_of": null, "base_url": "https://api.lucidquery.com/v1", "catalog_only": true, "doc": "https://lucidquery.com/docs", "id": "lucidquery", "model_count": 4, "name": "LucidQuery" }, "lynkr": { "alias_of": null, "base_url": "http://127.0.0.1:8081/v1", "catalog_only": true, "doc": "https://github.com/Fast-Editor/Lynkr", "id": "lynkr", "model_count": 1, "name": "Lynkr" }, "meganova": { "alias_of": null, "base_url": "https://api.meganova.ai/v1", "catalog_only": true, "doc": "https://docs.meganova.ai", "id": "meganova", "model_count": 19, "name": "Meganova" }, "melious": { "alias_of": null, "base_url": "https://api.melious.ai/v1", "catalog_only": true, "doc": "https://melious.ai/docs/reference/models", "id": "melious", "model_count": 15, "name": "Melious" }, "merge_gateway": { "alias_of": null, "base_url": "https://api-gateway.merge.dev/v1/ai-sdk", "catalog_only": true, "doc": "https://docs.merge.dev/merge-gateway", "id": "merge_gateway", "model_count": 191, "name": "Merge Gateway" }, "meta": { "alias_of": null, "base_url": "https://api.meta.ai/v1", "catalog_only": true, "doc": "https://dev.meta.ai/docs", "id": "meta", "model_count": 5, "name": "Meta" }, "minimax": { "alias_of": null, "base_url": "https://api.minimax.io/v1", "catalog_only": false, "doc": "https://platform.minimax.io/docs/guides/quickstart", "id": "minimax", "model_count": 8, "name": "MiniMax" }, "minimax_cn": { "alias_of": null, "base_url": "https://api.minimax.cn/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimaxi.com/docs/guides/quickstart", "id": "minimax_cn", "model_count": 8, "name": "MiniMax (minimax.cn)" }, "minimax_cn_coding_plan": { "alias_of": null, "base_url": "https://api.minimax.cn/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimaxi.com/docs/token-plan/intro", "id": "minimax_cn_coding_plan", "model_count": 8, "name": "MiniMax Token Plan (minimax.cn)" }, "minimax_coding_plan": { "alias_of": null, "base_url": "https://api.minimax.io/anthropic/v1", "catalog_only": true, "doc": "https://platform.minimax.io/docs/token-plan/intro", "id": "minimax_coding_plan", "model_count": 8, "name": "MiniMax Token Plan (minimax.io)" }, "mistral": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.mistral.ai/getting-started/models/", "id": "mistral", "model_count": 35, "name": "Mistral" }, "mixlayer": { "alias_of": null, "base_url": "https://models.mixlayer.ai/v1", "catalog_only": true, "doc": "https://docs.mixlayer.com", "id": "mixlayer", "model_count": 5, "name": "Mixlayer" }, "moark": { "alias_of": null, "base_url": "https://moark.com/v1", "catalog_only": true, "doc": "https://moark.com/docs/openapi/v1#tag/%E6%96%87%E6%9C%AC%E7%94%9F%E6%88%90", "id": "moark", "model_count": 2, "name": "Moark" }, "modal": { "alias_of": null, "base_url": "https://inference.us-west.modal.direct/v1", "catalog_only": true, "doc": "https://modal.com/docs/guide/endpoints", "id": "modal", "model_count": 4, "name": "Modal" }, "model_oracle_ai": { "alias_of": null, "base_url": "https://api.modeloracle.com/api/v1", "catalog_only": true, "doc": "https://modeloracle.com/setup/", "id": "model_oracle_ai", "model_count": 15, "name": "Model Oracle AI" }, "modelis": { "alias_of": null, "base_url": "https://modelishub.com/v1", "catalog_only": true, "doc": "https://modelishub.com/pricing", "id": "modelis", "model_count": 9, "name": "Modelis" }, "modelscope": { "alias_of": null, "base_url": "https://api-inference.modelscope.cn/v1", "catalog_only": true, "doc": "https://modelscope.cn/docs/model-service/API-Inference/intro", "id": "modelscope", "model_count": 7, "name": "ModelScope" }, "moonshotai": { "alias_of": null, "base_url": "https://api.moonshot.ai/v1", "catalog_only": false, "doc": "https://platform.kimi.ai/docs/api/chat", "id": "moonshotai", "model_count": 10, "name": "Moonshot AI" }, "moonshotai_cn": { "alias_of": null, "base_url": "https://api.moonshot.cn/v1", "catalog_only": false, "doc": "https://platform.moonshot.cn/docs/api/chat", "id": "moonshotai_cn", "model_count": 10, "name": "Moonshot AI (China)" }, "morph": { "alias_of": null, "base_url": "https://api.morphllm.com/v1", "catalog_only": true, "doc": "https://docs.morphllm.com/api-reference/introduction", "id": "morph", "model_count": 3, "name": "Morph" }, "nan": { "alias_of": null, "base_url": "https://api.nan.builders/v1", "catalog_only": true, "doc": "https://nan.builders/docs/models", "id": "nan", "model_count": 7, "name": "NaN" }, "nano_gpt": { "alias_of": null, "base_url": "https://nano-gpt.com/api/v1", "catalog_only": true, "doc": "https://docs.nano-gpt.com", "id": "nano_gpt", "model_count": 593, "name": "NanoGPT" }, "nearai": { "alias_of": null, "base_url": "https://cloud-api.near.ai/v1", "catalog_only": false, "doc": "https://docs.near.ai/", "id": "nearai", "model_count": 36, "name": "NEAR AI Cloud" }, "nebius": { "alias_of": null, "base_url": "https://api.tokenfactory.nebius.com/v1", "catalog_only": true, "doc": "https://docs.tokenfactory.nebius.com/", "id": "nebius", "model_count": 20, "name": "Nebius Token Factory" }, "neon": { "alias_of": null, "base_url": "${NEON_AI_GATEWAY_BASE_URL}/v1", "catalog_only": true, "doc": "https://neon.com/docs", "id": "neon", "model_count": 46, "name": "Neon" }, "neosmith": { "alias_of": null, "base_url": "https://router.neosmith.ai/v1", "catalog_only": true, "doc": "https://neosmith.ai/docs", "id": "neosmith", "model_count": 4, "name": "NeoSmith" }, "neuralwatt": { "alias_of": null, "base_url": "https://api.neuralwatt.com/v1", "catalog_only": true, "doc": "https://portal.neuralwatt.com/docs", "id": "neuralwatt", "model_count": 29, "name": "Neuralwatt" }, "nova": { "alias_of": null, "base_url": "https://api.nova.amazon.com/v1", "catalog_only": true, "doc": "https://nova.amazon.com/dev/documentation", "id": "nova", "model_count": 2, "name": "Nova" }, "novita_ai": { "alias_of": null, "base_url": "https://api.novita.ai/openai", "catalog_only": true, "doc": "https://novita.ai/docs/guides/introduction", "id": "novita_ai", "model_count": 110, "name": "NovitaAI" }, "nvidia": { "alias_of": null, "base_url": "https://integrate.api.nvidia.com/v1", "catalog_only": true, "doc": "https://docs.api.nvidia.com/nim/", "id": "nvidia", "model_count": 105, "name": "Nvidia" }, "oci": { "alias_of": null, "base_url": "https://inference.generativeai.us-chicago-1.oci.oraclecloud.com/openai/v1", "catalog_only": true, "doc": "https://docs.oracle.com/en-us/iaas/Content/generative-ai/pretrained-models.htm", "id": "oci", "model_count": 9, "name": "OCI Generative AI" }, "ofox": { "alias_of": null, "base_url": "https://api.ofox.ai/v1", "catalog_only": true, "doc": "https://ofox.ai/docs", "id": "ofox", "model_count": 148, "name": "Ofox" }, "ollama_cloud": { "alias_of": null, "base_url": "https://ollama.com/v1", "catalog_only": false, "doc": "https://docs.ollama.com/cloud", "id": "ollama_cloud", "model_count": 24, "name": "Ollama Cloud" }, "openai": { "alias_of": null, "base_url": "https://api.openai.com/v1", "catalog_only": false, "doc": "https://platform.openai.com/docs", "id": "openai", "model_count": 154, "name": "OpenAI" }, "opencode": { "alias_of": null, "base_url": "https://opencode.ai/zen/v1", "catalog_only": false, "doc": "https://opencode.ai/docs/zen/", "id": "opencode", "model_count": 111, "name": "OpenCode Zen" }, "opencode_go": { "alias_of": null, "base_url": "https://opencode.ai/zen/go/v1", "catalog_only": true, "doc": "https://opencode.ai/docs/go", "id": "opencode_go", "model_count": 41, "name": "OpenCode Go" }, "openreason": { "alias_of": null, "base_url": "https://api.openreason.app/v1", "catalog_only": true, "doc": "https://openreason.app/docs", "id": "openreason", "model_count": 3, "name": "OpenReason" }, "openrouter": { "alias_of": null, "base_url": "https://openrouter.ai/api/v1", "catalog_only": false, "doc": "https://openrouter.ai/docs", "id": "openrouter", "model_count": 626, "name": "OpenRouter" }, "opper": { "alias_of": null, "base_url": "https://api.opper.ai/v3/compat", "catalog_only": true, "doc": "https://opper.ai/models", "id": "opper", "model_count": 57, "name": "Opper" }, "orcarouter": { "alias_of": null, "base_url": "https://api.orcarouter.ai/v1", "catalog_only": true, "doc": "https://docs.orcarouter.ai", "id": "orcarouter", "model_count": 117, "name": "OrcaRouter" }, "ovhcloud": { "alias_of": null, "base_url": "https://oai.endpoints.kepler.ai.cloud.ovh.net/v1", "catalog_only": true, "doc": "https://www.ovhcloud.com/en/public-cloud/ai-endpoints/catalog//", "id": "ovhcloud", "model_count": 14, "name": "OVHcloud AI Endpoints" }, "pendra": { "alias_of": null, "base_url": "https://api.pendra.ai/api/v1", "catalog_only": true, "doc": "https://pendra.ai/docs/integrations/opencode", "id": "pendra", "model_count": 6, "name": "Pendra" }, "perplexity": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.perplexity.ai", "id": "perplexity", "model_count": 4, "name": "Perplexity" }, "perplexity_agent": { "alias_of": null, "base_url": "https://api.perplexity.ai/v1", "catalog_only": true, "doc": "https://docs.perplexity.ai/docs/agent-api/models", "id": "perplexity_agent", "model_count": 22, "name": "Perplexity Agent" }, "pioneer": { "alias_of": null, "base_url": "https://api.pioneer.ai/v1", "catalog_only": true, "doc": "https://agent.pioneer.ai/llms.txt", "id": "pioneer", "model_count": 114, "name": "Pioneer" }, "poe": { "alias_of": null, "base_url": "https://api.poe.com/v1", "catalog_only": true, "doc": "https://creator.poe.com/docs/external-applications/openai-compatible-api", "id": "poe", "model_count": 137, "name": "Poe" }, "poolside": { "alias_of": null, "base_url": "https://inference.poolside.ai/v1", "catalog_only": true, "doc": "https://platform.poolside.ai", "id": "poolside", "model_count": 3, "name": "Poolside" }, "privatemode_ai": { "alias_of": null, "base_url": "http://localhost:8080/v1", "catalog_only": true, "doc": "https://docs.privatemode.ai/api/overview", "id": "privatemode_ai", "model_count": 11, "name": "Privatemode AI" }, "qihang_ai": { "alias_of": null, "base_url": "https://api.qhaigc.net/v1", "catalog_only": true, "doc": "https://www.qhaigc.net/docs", "id": "qihang_ai", "model_count": 9, "name": "QiHang" }, "qiniu_ai": { "alias_of": null, "base_url": "https://api.qnaigc.com/v1", "catalog_only": true, "doc": "https://developer.qiniu.com/aitokenapi", "id": "qiniu_ai", "model_count": 91, "name": "Qiniu" }, "qvac": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://www.npmjs.com/package/@qvac/ai-sdk-provider", "id": "qvac", "model_count": 9, "name": "QVAC" }, "regolo_ai": { "alias_of": null, "base_url": "https://api.regolo.ai/v1", "catalog_only": true, "doc": "https://docs.regolo.ai/", "id": "regolo_ai", "model_count": 18, "name": "Regolo AI" }, "requesty": { "alias_of": null, "base_url": "https://router.requesty.ai/v1", "catalog_only": false, "doc": "https://requesty.ai/solution/llm-routing/models", "id": "requesty", "model_count": 159, "name": "Requesty" }, "routing_run": { "alias_of": null, "base_url": "https://api.routing.run/v1", "catalog_only": true, "doc": "https://docs.routing.run/api-reference/models", "id": "routing_run", "model_count": 15, "name": "routing.run" }, "runinfra": { "alias_of": null, "base_url": "https://api.runinfra.ai/v1", "catalog_only": true, "doc": "https://runinfra.ai/docs", "id": "runinfra", "model_count": 7, "name": "RunInfra" }, "sakana": { "alias_of": null, "base_url": "https://api.sakana.ai/v1", "catalog_only": true, "doc": "https://console.sakana.ai/models", "id": "sakana", "model_count": 4, "name": "Sakana AI" }, "salad_cloud": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://docs.salad.com/ai-gateway/explanation/overview", "id": "salad_cloud", "model_count": 1, "name": "SaladCloud AI Gateway" }, "sap_ai_core": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://help.sap.com/docs/sap-ai-core", "id": "sap_ai_core", "model_count": 49, "name": "SAP AI Core" }, "sarvam": { "alias_of": null, "base_url": "https://api.sarvam.ai/v1", "catalog_only": true, "doc": "https://docs.sarvam.ai/api-reference-docs/getting-started/models", "id": "sarvam", "model_count": 2, "name": "Sarvam AI" }, "scaleway": { "alias_of": null, "base_url": "https://api.scaleway.ai/v1", "catalog_only": true, "doc": "https://www.scaleway.com/en/docs/generative-apis/", "id": "scaleway", "model_count": 15, "name": "Scaleway" }, "scnet_token_plan": { "alias_of": null, "base_url": "https://api.scnet.cn/api/llm/v1", "catalog_only": true, "doc": "https://www.scnet.cn/ac/openapi/doc/2.0/moduleapi/plans/token-plan.html", "id": "scnet_token_plan", "model_count": 19, "name": "SCNet Token Plan" }, "scx_ai": { "alias_of": null, "base_url": "https://api.scx.ai/v1", "catalog_only": true, "doc": "https://platform.scx.ai/docs", "id": "scx_ai", "model_count": 4, "name": "SCX.ai" }, "sensenova": { "alias_of": null, "base_url": "https://token.sensenova.cn/v1", "catalog_only": true, "doc": "https://platform.sensenova.cn/docs", "id": "sensenova", "model_count": 5, "name": "SenseNova (China)" }, "siliconflow": { "alias_of": null, "base_url": "https://api.siliconflow.com/v1", "catalog_only": true, "doc": "https://cloud.siliconflow.com/models", "id": "siliconflow", "model_count": 57, "name": "SiliconFlow" }, "siliconflow_cn": { "alias_of": null, "base_url": "https://api.siliconflow.cn/v1", "catalog_only": true, "doc": "https://cloud.siliconflow.com/models", "id": "siliconflow_cn", "model_count": 44, "name": "SiliconFlow (China)" }, "snowflake_cortex": { "alias_of": null, "base_url": "https://${SNOWFLAKE_ACCOUNT}.snowflakecomputing.com/api/v2/cortex/v1", "catalog_only": true, "doc": "https://docs.snowflake.com/en/user-guide/snowflake-cortex/cortex-rest-api", "id": "snowflake_cortex", "model_count": 25, "name": "Snowflake Cortex" }, "stackit": { "alias_of": null, "base_url": "https://api.openai-compat.model-serving.eu01.onstackit.cloud/v1", "catalog_only": true, "doc": "https://docs.stackit.cloud/products/data-and-ai/ai-model-serving/basics/available-shared-models", "id": "stackit", "model_count": 8, "name": "STACKIT" }, "standardcompute": { "alias_of": null, "base_url": "https://api.stdcmpt.com/v1", "catalog_only": true, "doc": "https://standardcompute.com/models", "id": "standardcompute", "model_count": 1, "name": "Standard Compute" }, "stepfun": { "alias_of": null, "base_url": "https://api.stepfun.com/v1", "catalog_only": true, "doc": "https://platform.stepfun.com/docs/zh/overview/concept", "id": "stepfun", "model_count": 9, "name": "StepFun (China)" }, "stepfun_ai": { "alias_of": null, "base_url": "https://api.stepfun.ai/v1", "catalog_only": true, "doc": "https://platform.stepfun.ai/docs/en/overview/concept", "id": "stepfun_ai", "model_count": 9, "name": "StepFun (Global)" }, "stepfun_ai_step_plan": { "alias_of": null, "base_url": "https://api.stepfun.ai/step_plan/v1", "catalog_only": true, "doc": "https://platform.stepfun.ai/docs/en/step-plan/integrations/reasoning-api", "id": "stepfun_ai_step_plan", "model_count": 4, "name": "StepFun Step Plan (Global)" }, "stepfun_step_plan": { "alias_of": null, "base_url": "https://api.stepfun.com/step_plan/v1", "catalog_only": true, "doc": "https://platform.stepfun.com/docs/zh/step-plan/integrations/reasoning-api", "id": "stepfun_step_plan", "model_count": 5, "name": "StepFun Step Plan (China)" }, "subconscious": { "alias_of": null, "base_url": "https://api.subconscious.dev/v1", "catalog_only": true, "doc": "https://docs.subconscious.dev", "id": "subconscious", "model_count": 2, "name": "Subconscious" }, "submodel": { "alias_of": null, "base_url": "https://llm.submodel.ai/v1", "catalog_only": true, "doc": "https://submodel.gitbook.io", "id": "submodel", "model_count": 9, "name": "submodel" }, "synthetic": { "alias_of": null, "base_url": "https://api.synthetic.new/openai/v1", "catalog_only": true, "doc": "https://synthetic.new/pricing", "id": "synthetic", "model_count": 10, "name": "Synthetic" }, "tempr": { "alias_of": null, "base_url": "https://api.temprhq.io/v1", "catalog_only": true, "doc": "https://temprhq.io/docs/gateway-reference.html", "id": "tempr", "model_count": 39, "name": "Tempr" }, "tencent_coding_plan": { "alias_of": null, "base_url": "https://api.lkeap.cloud.tencent.com/coding/v3", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1772/128947", "id": "tencent_coding_plan", "model_count": 8, "name": "Tencent Coding Plan (China)" }, "tencent_token_plan": { "alias_of": null, "base_url": "https://api.lkeap.cloud.tencent.com/plan/v3", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1823/130060", "id": "tencent_token_plan", "model_count": 2, "name": "Tencent Token Plan" }, "tencent_tokenhub": { "alias_of": null, "base_url": "https://tokenhub.tencentmaas.com/v1", "catalog_only": true, "doc": "https://cloud.tencent.com/document/product/1823/130050", "id": "tencent_tokenhub", "model_count": 3, "name": "Tencent TokenHub" }, "tensorx": { "alias_of": null, "base_url": "https://api.tensorx.ai/v1", "catalog_only": true, "doc": "https://docs.tensorx.ai/", "id": "tensorx", "model_count": 25, "name": "TensorX" }, "the_grid_ai": { "alias_of": null, "base_url": "https://api.thegrid.ai/v1", "catalog_only": true, "doc": "https://thegrid.ai/docs", "id": "the_grid_ai", "model_count": 9, "name": "The Grid AI" }, "thinkingmachines": { "alias_of": null, "base_url": "https://tinker.thinkingmachines.dev/services/tinker-prod/anthropic/api/v1", "catalog_only": true, "doc": "https://tinker-docs.thinkingmachines.ai/tinker/compatible-apis/anthropic/", "id": "thinkingmachines", "model_count": 2, "name": "Thinking Machines" }, "tinfoil": { "alias_of": null, "base_url": "https://inference.tinfoil.sh/v1", "catalog_only": true, "doc": "https://docs.tinfoil.sh", "id": "tinfoil", "model_count": 9, "name": "Tinfoil" }, "togetherai": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.together.ai/docs/serverless-models", "id": "togetherai", "model_count": 39, "name": "Together AI" }, "tokengo": { "alias_of": null, "base_url": "https://api.tokengo.com/v1", "catalog_only": true, "doc": "https://www.tokengo.com/docs", "id": "tokengo", "model_count": 13, "name": "TokenGo" }, "tokenrouter": { "alias_of": null, "base_url": "https://api.tokenrouter.com/v1", "catalog_only": true, "doc": "https://www.tokenrouter.com/docs/tokenrouter-feature-guide/", "id": "tokenrouter", "model_count": 1, "name": "TokenRouter" }, "trustedrouter": { "alias_of": null, "base_url": "https://api.trustedrouter.com/v1", "catalog_only": true, "doc": "https://trustedrouter.com/docs", "id": "trustedrouter", "model_count": 7, "name": "TrustedRouter" }, "typesafe": { "alias_of": null, "base_url": "https://api.typesafe.ai", "catalog_only": false, "doc": "https://docs.typesafe.ai/api", "id": "typesafe", "model_count": 3, "name": "TypeSafe AI" }, "umans_ai": { "alias_of": null, "base_url": "https://api.code.umans.ai/v1", "catalog_only": true, "doc": "https://app.umans.ai/offers/code/docs/orgs", "id": "umans_ai", "model_count": 6, "name": "Umans AI" }, "umans_ai_coding_plan": { "alias_of": null, "base_url": "https://api.code.umans.ai/v1", "catalog_only": true, "doc": "https://app.umans.ai/offers/code/docs", "id": "umans_ai_coding_plan", "model_count": 7, "name": "Umans AI Coding Plan" }, "unorouter": { "alias_of": null, "base_url": "https://api.unorouter.com/v1", "catalog_only": true, "doc": "https://unorouter.com/models", "id": "unorouter", "model_count": 23, "name": "UnoRouter" }, "upstage": { "alias_of": null, "base_url": "https://api.upstage.ai/v1/solar", "catalog_only": true, "doc": "https://developers.upstage.ai/docs/apis/chat", "id": "upstage", "model_count": 4, "name": "Upstage" }, "v0": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://sdk.vercel.ai/providers/ai-sdk-providers/vercel", "id": "v0", "model_count": 3, "name": "v0" }, "vancine": { "alias_of": null, "base_url": "https://vancine.com/v1", "catalog_only": true, "doc": "https://vancine.com/docs", "id": "vancine", "model_count": 8, "name": "Vancine" }, "venice": { "alias_of": null, "base_url": null, "catalog_only": false, "doc": "https://docs.venice.ai", "id": "venice", "model_count": 111, "name": "Venice AI" }, "vercel": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://github.com/vercel/ai/tree/5eb85cc45a259553501f535b8ac79a77d0e79223/packages/gateway", "id": "vercel", "model_count": 389, "name": "Vercel AI Gateway" }, "vispark": { "alias_of": null, "base_url": "https://api.lab.vispark.in/v1", "catalog_only": true, "doc": "https://lab.vispark.in/#vision", "id": "vispark", "model_count": 3, "name": "Vispark" }, "vivgrid": { "alias_of": null, "base_url": "https://api.vivgrid.com/v1", "catalog_only": true, "doc": "https://docs.vivgrid.com/models", "id": "vivgrid", "model_count": 34, "name": "Vivgrid" }, "volcengine": { "alias_of": null, "base_url": "https://ark.cn-beijing.volces.com/api/v3", "catalog_only": true, "doc": "https://www.volcengine.com/docs/82379/1330310", "id": "volcengine", "model_count": 16, "name": "Volcengine Ark" }, "volcengine_coding_plan": { "alias_of": null, "base_url": "https://ark.cn-beijing.volces.com/api/coding/v3", "catalog_only": true, "doc": "https://www.volcengine.com/docs/82379/1928261", "id": "volcengine_coding_plan", "model_count": 10, "name": "Volcengine Ark Coding Plan" }, "vultr": { "alias_of": null, "base_url": "https://api.vultrinference.com/v1", "catalog_only": true, "doc": "https://api.vultrinference.com/", "id": "vultr", "model_count": 10, "name": "Vultr" }, "wafer_ai": { "alias_of": null, "base_url": "https://pass.wafer.ai/v1", "catalog_only": true, "doc": "https://docs.wafer.ai/wafer-pass", "id": "wafer_ai", "model_count": 5, "name": "Wafer" }, "wallaby": { "alias_of": null, "base_url": "https://api.wallabytoken.com/v1", "catalog_only": true, "doc": "https://wallabytoken.com/docs", "id": "wallaby", "model_count": 1, "name": "Wallaby" }, "wandb": { "alias_of": null, "base_url": "https://api.inference.wandb.ai/v1", "catalog_only": true, "doc": "https://docs.wandb.ai/inference", "id": "wandb", "model_count": 29, "name": "CoreWeave" }, "watsonx": { "alias_of": null, "base_url": null, "catalog_only": true, "doc": "https://www.ibm.com/docs/en/watsonx/saas?topic=solutions-supported-foundation-models", "id": "watsonx", "model_count": 5, "name": "watsonx.ai" }, "xai": { "alias_of": null, "base_url": "https://api.x.ai/v1", "catalog_only": false, "doc": "https://docs.x.ai", "id": "xai", "model_count": 27, "name": "xAI" }, "xiaomi": { "alias_of": null, "base_url": "https://api.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi", "model_count": 9, "name": "Xiaomi" }, "xiaomi_token_plan_ams": { "alias_of": null, "base_url": "https://token-plan-ams.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_ams", "model_count": 9, "name": "Xiaomi Token Plan (Europe)" }, "xiaomi_token_plan_cn": { "alias_of": null, "base_url": "https://token-plan-cn.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_cn", "model_count": 9, "name": "Xiaomi Token Plan (China)" }, "xiaomi_token_plan_sgp": { "alias_of": null, "base_url": "https://token-plan-sgp.xiaomimimo.com/v1", "catalog_only": true, "doc": "https://platform.xiaomimimo.com/#/docs", "id": "xiaomi_token_plan_sgp", "model_count": 9, "name": "Xiaomi Token Plan (Singapore)" }, "xpersona": { "alias_of": null, "base_url": "https://www.xpersona.co/v1", "catalog_only": true, "doc": "https://www.xpersona.co/docs", "id": "xpersona", "model_count": 13, "name": "Xpersona" }, "zai": { "alias_of": null, "base_url": "https://api.z.ai/api/paas/v4", "catalog_only": false, "doc": "https://docs.z.ai/guides/overview/pricing", "id": "zai", "model_count": 18, "name": "Z.AI" }, "zai_coder": { "alias_of": null, "base_url": "https://api.zai.chat", "catalog_only": true, "doc": "https://docs.zai.chat", "id": "zai_coder", "model_count": 5, "name": "Z.AI Coder" }, "zai_coding_plan": { "alias_of": null, "base_url": "https://api.z.ai/api/coding/paas/v4", "catalog_only": true, "doc": "https://docs.z.ai/devpack/overview", "id": "zai_coding_plan", "model_count": 7, "name": "Z.AI Coding Plan" }, "zeldoc": { "alias_of": null, "base_url": "https://api.zeldoc.ai/v1", "catalog_only": true, "doc": "https://docs.zeldoc.ai", "id": "zeldoc", "model_count": 1, "name": "Zeldoc" }, "zenifra": { "alias_of": null, "base_url": "https://ai.zenifra.com/v1", "catalog_only": true, "doc": "https://docs.zenifra.com", "id": "zenifra", "model_count": 1, "name": "Zenifra" }, "zenmux": { "alias_of": null, "base_url": "https://zenmux.ai/api/v1", "catalog_only": false, "doc": "https://docs.zenmux.ai", "id": "zenmux", "model_count": 237, "name": "Zenmux" }, "zhipuai": { "alias_of": null, "base_url": "https://open.bigmodel.cn/api/paas/v4", "catalog_only": true, "doc": "https://docs.z.ai/guides/overview/pricing", "id": "zhipuai", "model_count": 17, "name": "Zhipu AI" }, "zhipuai_coding_plan": { "alias_of": null, "base_url": "https://open.bigmodel.cn/api/coding/paas/v4", "catalog_only": true, "doc": "https://docs.bigmodel.cn/cn/coding-plan/overview", "id": "zhipuai_coding_plan", "model_count": 4, "name": "Zhipu AI Coding Plan" } }, "snapshot_id": "054d04e3201b103c8790af90c53a06c4af7ba633b579646c6534370d4e9844e7", "snapshot_schema_version": 1 };