@agentjido/llmdb 2026.9.2 → 2026.9.4
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/full.js +219 -215
- package/dist/generated/compact-keys.js +1 -1
- package/dist/generated/compact-strings.js +1 -1
- package/dist/generated/manifest.d.ts +1 -1
- package/dist/generated/manifest.js +1 -1
- package/dist/generated/provider-loaders.js +4 -0
- package/dist/providers/302ai.js +1 -1
- package/dist/providers/abacus.js +1 -1
- package/dist/providers/abliteration_ai.js +1 -1
- package/dist/providers/above.js +1 -1
- package/dist/providers/agentrouter.js +1 -1
- package/dist/providers/agnes.js +1 -1
- package/dist/providers/ai21.d.ts +11 -0
- package/dist/providers/ai21.js +8 -0
- package/dist/providers/ai_router.js +1 -1
- package/dist/providers/aiand.js +1 -1
- package/dist/providers/aihubmix.js +1 -1
- package/dist/providers/aixy.js +1 -1
- package/dist/providers/aki_io.js +1 -1
- package/dist/providers/alibaba.d.ts +1 -1
- package/dist/providers/alibaba.js +1 -1
- package/dist/providers/alibaba_cn.d.ts +1 -1
- package/dist/providers/alibaba_cn.js +1 -1
- package/dist/providers/alibaba_coding_plan.js +1 -1
- package/dist/providers/alibaba_coding_plan_cn.js +1 -1
- package/dist/providers/alibaba_token_plan.d.ts +1 -1
- package/dist/providers/alibaba_token_plan.js +1 -1
- package/dist/providers/alibaba_token_plan_cn.js +1 -1
- package/dist/providers/amazon_bedrock.js +1 -1
- package/dist/providers/ambient.js +1 -1
- package/dist/providers/amd.js +1 -1
- package/dist/providers/anthropic.js +1 -1
- package/dist/providers/anyapi.js +1 -1
- package/dist/providers/arcee.js +1 -1
- package/dist/providers/atomic_chat.js +1 -1
- package/dist/providers/auriko.js +1 -1
- package/dist/providers/azure.js +1 -1
- package/dist/providers/azure_cognitive_services.js +1 -1
- package/dist/providers/bailing.js +1 -1
- package/dist/providers/baseten.js +1 -1
- package/dist/providers/berget.js +1 -1
- package/dist/providers/blueclaw.js +1 -1
- package/dist/providers/bothub.js +1 -1
- package/dist/providers/cerebras.js +1 -1
- package/dist/providers/chutes.js +1 -1
- package/dist/providers/clarifai.js +1 -1
- package/dist/providers/claudinio.js +1 -1
- package/dist/providers/cline_pass.js +1 -1
- package/dist/providers/cloudferro_sherlock.js +1 -1
- package/dist/providers/cloudflare_ai_gateway.js +1 -1
- package/dist/providers/cloudflare_workers_ai.d.ts +1 -1
- package/dist/providers/cloudflare_workers_ai.js +1 -1
- package/dist/providers/cohere.js +1 -1
- package/dist/providers/coralbricks.js +1 -1
- package/dist/providers/cortecs.js +1 -1
- package/dist/providers/crof.js +1 -1
- package/dist/providers/crossmodel.js +1 -1
- package/dist/providers/crusoe.js +1 -1
- package/dist/providers/daoxe.js +1 -1
- package/dist/providers/databricks.js +1 -1
- package/dist/providers/deepinfra.js +1 -1
- package/dist/providers/deepseek.js +1 -1
- package/dist/providers/digitalocean.js +1 -1
- package/dist/providers/dinference.js +1 -1
- package/dist/providers/drun.js +1 -1
- package/dist/providers/ebcloud.js +1 -1
- package/dist/providers/echo.js +1 -1
- package/dist/providers/edenai.js +1 -1
- package/dist/providers/elevenlabs.js +1 -1
- package/dist/providers/empiriolabs.js +1 -1
- package/dist/providers/evroc.js +1 -1
- package/dist/providers/fastrouter.js +1 -1
- package/dist/providers/fireworks_ai.d.ts +1 -1
- package/dist/providers/fireworks_ai.js +1 -1
- package/dist/providers/freemodel.js +1 -1
- package/dist/providers/friendli.js +1 -1
- package/dist/providers/frogbot.js +1 -1
- package/dist/providers/github_copilot.js +1 -1
- package/dist/providers/github_models.js +1 -1
- package/dist/providers/gitlab.js +1 -1
- package/dist/providers/gmicloud.js +1 -1
- package/dist/providers/google.d.ts +1 -1
- package/dist/providers/google.js +1 -1
- package/dist/providers/google_vertex.d.ts +1 -1
- package/dist/providers/google_vertex.js +1 -1
- package/dist/providers/google_vertex_anthropic.js +1 -1
- package/dist/providers/greenpt.js +1 -1
- package/dist/providers/groq.js +1 -1
- package/dist/providers/helicone.js +1 -1
- package/dist/providers/hetzner.js +1 -1
- package/dist/providers/hpc_ai.js +1 -1
- package/dist/providers/huggingface.js +1 -1
- package/dist/providers/hyper.d.ts +1 -1
- package/dist/providers/hyper.js +1 -1
- package/dist/providers/iflowcn.js +1 -1
- package/dist/providers/impossibl.js +1 -1
- package/dist/providers/inception.js +1 -1
- package/dist/providers/inceptron.js +1 -1
- package/dist/providers/inco.d.ts +11 -0
- package/dist/providers/inco.js +8 -0
- package/dist/providers/infer.js +1 -1
- package/dist/providers/inference.js +1 -1
- package/dist/providers/inferx.js +1 -1
- package/dist/providers/infomaniak.js +1 -1
- package/dist/providers/io_net.js +1 -1
- package/dist/providers/iteracompute.d.ts +1 -1
- package/dist/providers/iteracompute.js +1 -1
- package/dist/providers/jalapeno.js +1 -1
- package/dist/providers/jiekou.js +1 -1
- package/dist/providers/kenari.js +1 -1
- package/dist/providers/kilo.d.ts +1 -1
- package/dist/providers/kilo.js +1 -1
- package/dist/providers/kimi_for_coding.js +1 -1
- package/dist/providers/klokintegration.js +1 -1
- package/dist/providers/kosmik.js +1 -1
- package/dist/providers/kuae_cloud_coding_plan.js +1 -1
- package/dist/providers/lilac.js +1 -1
- package/dist/providers/llama.js +1 -1
- package/dist/providers/llmgateway.js +1 -1
- package/dist/providers/llmgateway_providers.d.ts +1 -1
- package/dist/providers/llmgateway_providers.js +1 -1
- package/dist/providers/llmtech.js +1 -1
- package/dist/providers/llmtr.js +1 -1
- package/dist/providers/lmstudio.js +1 -1
- package/dist/providers/longcat.js +1 -1
- package/dist/providers/lucidquery.js +1 -1
- package/dist/providers/lynkr.js +1 -1
- package/dist/providers/meganova.js +1 -1
- package/dist/providers/melious.js +1 -1
- package/dist/providers/merge_gateway.js +1 -1
- package/dist/providers/meta.js +1 -1
- package/dist/providers/minimax.js +1 -1
- package/dist/providers/minimax_cn.js +1 -1
- package/dist/providers/minimax_cn_coding_plan.js +1 -1
- package/dist/providers/minimax_coding_plan.js +1 -1
- package/dist/providers/mistral.js +1 -1
- package/dist/providers/mixlayer.js +1 -1
- package/dist/providers/moark.js +1 -1
- package/dist/providers/modal.js +1 -1
- package/dist/providers/model_oracle_ai.js +1 -1
- package/dist/providers/modelis.js +1 -1
- package/dist/providers/modelscope.js +1 -1
- package/dist/providers/moonshotai.js +1 -1
- package/dist/providers/moonshotai_cn.js +1 -1
- package/dist/providers/morph.js +1 -1
- package/dist/providers/nan.js +1 -1
- package/dist/providers/nano_gpt.d.ts +1 -1
- package/dist/providers/nano_gpt.js +1 -1
- package/dist/providers/nearai.js +1 -1
- package/dist/providers/nebius.js +1 -1
- package/dist/providers/neon.js +1 -1
- package/dist/providers/neosmith.js +1 -1
- package/dist/providers/neuralwatt.js +1 -1
- package/dist/providers/nova.js +1 -1
- package/dist/providers/novita_ai.js +1 -1
- package/dist/providers/nvidia.d.ts +1 -1
- package/dist/providers/nvidia.js +1 -1
- package/dist/providers/oci.d.ts +11 -0
- package/dist/providers/oci.js +8 -0
- package/dist/providers/ofox.js +1 -1
- package/dist/providers/ollama_cloud.js +1 -1
- package/dist/providers/openai.js +1 -1
- package/dist/providers/opencode.d.ts +1 -1
- package/dist/providers/opencode.js +1 -1
- package/dist/providers/opencode_go.d.ts +1 -1
- package/dist/providers/opencode_go.js +1 -1
- package/dist/providers/openreason.js +1 -1
- package/dist/providers/openrouter.d.ts +1 -1
- package/dist/providers/openrouter.js +1 -1
- package/dist/providers/opper.js +1 -1
- package/dist/providers/orcarouter.js +1 -1
- package/dist/providers/ovhcloud.js +1 -1
- package/dist/providers/pendra.js +1 -1
- package/dist/providers/perplexity.js +1 -1
- package/dist/providers/perplexity_agent.js +1 -1
- package/dist/providers/pioneer.js +1 -1
- package/dist/providers/poe.js +1 -1
- package/dist/providers/poolside.js +1 -1
- package/dist/providers/privatemode_ai.d.ts +1 -1
- package/dist/providers/privatemode_ai.js +1 -1
- package/dist/providers/qihang_ai.js +1 -1
- package/dist/providers/qiniu_ai.js +1 -1
- package/dist/providers/qvac.js +1 -1
- package/dist/providers/regolo_ai.js +1 -1
- package/dist/providers/requesty.js +1 -1
- package/dist/providers/routing_run.js +1 -1
- package/dist/providers/runinfra.js +1 -1
- package/dist/providers/sakana.js +1 -1
- package/dist/providers/salad_cloud.js +1 -1
- package/dist/providers/sap_ai_core.js +1 -1
- package/dist/providers/sarvam.js +1 -1
- package/dist/providers/scaleway.js +1 -1
- package/dist/providers/scnet_token_plan.js +1 -1
- package/dist/providers/scx_ai.js +1 -1
- package/dist/providers/sensenova.js +1 -1
- package/dist/providers/siliconflow.js +1 -1
- package/dist/providers/siliconflow_cn.js +1 -1
- package/dist/providers/snowflake_cortex.js +1 -1
- package/dist/providers/stackit.js +1 -1
- package/dist/providers/standardcompute.js +1 -1
- package/dist/providers/stepfun.js +1 -1
- package/dist/providers/stepfun_ai.js +1 -1
- package/dist/providers/stepfun_ai_step_plan.js +1 -1
- package/dist/providers/stepfun_step_plan.js +1 -1
- package/dist/providers/subconscious.js +1 -1
- package/dist/providers/submodel.js +1 -1
- package/dist/providers/synthetic.d.ts +1 -1
- package/dist/providers/synthetic.js +1 -1
- package/dist/providers/tencent_coding_plan.js +1 -1
- package/dist/providers/tencent_token_plan.js +1 -1
- package/dist/providers/tencent_tokenhub.js +1 -1
- package/dist/providers/tensorx.d.ts +1 -1
- package/dist/providers/tensorx.js +1 -1
- package/dist/providers/the_grid_ai.js +1 -1
- package/dist/providers/thinkingmachines.js +1 -1
- package/dist/providers/tinfoil.d.ts +1 -1
- package/dist/providers/tinfoil.js +1 -1
- package/dist/providers/togetherai.js +1 -1
- package/dist/providers/tokengo.js +1 -1
- package/dist/providers/tokenrouter.js +1 -1
- package/dist/providers/trustedrouter.js +1 -1
- package/dist/providers/typesafe.d.ts +11 -0
- package/dist/providers/typesafe.js +8 -0
- package/dist/providers/umans_ai.d.ts +1 -1
- package/dist/providers/umans_ai.js +1 -1
- package/dist/providers/umans_ai_coding_plan.d.ts +1 -1
- package/dist/providers/umans_ai_coding_plan.js +1 -1
- package/dist/providers/unorouter.js +1 -1
- package/dist/providers/upstage.js +1 -1
- package/dist/providers/v0.js +1 -1
- package/dist/providers/vancine.d.ts +1 -1
- package/dist/providers/vancine.js +1 -1
- package/dist/providers/venice.d.ts +1 -1
- package/dist/providers/venice.js +1 -1
- package/dist/providers/vercel.d.ts +1 -1
- package/dist/providers/vercel.js +1 -1
- package/dist/providers/vispark.js +1 -1
- package/dist/providers/vivgrid.js +1 -1
- package/dist/providers/volcengine.d.ts +1 -1
- package/dist/providers/volcengine.js +1 -1
- package/dist/providers/volcengine_coding_plan.js +1 -1
- package/dist/providers/vultr.js +1 -1
- package/dist/providers/wafer_ai.js +1 -1
- package/dist/providers/wallaby.js +1 -1
- package/dist/providers/wandb.js +1 -1
- package/dist/providers/watsonx.js +1 -1
- package/dist/providers/xai.js +1 -1
- package/dist/providers/xiaomi.js +1 -1
- package/dist/providers/xiaomi_token_plan_ams.js +1 -1
- package/dist/providers/xiaomi_token_plan_cn.js +1 -1
- package/dist/providers/xiaomi_token_plan_sgp.js +1 -1
- package/dist/providers/xpersona.js +1 -1
- package/dist/providers/zai.js +1 -1
- package/dist/providers/zai_coder.js +1 -1
- package/dist/providers/zai_coding_plan.js +1 -1
- package/dist/providers/zeldoc.js +1 -1
- package/dist/providers/zenifra.js +1 -1
- package/dist/providers/zenmux.d.ts +1 -1
- package/dist/providers/zenmux.js +1 -1
- package/dist/providers/zhipuai.js +1 -1
- package/dist/providers/zhipuai_coding_plan.js +1 -1
- package/dist/snapshot.js +436 -428
- package/dist/types.d.ts +3 -0
- package/package.json +1 -1
|
@@ -1,2 +1,2 @@
|
|
|
1
1
|
// Generated by scripts/generate.mjs. Do not edit.
|
|
2
|
-
export const compactKeys = ["output", "family", "input", "enabled", "provider", "release_date", "id", "base_url", "name", "extra", "cost", "pricing", "aliases", "capabilities", "deprecated", "knowledge", "last_updated", "lifecycle", "limits", "modalities", "model", "provider_model_id", "retired", "tags", "context", "description", "attachment", "open_weights", "reasoning", "streaming", "strict", "temperature", "chat", "embeddings", "json", "rerank", "tools", "schema", "tool_calls", "native", "catalog_only", "reasoning_options", "type", "structured_output", "cache_read", "values", "supported", "wire_protocol", "path", "text", "arena", "category", "elo", "rank", "win_rate", "cache_write", "execution", "object", "interleaved", "field", "repo", "source", "architecture", "context_length", "discovered", "gguf_sources", "hf_downloads", "hf_likes", "llmfit", "matched_hugging_face_id", "memory", "min_ram_gb", "min_vram_gb", "model_id", "parameter_count", "parameters_raw", "pipeline_tag", "quantization", "recommended_ram_gb", "use_case", "hugging_face_id", "details", "expiration_date", "knowledge_cutoff", "links", "supported_voices", "active_experts", "active_parameters", "is_moe", "moe", "num_experts", "status", "parallel", "created", "owned_by", "min", "npm", "mandatory", "kind", "per", "rate", "unit", "benchmarks", "design_arena", "max", "config_schema", "env", "doc", "alias_of", "exclude_models", "models", "pricing_defaults", "api", "components", "currency", "merge", "agentic_index", "artificial_analysis", "coding_index", "intelligence_index", "tool", "default_effort", "supported_efforts", "retires_at", "default_enabled", "input_audio", "transport", "replacement", "image", "protocol", "wire", "caching", "deprecated_at", "supported_generation_methods", "thinking", "body", "version", "experimental", "modes", "features", "min_dimensions", "size_class", "max_dimensions", "default_dimensions", "top_k", "top_p", "headers", "max_temperature", "shape", "embed", "notes", "fast", "audio", "constraints", "reasoning_effort", "service_tier", "token_limit_key", "doc_url", "glm-5.2", "openai/gpt-oss-120b", "deepseek-v4-flash", "deepseek-v4-pro", "auth", "default_headers", "default_query", "kimi-k3", "runtime", "batch", "glm-5.3", "kimi-k2.5", "kimi-k2.6", "transcription", "code_execution", "context_management", "effort", "types", "citations", "glm-5.3-flash", "claude-opus-4-8", "claude-sonnet-4-6", "kimi-k2.7-code", "claude-fable-5", "deepseek/deepseek-v4-pro", "glm-5.1", "gpt-5.4", "gpt-oss-120b", "openai/gpt-5.5", "gpt-5.5", "aws_lifecycle", "claude-sonnet-5", "eol_date", "gpt-5.6-luna", "gpt-5.6-sol", "mode", "openai/gpt-5.4", "openai/gpt-oss-20b", "output_audio", "claude-opus-4-7", "claude-opus-5", "deepseek/deepseek-v4-flash", "glm-5", "gpt-5.6-terra", "openai/gpt-5", "openai/gpt-5-mini", "openai/gpt-5.4-mini", "openai/gpt-5.4-nano", "claude-opus-4-6", "gemini-2.5-flash", "openai/gpt-4.1", "openai/gpt-4.1-mini", "openai/gpt-5.2", "required", "anthropic-beta", "applies_when", "deepseek-v4-flash-0731", "gemini-2.5-pro", "gpt-5-mini", "gpt-5.4-mini", "moonshotai/kimi-k3", "openai/gpt-5-nano", "openai/gpt-5.1", "speed", "anthropic/claude-fable-5", "anthropic/claude-sonnet-5", "gemini-3.5-flash", "google/gemma-4-31b-it", "minimax/minimax-m2.5", "moonshotai/Kimi-K2.6", "openai/gpt-4o-mini", "openai/gpt-5.6-luna", "openai/gpt-5.6-sol", "openai/gpt-5.6-terra", "qwen3.6-plus", "qwen3.7-max", "qwen3.8-max", "zai-org/GLM-5.2", "anthropic/claude-opus-5", "deepseek-v3.2", "deepseek-v4-pro-0813", "google/gemini-2.5-flash", "google/gemini-2.5-pro", "google/gemini-3.5-flash", "moonshotai/kimi-k2.6", "openai/gpt-4.1-nano", "openai/gpt-4o", "openai/gpt-5.3-codex", "openai/gpt-5.5-pro", "openai/o4-mini", "priority", "qwen3.7-plus", "claude-haiku-4-5", "default_type", "disable_supported", "google/gemini-2.5-flash-lite", "google/gemini-3.1-flash-lite", "google/gemini-3.1-pro-preview", "gpt-5", "gpt-5-nano", "gpt-5.1", "gpt-5.2", "gpt-5.3-codex", "gpt-5.4-nano", "minimax/minimax-m2.7", "moonshotai/kimi-k2.7-code", "openai/gpt-5.4-pro", "openai/o3", "openai/o3-mini", "realtime", "tool_call", "adaptive", "anthropic/claude-opus-4.7", "anthropic/claude-opus-4.8", "anthropic/claude-sonnet-4.6", "clear_thinking_20251015", "clear_tool_uses_20250919", "compact_20260112", "deepseek-ai/DeepSeek-V4-Pro", "deepseek-v4.1-flash", "encrypted_supported", "gemini-3.7-flash", "gpt-4.1", "gpt-4.1-mini", "high", "image_input", "low", "medium", "MiniMax-M2.5", "minimax/minimax-m2.1", "minimax/minimax-m3", "moonshotai/kimi-k2.5", "pdf_input", "provider_capabilities", "qwen/qwen3.5-397b-a17b", "qwen/qwen3.7-max", "qwen/qwen3.7-plus", "qwen3.6-flash", "qwen3.8-flash", "raw_output_supported", "speech", "structured_outputs", "summary_supported", "xhigh", "z-ai/glm-5", "z-ai/glm-5.2", "anthropic/claude-opus-4.6", "claude-haiku-4-5-20251001", "deepseek/deepseek-v3.2", "deepseek/deepseek-v4-flash-0731", "gemini-3-flash-preview", "gemini-3.1-flash-lite", "gemini-3.1-pro-preview", "gemini-3.5-flash-lite", "gemini-3.6-flash", "glm-4.7", "google/gemini-3-flash-preview", "google/gemini-3.5-flash-lite", "google/gemini-3.6-flash", "google/gemma-4-31B-it", "gpt-4.1-nano", "gpt-4o", "gpt-5.1-codex", "gpt-5.2-codex", "gpt-oss-20b", "grok-4.5", "llama-3.3-70b-instruct", "mimo-v2.5-pro", "minimax-m3", "minimax/minimax-m2", "openai/gpt-3.5-turbo", "openai/gpt-4-turbo", "openai/gpt-5-pro", "openai/gpt-6-astra", "qwen/qwen3-max", "qwen/qwen3.5-122b-a10b", "qwen/qwen3.6-flash", "qwen/qwen3.8-27b", "supports_max_tokens", "z-ai/glm-5.1", "zai-org/GLM-5.3-Flash", "anthropic/claude-haiku-4.5", "anthropic/claude-opus-4-7", "anthropic/claude-opus-4.5", "anthropic/claude-sonnet-4-6", "anthropic/claude-sonnet-4.5", "claude-fable-5-1", "claude-sonnet-4-5", "deepseek-ai/DeepSeek-V3.1", "deepseek-ai/DeepSeek-V4-Flash", "deepseek/deepseek-v4-flash-vision-exp", "deepseek/deepseek-v4.1-flash", "gemini-3.8-flash", "gpt-6-astra", "grok-4.6", "kimi-k2-thinking", "meta/muse-spark-1.1", "meta/muse-spark-1.2", "minimax-m2.5", "minimax-m2.7", "MiniMax-M3", "moonshotai/Kimi-K2.7-Code", "moonshotai/Kimi-K3", "openai/gpt-5.1-codex-mini", "openai/gpt-5.2-codex", "openai/o1", "Qwen/Qwen3-235B-A22B-Instruct-2507", "qwen/qwen3.6-plus", "qwen3.5-397b-a17b", "thinkingmachines/inkling", "xiaomi/mimo-v2.5", "z-ai/glm-4.7", "z-ai/glm-5.3-flash", "anthropic/claude-opus-4-6", "charge_scope", "claude-opus-4-5", "deepseek-ai/DeepSeek-V4-Flash-0731", "deepseek/deepseek-v4-pro-0813", "gemini-2.5-flash-lite", "glm-4.6", "google/gemini-3.7-flash", "google/gemini-3.8-flash", "google/gemma-3-27b-it", "google/gemma-4-26b-a4b-it", "gpt-4o-mini", "gpt-5.1-codex-mini", "input_tokens", "mimo-v2.5", "minimax/minimax-m2.7-highspeed", "MiniMaxAI/MiniMax-M2.5", "MiniMaxAI/MiniMax-M3", "moonshotai/kimi-k2-thinking", "nvidia/nemotron-3-super-120b-a12b", "o3", "o4-mini", "openai/gpt-5.1-codex", "openai/gpt-5.2-pro", "qwen/qwen3-coder-next", "qwen/qwen3-coder-plus", "qwen/qwen3-next-80b-a3b-instruct", "qwen/qwen3-vl-235b-a22b-instruct", "qwen/qwen3.5-27b", "qwen/qwen3.5-35b-a3b", "Qwen/Qwen3.6-27B", "qwen3-32b", "qwen3-coder-30b-a3b-instruct", "qwen3.6-35b-a3b", "sakana/fugu-ultra", "tencent/hy3", "z-ai/glm-5.3", "zai-org/GLM-5.1", "anthropic/claude-fable-5.1", "anthropic/claude-haiku-4-5", "anthropic/claude-sonnet-4", "claude-opus-4-1", "claude-sonnet-4-5-20250929", "deepseek-ai/DeepSeek-V3.2", "deepseek-r1", "forced_choice", "glm-5v-turbo", "google/gemma-3-12b-it", "gpt-5-codex", "gpt-5-pro", "grok-4-1-fast-non-reasoning", "grok-4-1-fast-reasoning", "grok-4.3", "long_context_threshold", "meta-llama/Llama-3.3-70B-Instruct", "MiniMax-M2.7", "MiniMaxAI/MiniMax-M2.7", "o3-mini", "qwen/qwen3-coder-flash", "qwen/qwen3-next-80b-a3b-thinking", "qwen/qwen3-vl-235b-a22b-thinking", "Qwen/Qwen3.5-397B-A17B", "qwen/qwen3.6-27b", "Qwen/Qwen3.6-35B-A3B", "qwen/qwen3.8-flash", "qwen/qwen3.8-max", "qwen3-235b-a22b-instruct-2507", "qwen3-max", "qwen3-next-80b-a3b-instruct", "qwen3.5-9b", "qwen3.5-plus", "thinkingmachines/inkling-small", "x-ai/grok-4.3", "xai/grok-4.3", "xiaomi/mimo-v2.5-pro", "z-ai/glm-4.6", "z-ai/glm-5-turbo", "zai-org/GLM-5", "anthropic/claude-opus-4", "anthropic/claude-opus-4-8", "anthropic/claude-opus-4.1", "anthropic/claude-sonnet-4-5", "claude-opus-4-1-20250805", "claude-opus-4-5-20251101", "cohere/command-r-plus-08-2024", "deepseek-ai/DeepSeek-R1-0528", "deepseek-ai/DeepSeek-V3", "deepseek-ai/DeepSeek-V4-Pro-0813", "deepseek-r1-0528", "deepseek/deepseek-chat", "deepseek/deepseek-r1-0528", "deepseek/deepseek-v3.1", "gemma-4-26b-a4b-it", "gemma-4-31b-it", "glm-4.5", "glm-5-turbo", "google/gemini-3.1-flash-lite-preview", "google/gemini-3.1-pro-preview-customtools", "google/gemma-3-4b-it", "gpt-5.1-codex-max", "gpt-5.4-pro", "hy3", "meta/muse-spark-1.3", "meter", "mimo-v2-pro", "MiniMax-M2.1", "minimax/minimax-m2.5-highspeed", "moonshotai/kimi-k2-0905", "moonshotai/Kimi-K2.5", "nvidia/nemotron-3-ultra-550b-a55b", "openai/gpt-4", "openai/gpt-4o-2024-08-06", "openai/gpt-4o-2024-11-20", "openai/gpt-5.1-codex-max", "openai/gpt-oss-safeguard-20b", "perplexity/sonar", "perplexity/sonar-pro", "perplexity/sonar-reasoning-pro", "poolside/laguna-s-2.1", "qwen/qwen-plus", "Qwen/Qwen3-235B-A22B-Thinking-2507", "Qwen/Qwen3-30B-A3B-Instruct-2507", "Qwen/Qwen3-32B", "qwen/qwen3-coder-30b-a3b-instruct", "qwen/qwen3.5-9b", "Qwen/Qwen3.5-9B", "qwen/qwen3.5-plus", "qwen/qwen3.6-max-preview", "qwen/qwen3.7-flash", "qwen/qwen3.8-max-0902", "qwen3-coder-plus", "qwen3.8-27b", "tencent/hy4-preview", "x-ai/grok-4.5", "x-ai/grok-4.6", "x-ai/grok-build-0.1", "xai/grok-4.6", "z-ai/glm-5v-turbo", "deepseek-ai/DeepSeek-R1", "deepseek-ai/DeepSeek-V3-0324", "deepseek-v3", "deepseek-v4-flash-vision-exp", "deepseek/deepseek-r1", "deepseek/deepseek-v3.1-terminus", "deepseek/deepseek-v3.2-exp", "gemini-2.5-flash-image", "gemini-3-pro-preview", "gemini-3.1-flash-lite-preview", "glm-4.5-air", "glm-4.5v", "glm-4.6v", "glm-4.7-flash", "google/gemini-2.5-flash-image", "google/gemini-3-pro-image", "google/gemini-3.1-flash-image", "google/gemini-3.1-flash-image-preview", "grok-4-fast-non-reasoning", "grok-4-fast-reasoning", "grok-code-fast-1", "hy4-preview", "inclusionai/ling-3.0-flash", "meta/muse-glimmer-30b", "meta/muse-spark-1.2-contributor", "meta/muse-spark-1.3-contributor", "MiniMax-M2", "moonshot/kimi-k3", "moonshotai/kimi-k2.7-code-highspeed", "nvidia/nemotron-3-nano-30b-a3b", "o1", "openai/gpt-image-2", "openai/o1-pro", "openai/o3-pro", "poolside/laguna-xs-2.1", "pro", "qwen/qwen3-235b-a22b-instruct-2507", "qwen/qwen3-235b-a22b-thinking-2507", "qwen/qwen3-coder", "qwen/qwen3-coder-480b-a35b-instruct", "Qwen/Qwen3.5-122B-A10B", "Qwen/Qwen3.5-35B-A3B", "qwen/qwen3.5-flash", "qwen/qwen3.6-35b-a3b", "qwen/qwen3.8-2.4t-a95b", "Qwen/Qwen3.8-27B", "qwen3-coder-480b-a35b-instruct", "qwen3-coder-next", "qwen3.6-27b", "qwen3.6-max-preview", "sakana/fugu-max", "step-3.7-flash", "stepfun/step-3.5-flash", "stepfun/step-3.7-flash", "thinkingmachines/Inkling", "xai/grok-4.5", "z-ai/glm-4.5", "zai-org/GLM-4.6", "zai-org/GLM-4.7", "zai-org/glm-5.2", "zai-org/GLM-5.3", "zai/glm-4.6", "zai/glm-4.7", "zai/glm-5", "zai/glm-5.1", "zai/glm-5.2", "anthropic/claude-fable-5-1", "anthropic/claude-opus-4-5", "arcee-ai/trinity-large-thinking", "availability", "baidu/ernie-4.5-vl-424b-a47b", "character_limit", "cohere/command-r-08-2024", "cohere/command-r7b-12-2024", "deepseek-ai/DeepSeek-V3.1-Terminus", "deepseek-ai/DeepSeek-V4.1-Flash", "deepseek-v3-0324", "DeepSeek-V4-Flash", "deepseek/deepseek-chat-v3.1", "default", "gemini-2.0-flash-lite", "gemini-2.5-flash-lite-preview-09-2025", "gemini-3-pro-image", "gemini-3-pro-image-preview", "gemini-3.1-flash-image", "glm-5-3-flash", "GLM-5.1", "GLM-5.2", "glm-5.2-fast", "google/gemini-3-pro-image-preview", "google/gemini-3.1-flash-lite-image", "google/gemini-flash-latest", "google/gemma-4-26B-A4B-it", "gpt-4-turbo", "gpt-5.1-chat-latest", "gpt-5.5-pro", "grok-4-6", "grok-build-0.1", "gt", "ibm-granite/granite-4.2-8b", "image-01", "inclusionai/ling-3.0-flash-vl", "inkling", "kimi-k2", "kimi-k2-6", "kimi-k2.7-code-highspeed", "languages_supported", "llama-4-maverick-17b-128e-instruct-fp8", "lte", "meta-llama/llama-3.1-8b-instruct", "meta-llama/llama-3.2-3b-instruct", "meta-llama/llama-3.3-70b-instruct", "microsoft/wizardlm-2-8x22b", "minimax-m2-7", "MiniMax-M2.5-highspeed", "MiniMax-M2.7-highspeed", "minimax/minimax-m2-her", "mistral-large-2512", "mistral-small-2603", "mistral-small-3.2-24b-instruct-2506", "mistral/devstral-2512", "mistral/mistral-large-2512", "mistralai/mistral-large-2512", "moonshot/kimi-k2.6", "moonshot/kimi-k2.7-code", "moonshotai/kimi-k2", "muse-spark-1.2", "muse-spark-1.3-contributor", "nvidia/nemotron-3.5-lightning", "openai/gpt-4o-2024-05-13", "openai/gpt-5-codex", "openai/o3-mini-high", "openai/whisper-large-v3", "qwen/qwen-2.5-72b-instruct", "qwen/qwen3-14b", "qwen/qwen3-235b-a22b", "qwen/qwen3-235b-a22b-2507", "qwen/qwen3-30b-a3b", "qwen/qwen3-32b", "Qwen/Qwen3-Coder-30B-A3B-Instruct", "Qwen/Qwen3-Coder-480B-A35B-Instruct", "Qwen/Qwen3-VL-235B-A22B-Instruct", "Qwen/Qwen3.5-27B", "qwen3-235b-a22b", "qwen3-235b-a22b-thinking-2507", "qwen3-coder-flash", "qwen3-next-80b-a3b-thinking", "qwen3-vl-235b-a22b", "qwen3-vl-plus", "qwen3.7-flash", "sakana/fugu-ultra-v2", "sonar", "sonar-pro", "step-3.5-flash", "step-3.5-flash-2603", "text-embedding-3-large", "text-embedding-3-small", "upstage/solar-pro4", "x-ai/grok-4.20", "xai/grok-4.20-0309-non-reasoning", "xai/grok-4.20-0309-reasoning", "xai/grok-build-0.1", "XiaomiMiMo/MiMo-V2.5-Pro", "z-ai/glm-4.5-air", "z-ai/glm-4.6v", "zai-org/GLM-4.5-Air", "zai/glm-4.5", "zai/glm-4.5-air", "zai/glm-5-turbo", "zai/glm-5.3", "zai/glm-5.3-flash", "aion-labs/aion-2.0", "aion-labs/aion-3.0", "aion-labs/aion-3.0-mini", "aion-labs/aion-rp-llama-3.1-8b", "alibaba/qwen3-max", "alibaba/qwen3.7-max", "alibaba/qwen3.7-plus", "alibaba/qwen3.8-max", "amazon/nova-2-lite-v1", "amazon/nova-lite-v1", "amazon/nova-pro-v1", "anthracite-org/magnum-v4-72b", "anthropic/claude-3-haiku", "anthropic/claude-opus-4-5-20251101", "auto", "baai/bge-m3", "bytedance-seed/seed-2-1-turbo", "bytedance-seed/seed-2.0-code", "bytedance-seed/seed-2.0-lite", "cache_operation", "cerebras/gpt-oss-120b", "claude-opus-4-20250514", "claude-sonnet-4", "claude-sonnet-4-20250514", "cohere/command-a", "deepseek-flash", "deepseek-r1-distill-llama-70b", "deepseek-v3.1", "deepseek-v4-1-flash", "DeepSeek-V4-Pro", "deepseek/deepseek-r1-distill-llama-70b", "devstral-2512", "fugu-ultra", "gemini-2.0-flash", "gemini-2.5-flash-preview-09-2025", "gemini-3-5-flash", "gemini-3-6-flash", "gemini-3-flash", "gemini-3.1-flash-image-preview", "gemini-3.1-pro", "gemini-3.1-pro-preview-customtools", "gemma4-31b", "glm-4.5-flash", "glm-4.7-flashx", "glm-5-2", "google/gemini-embedding-001", "google/gemini-embedding-2", "google/gemini-flash-lite-latest", "gpt-3.5-turbo-0125", "gpt-3.5-turbo-1106", "gpt-3.5-turbo-instruct", "gpt-5-4-mini", "gpt-5-5", "gpt-5-chat-latest", "gpt-5.2-chat-latest", "gpt-5.2-pro", "gpt-5.3-chat-latest", "gpt-image-2", "grok-4", "grok-4-0709", "grok-4-3", "grok-4-5", "grok-build-0-1", "gryphe/mythomax-l2-13b", "header_name", "inception/mercury-2", "inception/mercury-2.5", "inclusionai/ling-3.0-flash-fin", "inference-net/schematron-v2-small", "inference-net/schematron-v2-turbo", "kimi-k2-0905-preview", "kimi-k2-7-code", "kimi-k2-turbo-preview", "Kimi-K2.6", "kwaipilot/kat-coder-pro-v2", "kwaipilot/kat-coder-pro-v2.5", "llama-3.3-70b-versatile", "llama-4-maverick", "meituan/longcat-2.0", "mercury-2", "meta-llama/Llama-3.1-8B-Instruct", "meta-llama/llama-4-maverick", "meta-llama/Llama-4-Maverick-17B-128E-Instruct-FP8", "meta-llama/llama-4-scout", "meta/llama-3.1-8b-instruct", "meta/llama-3.3-70b-instruct", "mimo-v2-tts", "mimo-v2.5-tts", "mimo-v2.5-tts-voiceclone", "mimo-v2.5-tts-voicedesign", "MiniMax-M1", "minimax-m2-7-highspeed", "minimax-m2.1", "minimax/minimax-01", "minimax/minimax-m2.1-lightning", "ministral-14b-2512", "ministral-3b", "mistral-7b-instruct-v0.3", "mistral-large-2411", "mistral-medium-2505", "mistral-nemo", "mistral-nemo-instruct-2407", "mistral-small-2503", "mistral/mistral-large-latest", "mistralai/codestral-2508", "mistralai/devstral-2512", "mistralai/ministral-14b-2512", "mistralai/ministral-3b-2512", "mistralai/ministral-8b-2512", "mistralai/mistral-large", "mistralai/mistral-medium-3", "mistralai/mistral-medium-3.1", "mistralai/mistral-nemo", "mistralai/Mistral-Nemo-Instruct-2407", "mistralai/mistral-saba", "mistralai/mistral-small-24b-instruct-2501", "mistralai/mistral-small-3.1-24b-instruct", "mistralai/mistral-small-3.2-24b-instruct", "mistralai/mixtral-8x22b-instruct", "moonshot/kimi-k2.7-code-highspeed", "moonshotai/kimi-k2-instruct", "moonshotai/Kimi-K2-Instruct-0905", "moonshotai/Kimi-K2-Thinking", "morph/morph-v3-fast", "morph/morph-v3-large", "muse-spark-1.1", "muse-spark-1.2-contributor", "muse-spark-1.3", "o3-pro", "openai/gpt-3.5-turbo-instruct", "openai/gpt-4o-mini-transcribe", "openai/gpt-4o-transcribe", "openai/gpt-5.6-luna-pro", "openai/gpt-5.6-sol-pro", "openai/gpt-5.6-terra-pro", "openai/gpt-6-astra-pro", "openai/gpt-chat-latest", "openai/gpt-image-1", "openai/gpt-image-1-mini", "openai/gpt-image-1.5", "openai/gpt-image-2.5-flare", "openai/gpt-image-2.5-sunburst", "openai/o4-mini-high", "openai/text-embedding-3-large", "openai/text-embedding-3-small", "perceptron/perceptron-mk1", "perplexity/sonar-deep-research", "pixtral-12b-2409", "qvq-max", "qwen-flash", "qwen-max", "qwen-omni-turbo", "qwen-plus", "qwen-turbo", "qwen/qwen-flash", "qwen/qwen-max", "Qwen/Qwen2.5-72B-Instruct", "qwen/qwen2.5-vl-72b-instruct", "qwen/qwen3-30b-a3b-instruct-2507", "qwen/qwen3-8b", "Qwen/Qwen3-8B", "Qwen/Qwen3-Embedding-8B", "Qwen/Qwen3-Next-80B-A3B-Instruct", "Qwen/Qwen3-VL-235B-A22B-Thinking", "qwen/qwen3-vl-30b-a3b-instruct", "Qwen/Qwen3-VL-30B-A3B-Instruct", "qwen/qwen3-vl-30b-a3b-thinking", "qwen/qwen3-vl-8b-instruct", "qwen/qwen3-vl-plus", "Qwen/Qwen3.6-35B-A3B-FP8", "Qwen/Qwen3.7-Max", "Qwen/Qwen3.8-2.4T-A95B", "qwen2.5-vl-72b-instruct", "qwen3-30b-a3b", "qwen3.5-35b-a3b", "qwen3.8-2.4t-a95b", "Qwen3.8-27B", "sakana/sakana-namazu", "sonar-deep-research", "sonar-reasoning-pro", "source_urls", "stepfun-ai/Step-3.5-Flash", "tencent/hy3-preview", "text-embedding-ada-002", "thinkingmachines/Inkling-Small", "undi95/remm-slerp-l2-13b", "upstage/solar-pro-3", "whisper-large-v3", "writer/palmyra-x5", "x-ai/grok-4.1-fast", "x-ai/grok-4.20-multi-agent", "XiaomiMiMo/MiMo-V2.5", "z-ai/glm-4.5v", "z-ai/glm-4.7-flash", "zai-org/GLM-5.1-FP8", "zai-org/GLM-5.2-FP8", "zai/glm-4.5v", "zai/glm-4.7-flashx", "alibaba/qwen3-coder-plus", "alibaba/qwen3.6-35b-a3b", "alibaba/qwen3.6-plus", "alibaba/qwen3.7-flash", "alibaba/qwen3.8-flash", "amazon/nova-micro-v1", "amazon/nova-premier-v1", "anthropic/claude-fable-latest", "anthropic/claude-haiku-4-5-20251001", "anthropic/claude-opus-latest", "anthropic/claude-sonnet-4-5-20250929", "anthropic/claude-sonnet-latest", "apertus-70b", "azure/gpt-5.1-codex", "azure/gpt-5.1-codex-mini", "azure/gpt-5.2-codex", "baidu/ERNIE-4.5-300B-A47B", "baidu/ernie-4.5-300b-a47b-paddle", "bfl/flux-3-video", "bytedance-seed/seed-1.6", "bytedance-seed/seed-1.6-flash", "bytedance-seed/seed-2.0-mini", "ByteDance-Seed/Seed-OSS-36B-Instruct", "bytedance/doubao-seed-2.1-pro", "bytedance/doubao-seed-2.1-turbo", "bytedance/doubao-seed-character", "bytedance/seedance-2.0", "bytedance/seedance-2.0-fast", "bytedance/seedance-2.0-mini", "bytedance/seedance-2.5", "bytedance/ui-tars-1.5-7b", "cache_ttl", "claude-3-7-sonnet-20250219", "claude-3-haiku-20240307", "claude-3.5-haiku", "claude-3.7-sonnet", "claude-4.5-haiku", "claude-4.5-opus", "claude-4.5-sonnet", "claude-fable-5-1@default", "claude-fable-5.1", "claude-fable-5@default", "claude-haiku-4-5@20251001", "claude-haiku-4.5", "claude-mythos-5", "claude-opus-4-1@20250805", "claude-opus-4-5@20251101", "claude-opus-4-6@default", "claude-opus-4-7@default", "claude-opus-4-8@default", "claude-opus-4.8", "claude-opus-4@20250514", "claude-opus-5-fast", "claude-opus-5@default", "claude-sonnet-4-5@20250929", "claude-sonnet-4-6@default", "claude-sonnet-4@20250514", "claude-sonnet-5@default", "codestral-2501", "codestral-2508", "codex-mini", "cognitivecomputations/dolphin-mistral-24b-venice-edition", "cohere-command-a", "cohere-embed-v-4-0", "cohere-embed-v3-english", "cohere-embed-v3-multilingual", "cohere/command-a-03-2025", "cohere/north-mini-code:free", "cohere/rerank-v3.5", "command-a-reasoning-08-2025", "deepseek-ai/DeepSeek-V3.2-Exp", "deepseek-ai/deepseek-v4-flash", "deepseek-ai/deepseek-v4-flash-0731", "deepseek-ai/DeepSeek-V4-Flash-Vision-Exp", "deepseek-ai/deepseek-v4-pro", "deepseek-ocr-2", "deepseek-reasoner", "deepseek-v3.2-speciale", "deepseek-v4-flash:free", "deepseek/deepseek-chat-v3-0324", "deepseek/deepseek-v3-0324", "deepseek/deepseek-v4-flash-0423", "deepseek/deepseek-v4-flash-latest", "deepseek/deepseek-v4-pro-0423", "devstral-latest", "dots-studio/dots-3-note-preview:free", "doubao-1.5-pro-32k", "doubao-seed-1-6-vision-250815", "doubao-seed-2-0-code-preview-260215", "doubao-seed-2-0-lite-260428", "doubao-seed-2-0-mini-260428", "doubao-seed-2-0-pro-260215", "doubao-seed-2.0-lite", "doubao-seed-evolving", "fish-audio/s1", "fish-audio/s2-pro", "fish-audio/s2.1-pro", "fish-audio/transcribe-1", "gemini-2.5-flash-lite-preview-06-17", "gemini-2.5-flash-preview-05-20", "gemini-2.5-pro-preview-06-05", "gemini-3-1-flash-lite", "gemini-3-1-pro", "gemini-3-1-pro-preview", "gemini-3-5-flash-lite", "gemini-3-7-flash", "gemini-embedding-001", "gemini-embedding-2", "gemini-flash-latest", "gemini-pro-latest", "gemini/gemini-3-flash-preview", "gemini/gemini-3.1-pro-preview", "gemini/gemini-3.5-flash", "gemini/gemini-3.5-flash-lite", "gemma-3-27b-it", "gemma-4-31b-it:free", "gemma4", "GLM-4.7", "glm-5-1", "glm-5-3", "glm-5.2-highspeed", "glm-5.3-highspeed", "google/gemini-2.5-pro-preview", "google/gemini-3-flash", "google/gemini-3-pro-preview", "google/gemini-3.1-flash-tts-preview", "google/gemini-pro-latest", "google/gemma-2-27b-it", "google/gemma-4-E4B-it", "google/lyria-3-clip-preview", "google/lyria-3-pro-preview", "google/veo-3.1", "google/veo-3.1-fast", "gpt-3.5-turbo", "gpt-4", "gpt-4-turbo-vision", "gpt-4.1-mini-2025-04-14", "gpt-4o-2024-11-20", "gpt-4o-mini-transcribe", "gpt-4o-transcribe", "gpt-5-3-codex", "gpt-5-4-nano", "gpt-5-6-luna", "gpt-5-6-sol", "gpt-5-6-terra", "gpt-5.3-codex-spark", "gpt-5.6", "gpt-chat-latest", "gpt-image-1", "gpt-image-1.5", "gpt-oss-safeguard-120b", "gpt-oss:120b", "grok-3", "grok-3-mini", "grok-4-20-non-reasoning", "grok-4-20-reasoning", "groq/gpt-oss-120b", "groq/gpt-oss-20b", "happyhorse-1.1-i2v", "happyhorse-1.1-r2v", "happyhorse-1.1-t2v", "hy3-preview", "ibm-granite/granite-4.0-h-micro", "inclusionai/ling-2.6-1t", "inclusionai/ling-2.6-flash", "inclusionai/ling-3.0-flash-fin:free", "inclusionai/ling-3.0-flash-sante:free", "inclusionai/ling-3.0-flash-vl:free", "inclusionAI/Ling-flash-2.0", "inclusionai/ring-2.6-1t", "kimi-k2-0711-preview", "kimi-k2-0905", "kimi-k2-thinking-turbo", "Kimi-K2.5", "Kimi-K2.7-Code", "Kimi-K3", "kimi-k3-fast", "kimi/kimi-k2.5", "liquid/lfm-2.5-2.6b:free", "LiquidAI/LFM2-24B-A2B", "llama-3.1-8b-instant", "llama-3.1-8b-instruct", "llama-4-scout-17b-16e-instruct", "longcat-2.0", "mancer/weaver", "max_file_size_mb", "meta-llama/llama-3.1-70b-instruct", "meta-llama/llama-3.2-1b-instruct", "meta-llama/Llama-3.3-70B-Instruct-Turbo", "meta-llama/llama-4-maverick-17b-128e-instruct-fp8", "meta-llama/llama-guard-4-12b", "meta/llama-3.1-70b-instruct", "meta/llama-3.2-11b-vision-instruct", "meta/llama-3.2-1b-instruct", "meta/llama-3.2-3b-instruct", "meta/muse-image-1.0", "microsoft/phi-4", "mimo-v2-5", "mimo-v2-5-pro", "mimo-v2-flash", "mimo-v2-omni", "minimax-m2", "minimax-m2.7-highspeed", "minimax/minimax-h3", "minimax/minimax-h3-max", "minimax/minimax-m1", "minimax/minimax-m2.5-lightning", "minimaxai/minimax-m1-80k", "MiniMaxAI/MiniMax-M2.1", "minimaxai/minimax-m3", "ministral-3b-2512", "ministral-8b-2512", "mistral-large-latest", "mistral-medium-3.5", "mistral-medium-3.5-128b", "mistral-medium-latest", "mistral-small-2506", "mistral/codestral-latest", "mistral/devstral-medium-latest", "mistral/magistral-medium-latest", "mistral/mistral-medium-2505", "mistral/mistral-medium-latest", "mistral/mistral-small-2603", "mistral/mistral-small-latest", "mistral/voxtral-small-latest", "mistralai/Magistral-Small-2506", "mistralai/ministral-14b-instruct-2512", "mistralai/mistral-large-2407", "mistralai/mistral-large-3-675b-instruct-2512", "mistralai/mistral-medium-3-5", "mistralai/mistral-small-2603", "mistralai/Mistral-Small-3.2-24B-Instruct-2506", "mistralai/mistral-small-4-119b-2603", "mistralai/Mistral-Small-4-119B-2603", "mistralai/voxtral-small-24b-2507", "model-router", "moonshot/kimi-k2.5", "muse-glimmer-30b", "nemotron-3-nano-omni", "nemotron-3-super-120b-a12b", "nemotron-3-ultra", "nemotron-3-ultra-550b", "nemotron-3-ultra-550b-a55b", "nemotron-3-ultra-550b-a55b:free", "nex-agi/nex-n2.5-mini:free", "nex-agi/nex-n2.5-pro:free", "nousresearch/hermes-3-llama-3.1-405b", "nousresearch/hermes-3-llama-3.1-70b", "nousresearch/hermes-4-405b", "novita/deepseek-v3.2", "novita/glm-4.6", "novita/glm-4.6v", "novita/glm-4.7", "novita/glm-5", "novita/kimi-k2.6", "novita/minimax-m2.1", "nvidia-nemotron-3-nano-30b-a3b", "nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free", "nvidia/nemotron-3-super-120b-a12b:free", "nvidia/nemotron-3-ultra-550b-a55b:free", "nvidia/nemotron-3.5-content-safety", "nvidia/nemotron-3.5-content-safety:free", "nvidia/nemotron-3.5-lightning-30b-a3b", "nvidia/nemotron-3.5-lightning:free", "nvidia/nemotron-nano-12b-v2-vl", "nvidia/nemotron-nano-9b-v2", "nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B", "nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-BF16", "openai-gpt-4.1", "openai-gpt-5", "openai-gpt-5-mini", "openai-gpt-5-nano", "openai-gpt-5.2", "openai-gpt-5.4", "openai-gpt-5.5", "openai-gpt-5.6-luna", "openai-gpt-5.6-sol", "openai-gpt-5.6-terra", "openai-gpt-6-astra", "openai-gpt-oss-120b", "openai/gpt-3.5-turbo-0613", "openai/gpt-3.5-turbo-16k", "openai/gpt-4o-mini-2024-07-18", "openai/gpt-5-chat-latest", "openai/gpt-5-image", "openai/gpt-5-image-mini", "openai/gpt-5.1-chat-latest", "openai/gpt-5.2-chat", "openai/gpt-5.2-chat-latest", "openai/gpt-5.3-chat-latest", "openai/gpt-5.4-image-2", "openai/gpt-audio", "openai/gpt-audio-mini", "openai/gpt-latest", "openai/gpt-oss-safeguard-120b", "openai/gpt-realtime-1.5", "openai/gpt-transcribe", "openai/sora-2-pro", "openai/text-embedding-ada-002", "openai/whisper-1", "openai/whisper-large-v3-turbo", "openrouter/auto", "openrouter/bodybuilder", "openrouter/free", "openrouter/pareto-code", "perplexity/pplx-embed-v1-0.6b", "perplexity/pplx-embed-v1-4b", "perplexity/sonar-pro-search", "phi-4", "phi-4-mini", "phi-4-mini-reasoning", "phi-4-multimodal", "phi-4-reasoning", "phi-4-reasoning-plus", "poolside/laguna-s-2.1:free", "poolside/laguna-xs-2.1:free", "qwen-3-6-plus", "qwen-3.8-27b", "qwen-image-2.0", "qwen-image-2.0-pro", "qwen-mt-plus", "qwen-mt-turbo", "qwen-omni-turbo-realtime", "qwen-vl-max", "qwen-vl-ocr", "qwen-vl-plus", "qwen/qwen-2.5-7b-instruct", "qwen/qwen-2.5-coder-32b-instruct", "qwen/qwen-audio-3.0-tts-plus", "qwen/qwen-plus-2025-07-28", "qwen/qwen-turbo", "qwen/qwen-vl-max", "Qwen/Qwen2.5-7B-Instruct", "qwen/qwen2.5-coder-32b-instruct", "Qwen/Qwen2.5-VL-32B-Instruct", "Qwen/Qwen3-14B", "qwen/qwen3-235b-a22b-fp8", "Qwen/Qwen3-30B-A3B", "qwen/qwen3-30b-a3b-fp8", "qwen/qwen3-30b-a3b-thinking-2507", "qwen/qwen3-32b-fp8", "Qwen/Qwen3-Coder-480B-A35B-Instruct-FP8", "qwen/qwen3-embedding-4b", "qwen/qwen3-embedding-8b", "qwen/qwen3-max-thinking", "Qwen/Qwen3-VL-30B-A3B-Thinking", "qwen/qwen3-vl-32b-instruct", "Qwen/Qwen3-VL-32B-Instruct", "Qwen/Qwen3-VL-32B-Thinking", "Qwen/Qwen3-VL-8B-Instruct", "qwen/qwen3-vl-8b-thinking", "qwen/qwen3.5-flash-02-23", "qwen/qwen3.5-plus-02-15", "qwen/qwen3.5-plus-20260420", "qwen2-5-14b-instruct", "qwen2-5-32b-instruct", "qwen2-5-72b-instruct", "qwen2-5-7b-instruct", "qwen2-5-omni-7b", "qwen2-5-vl-72b-instruct", "qwen2-5-vl-7b-instruct", "qwen3-14b", "qwen3-30b-a3b-instruct-2507", "qwen3-5-35b-a3b", "qwen3-5-397b-a17b", "qwen3-5-9b", "qwen3-6-27b", "qwen3-6-35b-a3b", "qwen3-7-plus", "qwen3-8-max", "qwen3-8b", "qwen3-asr-flash", "qwen3-coder", "qwen3-embedding-8b", "qwen3-max-2026-01-23", "qwen3-max-preview", "qwen3-omni-flash", "qwen3-omni-flash-realtime", "qwen3-vl-235b-a22b-instruct", "qwen3-vl-30b-a3b", "qwen3.5-122b-a10b", "qwen3.5-27b", "qwen3.5-2b", "qwen3.6-35b", "qwen3.8-flash-next", "Qwen3.8-Max", "qwen3.8-max-preview", "qwen3guard-gen-0.6b", "qwen3guard-gen-8b", "qwq-plus", "reasoning_effort_values", "recraft/recraft-v3", "recraft/recraft-v4", "recraft/recraft-v4-pro", "recraft/recraft-v4.1", "recraft/recraft-v4.1-pro", "recraft/recraft-v4.1-utility", "recraft/recraft-v4.1-utility-pro", "rekaai/reka-edge", "rekaai/reka-flash-3", "relace/relace-apply-3", "relace/relace-search", "request", "sao10k/l3-lunaris-8b", "sao10k/l3.1-euryale-70b", "sao10k/l3.3-euryale-70b", "sarvam-105b", "seed-2-1-turbo", "step-1-32k", "step-2-16k", "step-3-7-flash", "step-tts-2", "stepaudio-2.5-asr", "stepaudio-2.5-tts", "stepfun-ai/step-3.5-flash", "stepfun-ai/Step-3.7-Flash", "supported_formats", "tencent/hunyuan-a13b-instruct", "tencent/Hunyuan-A13B-Instruct", "tencent/hy-mt2-1.8b", "tencent/hy-mt2-30b-a3b", "tencent/hy-mt2-7b", "tencent/hy-mt2-plus", "tencent/Hy3", "thedrummer/cydonia-24b-v4.1", "thedrummer/skyfall-36b-v2", "thedrummer/unslopnemo-12b", "thinkingmachines/inkling-small:free", "training_data_cutoff", "umans-coder", "umans-deepseek-v4-flash-0731", "umans-deepseek-v4-pro-0813", "umans-flash", "umans-glm-5.2", "umans-kimi-k2.7", "umans-kimi-k3", "volcengine/doubao-seed-2.0-code", "volcengine/doubao-seed-2.0-lite", "volcengine/doubao-seed-2.0-mini", "volcengine/doubao-seed-2.0-pro", "wan2.7-image", "wan2.7-image-pro", "whisper-large-v3-turbo", "x-ai/grok-4", "x-ai/grok-4-fast", "x-ai/grok-4.1-fast-non-reasoning", "x-ai/grok-code-fast-1", "x-ai/grok-imagine-image-2.0", "x-ai/grok-voice-tts-1.0", "xai/grok-4", "xai/grok-4.1-fast-non-reasoning", "xai/grok-4.1-fast-reasoning", "xiaomi/mimo-v2-flash", "xiaomimimo/mimo-v2-flash", "XiaomiMiMo/MiMo-V2-Flash", "z-ai/glm-4.7-flashx", "zai-org/glm-4.5", "zai-org/GLM-4.5", "zai-org/glm-4.5v", "zai-org/glm-4.7", "zai-org/glm-4.7-flash", "zai-org/GLM-4.7-Flash", "zai-org/glm-5.1", "zai-org/GLM-5.2-Fast", "zai/glm-4.6v", "zai/glm-4.7-flash", "zai/glm-5v-turbo"];
|
|
2
|
+
export const compactKeys = ["output", "family", "input", "enabled", "provider", "release_date", "id", "name", "base_url", "extra", "cost", "provider_model_id", "pricing", "aliases", "capabilities", "deprecated", "knowledge", "last_updated", "lifecycle", "limits", "modalities", "model", "retired", "tags", "context", "description", "attachment", "open_weights", "reasoning", "streaming", "strict", "temperature", "chat", "embeddings", "json", "rerank", "tools", "schema", "tool_calls", "native", "catalog_only", "reasoning_options", "type", "structured_output", "cache_read", "values", "supported", "wire_protocol", "path", "text", "arena", "category", "elo", "rank", "win_rate", "cache_write", "execution", "object", "interleaved", "field", "repo", "source", "architecture", "context_length", "discovered", "gguf_sources", "hf_downloads", "hf_likes", "llmfit", "matched_hugging_face_id", "memory", "min_ram_gb", "min_vram_gb", "model_id", "parameter_count", "parameters_raw", "pipeline_tag", "quantization", "recommended_ram_gb", "use_case", "hugging_face_id", "details", "expiration_date", "knowledge_cutoff", "links", "supported_voices", "active_experts", "active_parameters", "is_moe", "moe", "num_experts", "status", "parallel", "min", "created", "owned_by", "mandatory", "npm", "kind", "per", "rate", "unit", "benchmarks", "design_arena", "max", "config_schema", "env", "doc", "alias_of", "exclude_models", "models", "pricing_defaults", "api", "components", "currency", "merge", "agentic_index", "artificial_analysis", "coding_index", "intelligence_index", "tool", "default_effort", "supported_efforts", "default_enabled", "retires_at", "input_audio", "transport", "replacement", "image", "protocol", "wire", "caching", "deprecated_at", "body", "supported_generation_methods", "experimental", "modes", "thinking", "version", "features", "min_dimensions", "size_class", "max_dimensions", "top_k", "top_p", "default_dimensions", "max_temperature", "service_tier", "headers", "shape", "audio", "embed", "notes", "doc_url", "fast", "constraints", "reasoning_effort", "token_limit_key", "glm-5.2", "openai/gpt-oss-120b", "deepseek-v4-pro", "kimi-k3", "auth", "deepseek-v4-flash", "default_headers", "default_query", "runtime", "glm-5.3", "batch", "transcription", "code_execution", "context_management", "effort", "glm-5.3-flash", "kimi-k2.5", "kimi-k2.6", "types", "citations", "claude-opus-4-8", "claude-sonnet-4-6", "kimi-k2.7-code", "priority", "claude-fable-5", "deepseek/deepseek-v4-pro", "gpt-5.4", "gpt-oss-120b", "openai/gpt-5.5", "glm-5.1", "gpt-5.5", "alias_target", "aws_lifecycle", "claude-sonnet-5", "eol_date", "gpt-5.6-luna", "gpt-5.6-sol", "mode", "openai/gpt-5.4", "openai/gpt-oss-20b", "output_audio", "slug", "claude-opus-4-7", "claude-opus-5", "deepseek/deepseek-v4-flash", "gpt-5.6-terra", "openai/gpt-5", "openai/gpt-5-mini", "openai/gpt-5.4-mini", "openai/gpt-5.4-nano", "claude-opus-4-6", "gemini-2.5-flash", "glm-5", "moonshotai/kimi-k3", "openai/gpt-4.1", "openai/gpt-4.1-mini", "openai/gpt-5.2", "required", "anthropic-beta", "applies_when", "deepseek-v4-flash-0731", "gemini-2.5-pro", "gpt-5-mini", "gpt-5.4-mini", "openai/gpt-5-nano", "openai/gpt-5.1", "speed", "anthropic/claude-fable-5", "anthropic/claude-sonnet-5", "gemini-3.5-flash", "google/gemma-4-31b-it", "minimax/minimax-m2.5", "moonshotai/Kimi-K2.6", "openai/gpt-4o-mini", "openai/gpt-5.6-luna", "openai/gpt-5.6-sol", "openai/gpt-5.6-terra", "qwen3.7-max", "qwen3.8-max", "zai-org/GLM-5.2", "anthropic/claude-opus-5", "deepseek-v3.2", "deepseek-v4-pro-0813", "deepseek-v4.1-flash", "evaluate", "google/gemini-2.5-flash", "google/gemini-2.5-pro", "google/gemini-3.5-flash", "moonshotai/kimi-k2.6", "openai/gpt-4.1-nano", "openai/gpt-4o", "openai/gpt-5.3-codex", "openai/gpt-5.5-pro", "openai/o4-mini", "qwen3.6-plus", "qwen3.7-plus", "claude-haiku-4-5", "default_type", "disable_supported", "google/gemini-2.5-flash-lite", "google/gemini-3.1-flash-lite", "google/gemini-3.1-pro-preview", "gpt-5", "gpt-5-nano", "gpt-5.1", "gpt-5.2", "gpt-5.3-codex", "gpt-5.4-nano", "minimax/minimax-m2.7", "minimax/minimax-m3", "moonshotai/kimi-k2.7-code", "openai/gpt-5.4-pro", "openai/o3", "openai/o3-mini", "realtime", "tool_call", "adaptive", "anthropic/claude-opus-4.7", "anthropic/claude-opus-4.8", "anthropic/claude-sonnet-4.6", "clear_thinking_20251015", "clear_tool_uses_20250919", "compact_20260112", "deepseek-ai/DeepSeek-V4-Pro", "deepseek/deepseek-v4-flash-0731", "encrypted_supported", "gemini-3.7-flash", "gpt-4.1", "gpt-4.1-mini", "high", "image_input", "low", "medium", "MiniMax-M2.5", "minimax-m3", "minimax/minimax-m2.1", "moonshotai/kimi-k2.5", "pdf_input", "provider_capabilities", "qwen/qwen3.5-397b-a17b", "qwen/qwen3.7-max", "qwen/qwen3.7-plus", "qwen/qwen3.8-27b", "qwen3.8-flash", "raw_output_supported", "speech", "structured_outputs", "summary_supported", "xhigh", "z-ai/glm-5", "z-ai/glm-5.2", "anthropic/claude-opus-4.6", "claude-haiku-4-5-20251001", "deepseek/deepseek-v3.2", "deepseek/deepseek-v4.1-flash", "gemini-3-flash-preview", "gemini-3.1-flash-lite", "gemini-3.1-pro-preview", "gemini-3.5-flash-lite", "gemini-3.6-flash", "glm-4.7", "google/gemini-3-flash-preview", "google/gemini-3.5-flash-lite", "google/gemini-3.6-flash", "google/gemma-4-31B-it", "gpt-4.1-nano", "gpt-4o", "gpt-5.1-codex", "gpt-5.2-codex", "gpt-oss-20b", "grok-4.5", "mimo-v2.5-pro", "minimax/minimax-m2", "openai/gpt-3.5-turbo", "openai/gpt-4-turbo", "openai/gpt-5-pro", "openai/gpt-6-astra", "qwen/qwen3-max", "qwen/qwen3.5-122b-a10b", "qwen/qwen3.6-flash", "qwen3.6-flash", "supports_max_tokens", "z-ai/glm-5.1", "z-ai/glm-5.3", "z-ai/glm-5.3-flash", "zai-org/GLM-5.3-Flash", "anthropic/claude-haiku-4.5", "anthropic/claude-opus-4-7", "anthropic/claude-opus-4.5", "anthropic/claude-sonnet-4-6", "anthropic/claude-sonnet-4.5", "claude-fable-5-1", "claude-sonnet-4-5", "deepseek-ai/DeepSeek-V3.1", "deepseek-ai/DeepSeek-V4-Flash", "deepseek/deepseek-v4-pro-0813", "gemini-3.8-flash", "gpt-6-astra", "grok-4.6", "kimi-k2-thinking", "llama-3.3-70b-instruct", "meta/muse-spark-1.1", "meta/muse-spark-1.2", "minimax-m2.5", "minimax-m2.7", "MiniMax-M3", "moonshotai/Kimi-K2.7-Code", "moonshotai/Kimi-K3", "openai/gpt-5.1-codex-mini", "openai/gpt-5.2-codex", "openai/o1", "Qwen/Qwen3-235B-A22B-Instruct-2507", "qwen/qwen3.6-plus", "qwen3.5-397b-a17b", "thinkingmachines/inkling", "xiaomi/mimo-v2.5", "z-ai/glm-4.7", "anthropic/claude-opus-4-6", "charge_scope", "claude-opus-4-5", "deepseek-ai/DeepSeek-V4-Flash-0731", "deepseek/deepseek-v4-flash-vision-exp", "gemini-2.5-flash-lite", "glm-4.6", "google/gemini-3.7-flash", "google/gemini-3.8-flash", "google/gemma-3-27b-it", "google/gemma-4-26b-a4b-it", "gpt-4o-mini", "gpt-5.1-codex-mini", "input_tokens", "mimo-v2.5", "minimax/minimax-m2.7-highspeed", "MiniMaxAI/MiniMax-M2.5", "MiniMaxAI/MiniMax-M3", "moonshotai/kimi-k2-thinking", "nvidia/nemotron-3-super-120b-a12b", "o3", "o4-mini", "openai/gpt-5.1-codex", "openai/gpt-5.2-pro", "qwen/qwen3-coder-next", "qwen/qwen3-coder-plus", "qwen/qwen3-next-80b-a3b-instruct", "qwen/qwen3-vl-235b-a22b-instruct", "qwen/qwen3.5-27b", "qwen/qwen3.5-35b-a3b", "Qwen/Qwen3.6-27B", "qwen3-32b", "qwen3-coder-30b-a3b-instruct", "qwen3.6-35b-a3b", "sakana/fugu-ultra", "tencent/hy3", "zai-org/GLM-5.1", "anthropic/claude-fable-5.1", "anthropic/claude-haiku-4-5", "anthropic/claude-sonnet-4", "claude-opus-4-1", "claude-sonnet-4-5-20250929", "deepseek-ai/DeepSeek-V3.2", "deepseek-r1", "forced_choice", "glm-5v-turbo", "google/gemma-3-12b-it", "gpt-5-codex", "gpt-5-pro", "grok-4-1-fast-non-reasoning", "grok-4-1-fast-reasoning", "grok-4.3", "long_context_threshold", "meta-llama/Llama-3.3-70B-Instruct", "MiniMax-M2.7", "MiniMaxAI/MiniMax-M2.7", "o3-mini", "qwen/qwen3-coder-flash", "qwen/qwen3-next-80b-a3b-thinking", "qwen/qwen3-vl-235b-a22b-thinking", "Qwen/Qwen3.5-397B-A17B", "qwen/qwen3.6-27b", "Qwen/Qwen3.6-35B-A3B", "qwen/qwen3.8-flash", "qwen/qwen3.8-max", "qwen3-235b-a22b-instruct-2507", "qwen3-max", "qwen3.5-9b", "qwen3.5-plus", "thinkingmachines/inkling-small", "x-ai/grok-4.3", "xai/grok-4.3", "xiaomi/mimo-v2.5-pro", "z-ai/glm-4.6", "z-ai/glm-5-turbo", "zai-org/GLM-5", "anthropic/claude-opus-4", "anthropic/claude-opus-4-8", "anthropic/claude-opus-4.1", "anthropic/claude-sonnet-4-5", "claude-opus-4-1-20250805", "claude-opus-4-5-20251101", "cohere/command-r-plus-08-2024", "deepseek-ai/DeepSeek-R1-0528", "deepseek-ai/DeepSeek-V3", "deepseek-ai/DeepSeek-V4-Pro-0813", "deepseek-r1-0528", "deepseek/deepseek-chat", "deepseek/deepseek-r1-0528", "deepseek/deepseek-v3.1", "gemma-4-26b-a4b-it", "gemma-4-31b-it", "glm-4.5", "glm-5-turbo", "google/gemini-3.1-flash-lite-preview", "google/gemini-3.1-pro-preview-customtools", "google/gemma-3-4b-it", "gpt-5.1-codex-max", "gpt-5.4-pro", "hy3", "meta/muse-spark-1.3", "meter", "mimo-v2-pro", "MiniMax-M2.1", "minimax/minimax-m2.5-highspeed", "moonshotai/kimi-k2-0905", "moonshotai/Kimi-K2.5", "nvidia/nemotron-3-ultra-550b-a55b", "openai/gpt-4", "openai/gpt-4o-2024-08-06", "openai/gpt-4o-2024-11-20", "openai/gpt-5.1-codex-max", "openai/gpt-oss-safeguard-20b", "perplexity/sonar", "perplexity/sonar-pro", "perplexity/sonar-reasoning-pro", "poolside/laguna-s-2.1", "qwen/qwen-plus", "Qwen/Qwen3-235B-A22B-Thinking-2507", "Qwen/Qwen3-30B-A3B-Instruct-2507", "Qwen/Qwen3-32B", "qwen/qwen3-coder-30b-a3b-instruct", "qwen/qwen3.5-9b", "Qwen/Qwen3.5-9B", "qwen/qwen3.5-plus", "qwen/qwen3.6-max-preview", "qwen/qwen3.7-flash", "qwen/qwen3.8-2.4t-a95b", "qwen/qwen3.8-max-0902", "qwen3-coder-plus", "qwen3-next-80b-a3b-instruct", "qwen3.8-27b", "tencent/hy4-preview", "x-ai/grok-4.5", "x-ai/grok-4.6", "x-ai/grok-build-0.1", "xai/grok-4.6", "z-ai/glm-5v-turbo", "deepseek-ai/DeepSeek-R1", "deepseek-ai/DeepSeek-V3-0324", "deepseek-v3", "deepseek-v4-flash-vision-exp", "deepseek/deepseek-r1", "deepseek/deepseek-v3.1-terminus", "deepseek/deepseek-v3.2-exp", "gemini-2.5-flash-image", "gemini-3-pro-preview", "gemini-3.1-flash-lite-preview", "glm-4.5-air", "glm-4.5v", "glm-4.6v", "glm-4.7-flash", "google/gemini-2.5-flash-image", "google/gemini-3-pro-image", "google/gemini-3.1-flash-image", "google/gemini-3.1-flash-image-preview", "grok-4-fast-non-reasoning", "grok-4-fast-reasoning", "grok-code-fast-1", "hy4-preview", "inclusionai/ling-3.0-flash", "inclusionai/ling-3.0-flash-vl", "meta/muse-glimmer-30b", "meta/muse-spark-1.2-contributor", "meta/muse-spark-1.3-contributor", "MiniMax-M2", "moonshot/kimi-k3", "moonshotai/kimi-k2.7-code-highspeed", "nvidia/nemotron-3-nano-30b-a3b", "o1", "openai/gpt-image-2", "openai/o1-pro", "openai/o3-pro", "poolside/laguna-xs-2.1", "pro", "qwen/qwen3-235b-a22b-instruct-2507", "qwen/qwen3-235b-a22b-thinking-2507", "qwen/qwen3-coder", "qwen/qwen3-coder-480b-a35b-instruct", "Qwen/Qwen3.5-122B-A10B", "Qwen/Qwen3.5-35B-A3B", "qwen/qwen3.5-flash", "qwen/qwen3.6-35b-a3b", "Qwen/Qwen3.8-27B", "qwen3-coder-480b-a35b-instruct", "qwen3-coder-next", "qwen3.6-27b", "qwen3.6-max-preview", "sakana/fugu-max", "step-3.7-flash", "stepfun/step-3.5-flash", "stepfun/step-3.7-flash", "thinkingmachines/Inkling", "xai/grok-4.5", "z-ai/glm-4.5", "zai-org/GLM-4.6", "zai-org/GLM-4.7", "zai-org/glm-5.2", "zai-org/GLM-5.3", "zai/glm-4.6", "zai/glm-4.7", "zai/glm-5", "zai/glm-5.1", "zai/glm-5.2", "anthropic/claude-fable-5-1", "anthropic/claude-opus-4-5", "arcee-ai/trinity-large-thinking", "availability", "baidu/ernie-4.5-vl-424b-a47b", "character_limit", "cohere/command-r-08-2024", "cohere/command-r7b-12-2024", "deepseek-ai/DeepSeek-V3.1-Terminus", "deepseek-ai/DeepSeek-V4.1-Flash", "deepseek-v3-0324", "DeepSeek-V4-Flash", "deepseek/deepseek-chat-v3.1", "default", "gemini-2.0-flash-lite", "gemini-2.5-flash-lite-preview-09-2025", "gemini-3-pro-image", "gemini-3-pro-image-preview", "gemini-3.1-flash-image", "glm-5-3-flash", "GLM-5.1", "GLM-5.2", "glm-5.2-fast", "google/gemini-3-pro-image-preview", "google/gemini-3.1-flash-lite-image", "google/gemini-flash-latest", "google/gemma-4-26B-A4B-it", "gpt-4-turbo", "gpt-5.1-chat-latest", "gpt-5.5-pro", "grok-4-6", "grok-build-0.1", "gt", "ibm-granite/granite-4.2-8b", "image-01", "inference-net/schematron-v2-small", "inference-net/schematron-v2-turbo", "inkling", "kimi-k2", "kimi-k2-6", "kimi-k2.7-code-highspeed", "languages_supported", "lte", "meta-llama/llama-3.1-8b-instruct", "meta-llama/llama-3.2-3b-instruct", "meta-llama/llama-3.3-70b-instruct", "microsoft/wizardlm-2-8x22b", "minimax-m2-7", "MiniMax-M2.5-highspeed", "MiniMax-M2.7-highspeed", "minimax/minimax-m2-her", "mistral-large-2512", "mistral-small-2603", "mistral-small-3.2-24b-instruct-2506", "mistral/devstral-2512", "mistral/mistral-large-2512", "mistralai/mistral-large-2512", "moonshot/kimi-k2.6", "moonshot/kimi-k2.7-code", "moonshotai/kimi-k2", "muse-spark-1.2", "muse-spark-1.3-contributor", "nvidia/nemotron-3.5-lightning", "openai/gpt-4o-2024-05-13", "openai/gpt-5-codex", "openai/o3-mini-high", "openai/whisper-large-v3", "qwen/qwen-2.5-72b-instruct", "qwen/qwen3-14b", "qwen/qwen3-235b-a22b", "qwen/qwen3-235b-a22b-2507", "qwen/qwen3-30b-a3b", "qwen/qwen3-32b", "Qwen/Qwen3-Coder-30B-A3B-Instruct", "Qwen/Qwen3-Coder-480B-A35B-Instruct", "Qwen/Qwen3-VL-235B-A22B-Instruct", "Qwen/Qwen3.5-27B", "qwen3-235b-a22b", "qwen3-235b-a22b-thinking-2507", "qwen3-coder-flash", "qwen3-next-80b-a3b-thinking", "qwen3-vl-235b-a22b", "qwen3-vl-plus", "qwen3.7-flash", "sakana/fugu-ultra-v2", "sonar", "sonar-pro", "step-3.5-flash", "step-3.5-flash-2603", "text-embedding-3-large", "text-embedding-3-small", "upstage/solar-pro4", "x-ai/grok-4.20", "xai/grok-4.20-0309-non-reasoning", "xai/grok-4.20-0309-reasoning", "xai/grok-build-0.1", "XiaomiMiMo/MiMo-V2.5-Pro", "z-ai/glm-4.5-air", "z-ai/glm-4.6v", "zai-org/GLM-4.5-Air", "zai/glm-4.5", "zai/glm-4.5-air", "zai/glm-5-turbo", "zai/glm-5.3", "zai/glm-5.3-flash", "aion-labs/aion-2.0", "aion-labs/aion-3.0", "aion-labs/aion-3.0-mini", "aion-labs/aion-rp-llama-3.1-8b", "alibaba/qwen3-max", "alibaba/qwen3.7-max", "alibaba/qwen3.7-plus", "alibaba/qwen3.8-max", "amazon/nova-2-lite-v1", "amazon/nova-lite-v1", "amazon/nova-pro-v1", "anthracite-org/magnum-v4-72b", "anthropic/claude-3-haiku", "anthropic/claude-opus-4-5-20251101", "auto", "baai/bge-m3", "bytedance-seed/seed-2-1-turbo", "bytedance-seed/seed-2.0-code", "bytedance-seed/seed-2.0-lite", "cache_operation", "cerebras/gpt-oss-120b", "claude-opus-4-20250514", "claude-sonnet-4", "claude-sonnet-4-20250514", "cohere/command-a", "deepseek-r1-distill-llama-70b", "deepseek-v3.1", "deepseek-v4-1-flash", "DeepSeek-V4-Pro", "deepseek/deepseek-r1-distill-llama-70b", "devstral-2512", "fugu-ultra", "gemini-2.0-flash", "gemini-2.5-flash-preview-09-2025", "gemini-3-5-flash", "gemini-3-6-flash", "gemini-3-flash", "gemini-3.1-flash-image-preview", "gemini-3.1-pro", "gemini-3.1-pro-preview-customtools", "gemma4-31b", "glm-4.5-flash", "glm-4.7-flashx", "glm-5-2", "google/gemini-embedding-001", "google/gemini-embedding-2", "google/gemini-flash-lite-latest", "gpt-3.5-turbo-0125", "gpt-3.5-turbo-1106", "gpt-3.5-turbo-instruct", "gpt-5-4-mini", "gpt-5-5", "gpt-5-chat-latest", "gpt-5.2-chat-latest", "gpt-5.2-pro", "gpt-5.3-chat-latest", "gpt-image-2", "grok-4", "grok-4-0709", "grok-4-3", "grok-4-5", "grok-build-0-1", "gryphe/mythomax-l2-13b", "header_name", "inception/mercury-2", "inception/mercury-2.5", "inclusionai/ling-3.0-flash-fin", "kimi-k2-0905-preview", "kimi-k2-7-code", "kimi-k2-turbo-preview", "Kimi-K2.6", "llama-3.3-70b-versatile", "llama-4-maverick", "llama-4-maverick-17b-128e-instruct-fp8", "meituan/longcat-2.0", "mercury-2", "meta-llama/Llama-3.1-8B-Instruct", "meta-llama/llama-4-maverick", "meta-llama/Llama-4-Maverick-17B-128E-Instruct-FP8", "meta-llama/llama-4-scout", "meta/llama-3.1-8b-instruct", "meta/llama-3.3-70b-instruct", "mimo-v2-tts", "mimo-v2.5-tts", "mimo-v2.5-tts-voiceclone", "mimo-v2.5-tts-voicedesign", "MiniMax-M1", "minimax-m2-7-highspeed", "minimax-m2.1", "minimax/minimax-01", "minimax/minimax-m2.1-lightning", "ministral-14b-2512", "ministral-3b", "mistral-7b-instruct-v0.3", "mistral-large-2411", "mistral-medium-2505", "mistral-nemo", "mistral-nemo-instruct-2407", "mistral-small-2503", "mistral/mistral-large-latest", "mistralai/codestral-2508", "mistralai/devstral-2512", "mistralai/ministral-14b-2512", "mistralai/ministral-3b-2512", "mistralai/ministral-8b-2512", "mistralai/mistral-large", "mistralai/mistral-medium-3", "mistralai/mistral-medium-3.1", "mistralai/mistral-nemo", "mistralai/Mistral-Nemo-Instruct-2407", "mistralai/mistral-saba", "mistralai/mistral-small-24b-instruct-2501", "mistralai/mistral-small-3.1-24b-instruct", "mistralai/mistral-small-3.2-24b-instruct", "mistralai/mixtral-8x22b-instruct", "moonshot/kimi-k2.7-code-highspeed", "moonshotai/kimi-k2-instruct", "moonshotai/Kimi-K2-Instruct-0905", "moonshotai/Kimi-K2-Thinking", "morph/morph-v3-fast", "morph/morph-v3-large", "muse-spark-1.1", "muse-spark-1.2-contributor", "muse-spark-1.3", "o3-pro", "openai/gpt-3.5-turbo-instruct", "openai/gpt-4o-mini-transcribe", "openai/gpt-4o-transcribe", "openai/gpt-5.6-luna-pro", "openai/gpt-5.6-sol-pro", "openai/gpt-5.6-terra-pro", "openai/gpt-6-astra-pro", "openai/gpt-chat-latest", "openai/gpt-image-1", "openai/gpt-image-1-mini", "openai/gpt-image-1.5", "openai/gpt-image-2.5-flare", "openai/gpt-image-2.5-sunburst", "openai/o4-mini-high", "openai/text-embedding-3-large", "openai/text-embedding-3-small", "perceptron/perceptron-mk1", "perplexity/sonar-deep-research", "pixtral-12b-2409", "qvq-max", "qwen-flash", "qwen-max", "qwen-omni-turbo", "qwen-plus", "qwen-turbo", "qwen/qwen-flash", "qwen/qwen-max", "Qwen/Qwen2.5-72B-Instruct", "qwen/qwen2.5-vl-72b-instruct", "qwen/qwen3-30b-a3b-instruct-2507", "qwen/qwen3-8b", "Qwen/Qwen3-8B", "Qwen/Qwen3-Embedding-8B", "Qwen/Qwen3-Next-80B-A3B-Instruct", "Qwen/Qwen3-VL-235B-A22B-Thinking", "qwen/qwen3-vl-30b-a3b-instruct", "Qwen/Qwen3-VL-30B-A3B-Instruct", "qwen/qwen3-vl-30b-a3b-thinking", "qwen/qwen3-vl-8b-instruct", "qwen/qwen3-vl-plus", "Qwen/Qwen3.6-35B-A3B-FP8", "Qwen/Qwen3.7-Max", "Qwen/Qwen3.8-2.4T-A95B", "qwen2.5-vl-72b-instruct", "qwen3-30b-a3b", "qwen3.5-35b-a3b", "qwen3.8-2.4t-a95b", "Qwen3.8-27B", "sakana/sakana-namazu", "sonar-deep-research", "sonar-reasoning-pro", "source_urls", "stepfun-ai/Step-3.5-Flash", "tencent/hy3-preview", "text-embedding-ada-002", "thinkingmachines/Inkling-Small", "undi95/remm-slerp-l2-13b", "upstage/solar-pro-3", "whisper-large-v3", "writer/palmyra-x5", "x-ai/grok-4.1-fast", "x-ai/grok-4.20-multi-agent", "XiaomiMiMo/MiMo-V2.5", "z-ai/glm-4.5v", "z-ai/glm-4.7-flash", "zai-org/GLM-5.1-FP8", "zai-org/GLM-5.2-FP8", "zai/glm-4.5v", "zai/glm-4.7-flashx", "~anthropic/claude-fable-latest", "~anthropic/claude-haiku-latest", "~anthropic/claude-opus-latest", "~anthropic/claude-sonnet-latest", "~deepseek/deepseek-flash-latest", "~deepseek/deepseek-pro-latest", "~deepseek/deepseek-v4-flash-latest", "~google/gemini-flash-latest", "~google/gemini-pro-latest", "~moonshotai/kimi-latest", "~openai/gpt-astra-latest", "~openai/gpt-luna-latest", "~openai/gpt-mini-latest", "~openai/gpt-sol-latest", "~openai/gpt-terra-latest", "~x-ai/grok-latest", "~z-ai/glm-flash-latest", "~z-ai/glm-latest", "alibaba/qwen3-coder-plus", "alibaba/qwen3.6-35b-a3b", "alibaba/qwen3.6-plus", "alibaba/qwen3.7-flash", "alibaba/qwen3.8-flash", "amazon/nova-micro-v1", "amazon/nova-premier-v1", "anthropic/claude-fable-latest", "anthropic/claude-haiku-4-5-20251001", "anthropic/claude-opus-latest", "anthropic/claude-sonnet-4-5-20250929", "anthropic/claude-sonnet-latest", "apertus-70b", "azure/gpt-5.1-codex", "azure/gpt-5.1-codex-mini", "azure/gpt-5.2-codex", "baidu/ERNIE-4.5-300B-A47B", "baidu/ernie-4.5-300b-a47b-paddle", "bfl/flux-3-video", "bytedance-seed/seed-1.6", "bytedance-seed/seed-1.6-flash", "bytedance-seed/seed-2.0-mini", "ByteDance-Seed/Seed-OSS-36B-Instruct", "bytedance/doubao-seed-2.1-pro", "bytedance/doubao-seed-2.1-turbo", "bytedance/doubao-seed-character", "bytedance/seedance-2.0", "bytedance/seedance-2.0-fast", "bytedance/seedance-2.0-mini", "bytedance/seedance-2.5", "bytedance/ui-tars-1.5-7b", "cache_ttl", "claude-3-7-sonnet-20250219", "claude-3-haiku-20240307", "claude-3.5-haiku", "claude-3.7-sonnet", "claude-4.5-haiku", "claude-4.5-opus", "claude-4.5-sonnet", "claude-fable-5-1@default", "claude-fable-5.1", "claude-fable-5@default", "claude-haiku-4-5@20251001", "claude-haiku-4.5", "claude-mythos-5", "claude-opus-4-1@20250805", "claude-opus-4-5@20251101", "claude-opus-4-6@default", "claude-opus-4-7@default", "claude-opus-4-8@default", "claude-opus-4.8", "claude-opus-4@20250514", "claude-opus-5-fast", "claude-opus-5@default", "claude-sonnet-4-5@20250929", "claude-sonnet-4-6@default", "claude-sonnet-4@20250514", "claude-sonnet-5@default", "codestral-2501", "codestral-2508", "codex-mini", "cognitivecomputations/dolphin-mistral-24b-venice-edition", "cohere-command-a", "cohere-embed-v-4-0", "cohere-embed-v3-english", "cohere-embed-v3-multilingual", "cohere/command-a-03-2025", "cohere/north-mini-code:free", "cohere/rerank-v3.5", "command-a-reasoning-08-2025", "deepseek-ai/DeepSeek-V3.2-Exp", "deepseek-ai/deepseek-v4-flash", "deepseek-ai/deepseek-v4-flash-0731", "deepseek-ai/DeepSeek-V4-Flash-Vision-Exp", "deepseek-ai/deepseek-v4-pro", "deepseek-flash", "deepseek-ocr-2", "deepseek-reasoner", "deepseek-v3.2-speciale", "deepseek-v4-flash:free", "deepseek/deepseek-chat-v3-0324", "deepseek/deepseek-v3-0324", "deepseek/deepseek-v4-flash-0423", "deepseek/deepseek-v4-flash-latest", "deepseek/deepseek-v4-pro-0423", "devstral-latest", "dots-studio/dots-3-note-preview:free", "doubao-1.5-pro-32k", "doubao-seed-1-6-vision-250815", "doubao-seed-2-0-code-preview-260215", "doubao-seed-2-0-lite-260428", "doubao-seed-2-0-mini-260428", "doubao-seed-2-0-pro-260215", "doubao-seed-2.0-lite", "doubao-seed-evolving", "fish-audio/s1", "fish-audio/s2-pro", "fish-audio/s2.1-pro", "fish-audio/transcribe-1", "gemini-2.5-flash-lite-preview-06-17", "gemini-2.5-flash-preview-05-20", "gemini-2.5-pro-preview-06-05", "gemini-3-1-flash-lite", "gemini-3-1-pro", "gemini-3-1-pro-preview", "gemini-3-5-flash-lite", "gemini-3-7-flash", "gemini-embedding-001", "gemini-embedding-2", "gemini-flash-latest", "gemini-pro-latest", "gemini/gemini-3-flash-preview", "gemini/gemini-3.1-pro-preview", "gemini/gemini-3.5-flash", "gemini/gemini-3.5-flash-lite", "gemma-3-27b-it", "gemma-4-31b-it:free", "gemma4", "GLM-4.7", "glm-5-1", "glm-5-3", "glm-5.2-highspeed", "glm-5.3-highspeed", "google/gemini-2.5-pro-preview", "google/gemini-3-flash", "google/gemini-3-pro-preview", "google/gemini-3.1-flash-tts-preview", "google/gemini-pro-latest", "google/gemma-2-27b-it", "google/gemma-4-E4B-it", "google/lyria-3-clip-preview", "google/lyria-3-pro-preview", "google/veo-3.1", "google/veo-3.1-fast", "gpt-3.5-turbo", "gpt-4", "gpt-4-turbo-vision", "gpt-4.1-mini-2025-04-14", "gpt-4o-2024-11-20", "gpt-4o-mini-transcribe", "gpt-4o-transcribe", "gpt-5-3-codex", "gpt-5-4-nano", "gpt-5-6-luna", "gpt-5-6-sol", "gpt-5-6-terra", "gpt-5.3-codex-spark", "gpt-5.6", "gpt-chat-latest", "gpt-image-1", "gpt-image-1.5", "gpt-oss-safeguard-120b", "gpt-oss:120b", "grok-3", "grok-3-mini", "grok-4-20-non-reasoning", "grok-4-20-reasoning", "groq/gpt-oss-120b", "groq/gpt-oss-20b", "happyhorse-1.1-i2v", "happyhorse-1.1-r2v", "happyhorse-1.1-t2v", "hy3-preview", "ibm-granite/granite-4.0-h-micro", "inclusionai/ling-2.6-1t", "inclusionai/ling-2.6-flash", "inclusionai/ling-3.0-flash-fin:free", "inclusionai/ling-3.0-flash-sante:free", "inclusionai/ling-3.0-flash-vl:free", "inclusionAI/Ling-flash-2.0", "inclusionai/ring-2.6-1t", "kimi-k2-0711-preview", "kimi-k2-0905", "kimi-k2-thinking-turbo", "Kimi-K2.5", "Kimi-K2.7-Code", "Kimi-K3", "kimi-k3-fast", "kimi/kimi-k2.5", "kwaipilot/kat-coder-pro-v2", "kwaipilot/kat-coder-pro-v2.5", "liquid/lfm-2.5-2.6b:free", "LiquidAI/LFM2-24B-A2B", "llama-3.1-8b-instant", "llama-3.1-8b-instruct", "llama-4-scout-17b-16e-instruct", "longcat-2.0", "mancer/weaver", "max_file_size_mb", "meta-llama/llama-3.1-70b-instruct", "meta-llama/llama-3.2-1b-instruct", "meta-llama/Llama-3.3-70B-Instruct-Turbo", "meta-llama/llama-4-maverick-17b-128e-instruct-fp8", "meta-llama/llama-guard-4-12b", "meta/llama-3.1-70b-instruct", "meta/llama-3.2-11b-vision-instruct", "meta/llama-3.2-1b-instruct", "meta/llama-3.2-3b-instruct", "meta/muse-image-1.0", "microsoft/phi-4", "mimo-v2-5", "mimo-v2-5-pro", "mimo-v2-flash", "mimo-v2-omni", "minimax-m2", "minimax-m2.7-highspeed", "minimax/minimax-h3", "minimax/minimax-h3-max", "minimax/minimax-m1", "minimax/minimax-m2.5-lightning", "minimaxai/minimax-m1-80k", "MiniMaxAI/MiniMax-M2.1", "minimaxai/minimax-m3", "ministral-3b-2512", "ministral-8b-2512", "mistral-large-latest", "mistral-medium-3.5", "mistral-medium-3.5-128b", "mistral-medium-latest", "mistral-small-2506", "mistral/codestral-latest", "mistral/devstral-medium-latest", "mistral/magistral-medium-latest", "mistral/mistral-medium-2505", "mistral/mistral-medium-latest", "mistral/mistral-small-2603", "mistral/mistral-small-latest", "mistral/voxtral-small-latest", "mistralai/Magistral-Small-2506", "mistralai/ministral-14b-instruct-2512", "mistralai/mistral-large-2407", "mistralai/mistral-large-3-675b-instruct-2512", "mistralai/mistral-medium-3-5", "mistralai/mistral-small-2603", "mistralai/Mistral-Small-3.2-24B-Instruct-2506", "mistralai/mistral-small-4-119b-2603", "mistralai/Mistral-Small-4-119B-2603", "mistralai/voxtral-small-24b-2507", "model-router", "moonshot/kimi-k2.5", "muse-glimmer-30b", "nemotron-3-nano-omni", "nemotron-3-super-120b-a12b", "nemotron-3-ultra", "nemotron-3-ultra-550b", "nemotron-3-ultra-550b-a55b", "nemotron-3-ultra-550b-a55b:free", "nex-agi/nex-n2.5-mini:free", "nex-agi/nex-n2.5-pro:free", "nousresearch/hermes-3-llama-3.1-405b", "nousresearch/hermes-3-llama-3.1-70b", "nousresearch/hermes-4-405b", "novita/deepseek-v3.2", "novita/glm-4.6", "novita/glm-4.6v", "novita/glm-4.7", "novita/glm-5", "novita/kimi-k2.6", "novita/minimax-m2.1", "nvidia-nemotron-3-nano-30b-a3b", "nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free", "nvidia/nemotron-3-super-120b-a12b:free", "nvidia/nemotron-3-ultra-550b-a55b:free", "nvidia/nemotron-3.5-content-safety", "nvidia/nemotron-3.5-content-safety:free", "nvidia/nemotron-3.5-lightning-30b-a3b", "nvidia/nemotron-3.5-lightning:free", "nvidia/nemotron-nano-12b-v2-vl", "nvidia/nemotron-nano-9b-v2", "nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B", "nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-BF16", "openai-gpt-4.1", "openai-gpt-5", "openai-gpt-5-mini", "openai-gpt-5-nano", "openai-gpt-5.2", "openai-gpt-5.4", "openai-gpt-5.5", "openai-gpt-5.6-luna", "openai-gpt-5.6-sol", "openai-gpt-5.6-terra", "openai-gpt-6-astra", "openai-gpt-oss-120b", "openai.gpt-oss-120b", "openai.gpt-oss-20b", "openai/gpt-3.5-turbo-0613", "openai/gpt-3.5-turbo-16k", "openai/gpt-4o-mini-2024-07-18", "openai/gpt-5-chat-latest", "openai/gpt-5-image", "openai/gpt-5-image-mini", "openai/gpt-5.1-chat-latest", "openai/gpt-5.2-chat", "openai/gpt-5.2-chat-latest", "openai/gpt-5.3-chat-latest", "openai/gpt-5.4-image-2", "openai/gpt-audio", "openai/gpt-audio-mini", "openai/gpt-latest", "openai/gpt-oss-safeguard-120b", "openai/gpt-realtime-1.5", "openai/gpt-transcribe", "openai/sora-2-pro", "openai/text-embedding-ada-002", "openai/whisper-1", "openai/whisper-large-v3-turbo", "openrouter/auto", "openrouter/bodybuilder", "openrouter/free", "openrouter/pareto-code", "ornith-ai/ornith-1.5-35b-a3b", "perplexity/pplx-embed-v1-0.6b", "perplexity/pplx-embed-v1-4b", "perplexity/sonar-pro-search", "phi-4", "phi-4-mini", "phi-4-mini-reasoning", "phi-4-multimodal", "phi-4-reasoning", "phi-4-reasoning-plus", "poolside/laguna-s-2.1:free", "poolside/laguna-xs-2.1:free", "qwen-3-6-plus", "qwen-3.8-27b", "qwen-image-2.0", "qwen-image-2.0-pro", "qwen-mt-plus", "qwen-mt-turbo", "qwen-omni-turbo-realtime", "qwen-vl-max", "qwen-vl-ocr", "qwen-vl-plus", "qwen/qwen-2.5-7b-instruct", "qwen/qwen-2.5-coder-32b-instruct", "qwen/qwen-audio-3.0-tts-plus", "qwen/qwen-plus-2025-07-28", "qwen/qwen-turbo", "qwen/qwen-vl-max", "Qwen/Qwen2.5-7B-Instruct", "qwen/qwen2.5-coder-32b-instruct", "Qwen/Qwen2.5-VL-32B-Instruct", "Qwen/Qwen3-14B", "qwen/qwen3-235b-a22b-fp8", "Qwen/Qwen3-30B-A3B", "qwen/qwen3-30b-a3b-fp8", "qwen/qwen3-30b-a3b-thinking-2507", "qwen/qwen3-32b-fp8", "Qwen/Qwen3-Coder-480B-A35B-Instruct-FP8", "qwen/qwen3-embedding-4b", "qwen/qwen3-embedding-8b", "qwen/qwen3-max-thinking", "Qwen/Qwen3-VL-30B-A3B-Thinking", "qwen/qwen3-vl-32b-instruct", "Qwen/Qwen3-VL-32B-Instruct", "Qwen/Qwen3-VL-32B-Thinking", "Qwen/Qwen3-VL-8B-Instruct", "qwen/qwen3-vl-8b-thinking", "qwen/qwen3.5-flash-02-23", "qwen/qwen3.5-plus-02-15", "qwen/qwen3.5-plus-20260420", "qwen2-5-14b-instruct", "qwen2-5-32b-instruct", "qwen2-5-72b-instruct", "qwen2-5-7b-instruct", "qwen2-5-omni-7b", "qwen2-5-vl-72b-instruct", "qwen2-5-vl-7b-instruct", "qwen3-14b", "qwen3-30b-a3b-instruct-2507", "qwen3-5-35b-a3b", "qwen3-5-397b-a17b", "qwen3-5-9b", "qwen3-6-27b", "qwen3-6-35b-a3b", "qwen3-7-plus", "qwen3-8-max", "qwen3-8b", "qwen3-asr-flash", "qwen3-coder", "qwen3-embedding-8b", "qwen3-max-2026-01-23", "qwen3-max-preview", "qwen3-omni-flash", "qwen3-omni-flash-realtime", "qwen3-vl-235b-a22b-instruct", "qwen3-vl-30b-a3b", "qwen3.5-122b-a10b", "qwen3.5-27b", "qwen3.5-2b", "qwen3.6-35b", "qwen3.8-flash-next", "Qwen3.8-Max", "qwen3.8-max-preview", "qwen3guard-gen-0.6b", "qwen3guard-gen-8b", "qwq-plus", "reasoning_effort_values", "recraft/recraft-v3", "recraft/recraft-v4", "recraft/recraft-v4-pro", "recraft/recraft-v4.1", "recraft/recraft-v4.1-pro", "recraft/recraft-v4.1-utility", "recraft/recraft-v4.1-utility-pro", "rekaai/reka-edge", "rekaai/reka-flash-3", "relace/relace-apply-3", "relace/relace-search", "request", "sao10k/l3-lunaris-8b", "sao10k/l3.1-euryale-70b", "sao10k/l3.3-euryale-70b", "sarvam-105b", "seed-2-1-turbo", "stealth/union-alpha", "step-1-32k", "step-2-16k", "step-3-7-flash", "step-tts-2", "stepaudio-2.5-asr", "stepaudio-2.5-tts", "stepfun-ai/step-3.5-flash", "stepfun-ai/Step-3.7-Flash", "supported_formats", "tencent/hunyuan-a13b-instruct", "tencent/Hunyuan-A13B-Instruct", "tencent/hy-mt2-1.8b", "tencent/hy-mt2-30b-a3b", "tencent/hy-mt2-7b", "tencent/hy-mt2-plus", "tencent/Hy3", "thedrummer/cydonia-24b-v4.1", "thedrummer/skyfall-36b-v2", "thedrummer/unslopnemo-12b", "thinkingmachines/inkling-small:free", "training_data_cutoff", "umans-coder", "umans-deepseek-v4-flash-0731", "umans-deepseek-v4-pro-0813", "umans-flash", "umans-glm-5.3-flash", "umans-kimi-k3", "union-alpha", "volcengine/doubao-seed-2.0-code", "volcengine/doubao-seed-2.0-lite", "volcengine/doubao-seed-2.0-mini", "volcengine/doubao-seed-2.0-pro", "wan2.7-image", "wan2.7-image-pro", "whisper-large-v3-turbo", "x-ai/grok-4", "x-ai/grok-4-fast", "x-ai/grok-4.1-fast-non-reasoning", "x-ai/grok-code-fast-1", "x-ai/grok-imagine-image-2.0", "x-ai/grok-voice-tts-1.0", "xai.grok-4.3", "xai.grok-4.6", "xai/grok-4", "xai/grok-4.1-fast-non-reasoning", "xai/grok-4.1-fast-reasoning", "xiaomi/mimo-v2-flash", "xiaomimimo/mimo-v2-flash", "XiaomiMiMo/MiMo-V2-Flash", "z-ai/glm-4.7-flashx", "z-ai/glm-5.2:free", "zai-org/glm-4.5", "zai-org/GLM-4.5", "zai-org/glm-4.5v", "zai-org/glm-4.7", "zai-org/glm-4.7-flash", "zai-org/GLM-4.7-Flash", "zai-org/glm-5.1", "zai-org/GLM-5.2-Fast", "zai/glm-4.6v", "zai/glm-4.7-flash", "zai/glm-5v-turbo"];
|
|
@@ -1,2 +1,2 @@
|
|
|
1
1
|
// Generated by scripts/generate.mjs. Do not edit.
|
|
2
|
-
export const compactStrings = ["openai_chat_compatible", "/chat/completions", "Qwen vision-language model for visual reasoning, documents, and agent tasks", "openai_chat", "reasoning_content", "Compact GPT model for low-latency assistance and high-volume workloads", "Qwen instruction model for multilingual chat, reasoning, and tool use", "Image model for prompt-driven generation, editing, and visual design workflows", "Flagship Claude model for deep reasoning, coding, and long-horizon agents", "General purpose text generation", "Fast Gemini model balancing multimodal reasoning, tool use, and cost", "Open flagship GLM for long-horizon coding agents and million-token context work", "Open Llama instruction model for multilingual chat, reasoning, and coding", "Coding-optimized GPT model for repository edits, reviews, and agentic software work", "Balanced Claude model for coding, analysis, agent workflows, and cost control", "Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work", "GPT model for general reasoning, writing, coding, and tool-assisted tasks", "Flagship GLM model for hybrid reasoning, coding, and agentic engineering", "Multimodal reasoning model for visual analysis, planning, and tool use", "Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding", "DeepSeek chat model for instruction following, coding, and analysis", "llmgateway_providers", "Open GPT reasoning model for self-hosted agents and controllable deployments", "Open Gemma instruction model for efficient chat and self-hosted deployments", "Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking", "budget_tokens", "Efficient model for low-latency assistance, extraction, and routine automation", "Fast DeepSeek V4 lane for economical reasoning, coding, and long-context work", "Qwen coding model for software agents, repository edits, and code reasoning", "Embedding model for semantic search, retrieval, clustering, and ranking pipelines", "Multimodal, vision and text", "Open MoE flagship with million-token context for coding and long agent runs", "text-generation", "Flagship GLM model for long-horizon coding, agents, and complex project delivery", "MiniMax model for chat, coding, office work, and agentic tasks", "Open-weight GPT model for self-hosted reasoning and instruction-following workloads", "Multimodal Kimi workhorse for agent loops, coding tasks, and visual context", "Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents", "Speech generation model for controllable voice, narration, and audio delivery", "Native multimodal GLM model for efficient coding and long-horizon agent tasks", "Strong GLM coding model for agentic engineering, terminals, and repository generation", "DeepSeek V4 Pro snapshot with million-token context and support for thinking and non-thinking modes", "Fast Claude model for responsive assistance, classification, and lightweight agents", "2.4-trillion-parameter MoE flagship for coding, professional work, multimodal understanding, and long-horizon agentic workflows", "Legacy model retained for compatibility with older integrations", "Low-latency Gemini model for high-volume multimodal and agent workloads", "Kimi multimodal agent model for visual understanding, coding, and planning", "Grok model for agentic tool use, reasoning, coding, and live assistance", "Frontier GPT-5.6 model for complex professional work, coding, and agentic workflows", "Claude workhorse for coding agents, careful analysis, and production cost control", "openai_responses_compatible", "effort", "Mistral model for multilingual chat, reasoning, and tool-assisted workflows", "Reasoning model for deliberate analysis, multi-step problem solving, and tool use", "Everyday Claude agent model for coding, planning, browsing, and general work", "O-series reasoning model for hard analysis, math, coding, and planning", "Claude model for creative writing, analysis, and controlled agent workflows", "Flagship model for demanding analysis, coding, and production agent workflows", "Default frontier GPT for coding, computer use, research, and knowledge work", "Advanced Gemini model for complex reasoning, coding, and multimodal analysis", "GLM vision model for visual reasoning, documents, and multimodal agents", "openai_responses", "Video model for prompt-guided generation, editing, and motion workflows", "Fast Grok model for responsive chat, reasoning, and tool-assisted work", "Tencent Hy reasoning model for coding, instruction following, and agent tasks", "MiniMax multimodal model for long-context coding, perception, and agent planning", "Open-weight instruction model for adaptable chat and self-hosted production workloads", "google_generate_content", "image-text-to-text", "openrouter", "/models/{provider_model_id}:generateContent", "medium", "Largest Gemma 4 instruction model for open, self-hosted chat and reasoning", "High-end Claude for difficult coding, planning, and slower expert reasoning", "DeepSeek reasoning model for multi-step analysis, math, coding, and tools", "Stronger Opus tier for advanced software work and high-stakes reasoning", "Google's most intelligent Flash model, engineered for long-horizon software engineering, autonomous agents, and complex enterprise workflows", "Strongest Claude Opus model for coding, agents, and professional work", "Frontier GPT model for professional reasoning, coding, and multimodal work", "Qwen frontier model tuned for agent frameworks, coding assistants, and long tasks", "Balanced GPT-5.6 model for capable, cost-efficient everyday work", "Open MiniMax flagship for coding agents, office automation, and complex environments", "Agent-ready GPT for coding and computer-use workflows at a lower cost", "Efficient GLM model for fast reasoning, coding, and agent workflows", "GPT-6 Astra is OpenAI's most capable model for complex reasoning, coding, computer use, research, and document creation.", "Muse Spark 1.2 is a coding-focused update to Muse Spark 1.1 with improvements in code generation, complex debugging, codebase understanding, and end-to-end developer workflows.", "General GLM flagship for coding, analysis, and tool-heavy engineering workflows", "Cost-efficient GPT-5.6 model for fast, high-volume workloads", "xAI's frontier model for long-running agents, coding, knowledge work, and visual projects", "DeepSeek V4.1 Flash model for reasoning and agentic coding", "Muse Glimmer is a 30-billion-parameter open-weight multimodal model from Meta Superintelligence Labs, distilled from Muse Spark for always-on local agents, tool use, coding, and image understanding.", "Chat-tuned GPT model for conversational assistance, writing, and tool workflows", "Flagship Qwen model for complex reasoning, coding, and agentic workflows", "Efficient Mistral model for fast chat, extraction, and production assistants", "Compact Mistral model for edge, latency-sensitive, and cost-efficient workloads", "Qwen reasoning model for deliberate problem solving, math, and coding", "Open multimodal Qwen MoE for local agents that need vision, audio, and code", "Strong small GPT for coding subagents, quick tool use, and high-volume work", "Safety model for policy screening, moderation, and risk-aware routing workflows", "Stronger MiMo Pro tier for multimodal reasoning and coding-agent execution", "High-efficiency Gemini model for agentic workflows, coding, and multimodal reasoning", "Multimodal Qwen workhorse for long-context agents, visual inputs, and coding", "Dense 27B vision-language model for coding, agent tasks, and image and video understanding", "deepseek-thinking", "Earlier Kimi frontier model for long-context agents, coding, and multimodal work", "Coding model for repository understanding, refactors, and agentic engineering tasks", "Reasoning-first Gemini preview for agentic coding and complex problem solving", "Open MiMo model for multimodal coding agents and long-context automation", "mradermacher", "claude-opus", "deepseek-flash", "Open-weight sparse MoE (2.4T total, 95B active), the open-weight twin of Qwen3.8 Max for coding, research, complex reasoning, and agentic workflows", "Nemotron middle tier for collaborative agents and high-volume reasoning workloads", "https://${GOOGLE_VERTEX_ENDPOINT}/v1/projects/${GOOGLE_VERTEX_PROJECT}/locations/${GOOGLE_VERTEX_LOCATION}/endpoints/openapi", "Popular open Llama workhorse for multilingual chat, coding, and self-hosting", "Speech transcription model for accurate audio-to-text and captioning workflows", "General-purpose chat model for instruction following, writing, and analysis", "nano_gpt", "Mature GLM model for dependable coding, reasoning, and structured agent tasks", "2025-08-05", "Long-lived GPT workhorse for coding, instruction following, and production apps", "Cheapest GPT-5.4 lane for simple routing, extraction, and bulk automation", "Kimi model for long-context chat, coding, and agentic reasoning", "2026-07-09", "Fast Claude lane for lightweight agents, office tasks, and responsive chat", "Large open Qwen multimodal MoE for visual agents and long technical tasks", "amazon_bedrock", "anthropic_messages", "Muse Spark 1.3 is a multimodal reasoning model from Meta for long-running agentic, multi-agent, and coding workflows. It improves long-horizon agent collaboration, instruction following, and coding efficiency relative to Muse Spark 1.2.", "tool.web_search", "Reliable GPT generation for broad coding, writing, and tool-assisted product work", "2026-04-24", "xAI's default Grok for chat, coding, agentic tools, and lower hallucination risk", "claude-sonnet", "Prior MiniMax coding model for agent workflows, office edits, and automation", "2025-08-31", "Multimodal MoE reasoning model (975B total, 41B active) for text, image, and audio", "Original GPT-5 workhorse for reasoning, coding, writing, and tool workflows", "xAI's Grok model for chat, coding, agentic tools, and lower hallucination risk", "gemini-flash", "Claude model for demanding reasoning and long-horizon agentic work", "Small GPT-5 for responsive agents, coding help, and everyday automation", "Largest Nemotron 3 model for maximum open-weight reasoning and agent accuracy", "merge_gateway", "Affordable GPT-4.1 lane for fast coding help and structured extraction", "Flagship Mistral model for advanced reasoning, coding, and multilingual work", "Tool-capable chat model for instruction following and agentic application workflows", "azure_cognitive_services", "New Gemini flash lane bringing frontier-style multimodal reasoning to cheaper runs", "deprecated", "Automatic model router for matching prompts to suitable backends and budgets", "Earlier Qwen multimodal workhorse for million-token agent and document tasks", "Small omni GPT for cheap multimodal assistance and production-scale traffic", "Tiny GPT-4.1 option for classification, routing, and very high-volume tasks", "codecategories", "Sharper GPT-5 generation for coding, product work, and tool-assisted tasks", "Multimodal model for analyzing text, images, documents, and rich media", "StepFun flash model for efficient multimodal reasoning, coding, and tool use", "Instruction following, chat", "2026-04-22", "openai/gpt-oss-120b", "Tiny GPT-5 lane for routing, extraction, classification, and bulk jobs", "Fast Gemini workhorse for multimodal apps where latency and price matter", "gemini-flash-lite", "Omni-era GPT for multimodal chat, practical coding, and general assistants", "Highest-accuracy GPT-5.5 tier for slower, precision-heavy reasoning and coding", "2026-06-13", "High-speed MiniMax model for low-latency coding and agent workflows", "Hybrid-reasoning DeepSeek model with thinking and non-thinking modes, sparse attention, and tool-use", "toggle", "Google's proven reasoning model for coding, math, and multimodal analysis", "Nemotron model for efficient reasoning, coding, and specialized AI agents", "Quality-first multi-agent model for hard research, analysis, and competitions", "minimal", "Deliberate o-series reasoner for hard math, coding, and multi-step analysis", "Newer StepFun flash model for faster agents, coding, and multimodal prompts", "Budget GLM lane for fast coding help, routing, and everyday automation", "models", "@ai-sdk/openai-compatible", "2026-04-02", "@ai-sdk/anthropic", "2025-01", "2026-02-16", "2025-04", "Experimental multimodal DeepSeek V4 Flash model for image understanding, coding, and agentic work", "merge_by_id", "Efficient Qwen model for fast chat, extraction, and high-volume workloads", "More exact GPT-5.4 tier for demanding professional reasoning and agent tasks", "Flagship Qwen3 model for coding agents, complex reasoning, and tool use", "Late GLM-4 workhorse for coding agents, reasoning, and structured tasks", "Fast o-series model for compact reasoning, coding, and tool use", "Kimi reasoning model for long-horizon research, planning, and tool use", "Muse Spark is a natively multimodal reasoning model with support for tool-use, visual chain of thought, and multi-agent orchestration.", "Smaller o-series reasoner for economical coding, math, and planning tasks", "openai_images", "Fast Grok coding model tuned for agentic engineering and iterative edits", "Open multimodal Llama model for strong reasoning and fast responses", "Flagship DeepSeek model for coding, reasoning, and agentic work", "Lean Gemini 2.5 lane for cheap multimodal traffic and quick agents", "2025-08-07", "2026-08-14", "Open Llama multimodal model for image understanding and text reasoning", "llmgateway", "@ai-sdk/openai", "Balanced Mistral model for enterprise assistants, multilingual work, and tools", "https://${AZURE_COGNITIVE_SERVICES_RESOURCE_NAME}.services.ai.azure.com/anthropic/v1", "openai_embeddings", "2026-08-26", "2026-06-12", "Small Nemotron 3 MoE for efficient coding, math, and long-context agents", "Classic open reasoning model for transparent math, coding, and deliberate problem solving", "Mistral's largest general model for enterprise agents, coding, and multilingual reasoning", "2025-04-14", "2026-04-21", "uicomponent", "/images/generations", "Fast DeepSeek model for efficient chat, coding help, and agent loops", "2025-12-01", "2026-02-12", "Fast Mistral production model for chat, extraction, and cost-sensitive agents", "Fast GLM vision model for screenshots, documents, and multimodal agent tasks", "2026-03-18", "A next-generation productivity model with significantly enhanced Agent and complex task execution capabilities.", "2026-07-16", "deepseek-ai/DeepSeek-V4-Flash-0731", "Qwen/Qwen3-235B-A22B-Instruct-2507", "Nano Banana Pro for higher-fidelity image generation and design-heavy edits", "2026-04-16", "2026-05-28", "Qwen omni model for text, vision, audio, and multimodal agent tasks", "openai_completion", "2025-01-01", "web_search", "Hybrid-reasoning GLM release that made the 4.5 line broadly useful", "deepseek-ai/DeepSeek-V4-Pro", "Earlier MiniMax agent model for practical coding and productivity tasks", "http_request", "Mistral's coding-agent model for repository work, terminal tasks, and software fixes", "Reranking model for improving retrieval quality in search and recommendation systems", "Lower-latency Kimi Code variant for interactive edits and coding-agent loops", "Code-specialist GPT for repository edits, reviews, and long-running software agents", "2026-01-31", "Lighter GLM-4.5 variant for fast coding assistance and cheaper agents", "Low-latency M2.7 variant for interactive coding plans and agent loops", "/responses", "2026-03-17", "2025-06-17", "MiniMax multimodal coding model for long-context reasoning and agent tasks", "Open multimodal Llama model for long-context analysis and efficient agents", "openai/gpt-oss-20b", "2026-08-12", "2026-02-23", "Updated large open Qwen3 MoE instruct model for multilingual chat, coding, and tool use", "Hosted Qwen coder for software agents, repo edits, and long-context code", "Smaller Qwen coder for efficient local agents and repo-level fixes", "Largest open Gemma 3 instruction model for multilingual text generation and visual understanding"];
|
|
2
|
+
export const compactStrings = ["openai_chat_compatible", "/chat/completions", "Qwen vision-language model for visual reasoning, documents, and agent tasks", "openai_chat", "reasoning_content", "Compact GPT model for low-latency assistance and high-volume workloads", "Qwen instruction model for multilingual chat, reasoning, and tool use", "Image model for prompt-driven generation, editing, and visual design workflows", "Flagship Claude model for deep reasoning, coding, and long-horizon agents", "General purpose text generation", "Fast Gemini model balancing multimodal reasoning, tool use, and cost", "Open flagship GLM for long-horizon coding agents and million-token context work", "Open Llama instruction model for multilingual chat, reasoning, and coding", "Balanced Claude model for coding, analysis, agent workflows, and cost control", "Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work", "Coding-optimized GPT model for repository edits, reviews, and agentic software work", "GPT model for general reasoning, writing, coding, and tool-assisted tasks", "Flagship GLM model for hybrid reasoning, coding, and agentic engineering", "Multimodal reasoning model for visual analysis, planning, and tool use", "Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding", "DeepSeek chat model for instruction following, coding, and analysis", "llmgateway_providers", "Open GPT reasoning model for self-hosted agents and controllable deployments", "Open Gemma instruction model for efficient chat and self-hosted deployments", "budget_tokens", "Efficient model for low-latency assistance, extraction, and routine automation", "Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking", "Fast DeepSeek V4 lane for economical reasoning, coding, and long-context work", "Flagship GLM model for long-horizon coding, agents, and complex project delivery", "Qwen coding model for software agents, repository edits, and code reasoning", "Embedding model for semantic search, retrieval, clustering, and ranking pipelines", "Open MoE flagship with million-token context for coding and long agent runs", "Native multimodal GLM model for efficient coding and long-horizon agent tasks", "Multimodal, vision and text", "text-generation", "Speech generation model for controllable voice, narration, and audio delivery", "MiniMax model for chat, coding, office work, and agentic tasks", "Open-weight GPT model for self-hosted reasoning and instruction-following workloads", "Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents", "Multimodal Kimi workhorse for agent loops, coding tasks, and visual context", "DeepSeek V4 Pro snapshot with million-token context and support for thinking and non-thinking modes", "Strong GLM coding model for agentic engineering, terminals, and repository generation", "2.4-trillion-parameter MoE flagship for coding, professional work, multimodal understanding, and long-horizon agentic workflows", "Fast Claude model for responsive assistance, classification, and lightweight agents", "Grok model for agentic tool use, reasoning, coding, and live assistance", "Kimi multimodal agent model for visual understanding, coding, and planning", "Legacy model retained for compatibility with older integrations", "Low-latency Gemini model for high-volume multimodal and agent workloads", "Frontier GPT-5.6 model for complex professional work, coding, and agentic workflows", "Claude workhorse for coding agents, careful analysis, and production cost control", "effort", "openai_responses_compatible", "MiniMax multimodal model for long-context coding, perception, and agent planning", "Claude model for creative writing, analysis, and controlled agent workflows", "Mistral model for multilingual chat, reasoning, and tool-assisted workflows", "Everyday Claude agent model for coding, planning, browsing, and general work", "O-series reasoning model for hard analysis, math, coding, and planning", "Advanced Gemini model for complex reasoning, coding, and multimodal analysis", "Reasoning model for deliberate analysis, multi-step problem solving, and tool use", "Flagship model for demanding analysis, coding, and production agent workflows", "GLM vision model for visual reasoning, documents, and multimodal agents", "Default frontier GPT for coding, computer use, research, and knowledge work", "openai_responses", "Video model for prompt-guided generation, editing, and motion workflows", "Fast Grok model for responsive chat, reasoning, and tool-assisted work", "Tencent Hy reasoning model for coding, instruction following, and agent tasks", "openrouter", "Open-weight instruction model for adaptable chat and self-hosted production workloads", "google_generate_content", "Stronger Opus tier for advanced software work and high-stakes reasoning", "image-text-to-text", "/models/{provider_model_id}:generateContent", "medium", "High-end Claude for difficult coding, planning, and slower expert reasoning", "Largest Gemma 4 instruction model for open, self-hosted chat and reasoning", "DeepSeek reasoning model for multi-step analysis, math, coding, and tools", "DeepSeek V4.1 Flash model for reasoning and agentic coding", "Google's most intelligent Flash model, engineered for long-horizon software engineering, autonomous agents, and complex enterprise workflows", "Strongest Claude Opus model for coding, agents, and professional work", "Frontier GPT model for professional reasoning, coding, and multimodal work", "Open MiniMax flagship for coding agents, office automation, and complex environments", "Qwen frontier model tuned for agent frameworks, coding assistants, and long tasks", "Balanced GPT-5.6 model for capable, cost-efficient everyday work", "xAI's frontier model for long-running agents, coding, knowledge work, and visual projects", "Agent-ready GPT for coding and computer-use workflows at a lower cost", "Efficient GLM model for fast reasoning, coding, and agent workflows", "GPT-6 Astra is OpenAI's most capable model for complex reasoning, coding, computer use, research, and document creation.", "Muse Spark 1.2 is a coding-focused update to Muse Spark 1.1 with improvements in code generation, complex debugging, codebase understanding, and end-to-end developer workflows.", "Cost-efficient GPT-5.6 model for fast, high-volume workloads", "Muse Glimmer is a 30-billion-parameter open-weight multimodal model from Meta Superintelligence Labs, distilled from Muse Spark for always-on local agents, tool use, coding, and image understanding.", "Chat-tuned GPT model for conversational assistance, writing, and tool workflows", "General GLM flagship for coding, analysis, and tool-heavy engineering workflows", "Efficient Mistral model for fast chat, extraction, and production assistants", "Compact Mistral model for edge, latency-sensitive, and cost-efficient workloads", "Flagship Qwen model for complex reasoning, coding, and agentic workflows", "Qwen reasoning model for deliberate problem solving, math, and coding", "Open multimodal Qwen MoE for local agents that need vision, audio, and code", "Strong small GPT for coding subagents, quick tool use, and high-volume work", "Safety model for policy screening, moderation, and risk-aware routing workflows", "Stronger MiMo Pro tier for multimodal reasoning and coding-agent execution", "High-efficiency Gemini model for agentic workflows, coding, and multimodal reasoning", "Multimodal Qwen workhorse for long-context agents, visual inputs, and coding", "Dense 27B vision-language model for coding, agent tasks, and image and video understanding", "deepseek-thinking", "Open-weight sparse MoE (2.4T total, 95B active), the open-weight twin of Qwen3.8 Max for coding, research, complex reasoning, and agentic workflows", "Earlier Kimi frontier model for long-context agents, coding, and multimodal work", "Reasoning-first Gemini preview for agentic coding and complex problem solving", "Open MiMo model for multimodal coding agents and long-context automation", "deepseek-flash", "https://${GOOGLE_VERTEX_ENDPOINT}/v1/projects/${GOOGLE_VERTEX_PROJECT}/locations/${GOOGLE_VERTEX_LOCATION}/endpoints/openapi", "claude-opus", "mradermacher", "Nemotron middle tier for collaborative agents and high-volume reasoning workloads", "Popular open Llama workhorse for multilingual chat, coding, and self-hosting", "Speech transcription model for accurate audio-to-text and captioning workflows", "General-purpose chat model for instruction following, writing, and analysis", "Mature GLM model for dependable coding, reasoning, and structured agent tasks", "nano_gpt", "2025-08-05", "tool.web_search", "Long-lived GPT workhorse for coding, instruction following, and production apps", "Cheapest GPT-5.4 lane for simple routing, extraction, and bulk automation", "Kimi model for long-context chat, coding, and agentic reasoning", "2026-07-09", "amazon_bedrock", "Fast Claude lane for lightweight agents, office tasks, and responsive chat", "xAI's default Grok for chat, coding, agentic tools, and lower hallucination risk", "Coding model for repository understanding, refactors, and agentic engineering tasks", "Large open Qwen multimodal MoE for visual agents and long technical tasks", "anthropic_messages", "Muse Spark 1.3 is a multimodal reasoning model from Meta for long-running agentic, multi-agent, and coding workflows. It improves long-horizon agent collaboration, instruction following, and coding efficiency relative to Muse Spark 1.2.", "claude-sonnet", "Reliable GPT generation for broad coding, writing, and tool-assisted product work", "2026-04-24", "2025-08-31", "Prior MiniMax coding model for agent workflows, office edits, and automation", "Multimodal MoE reasoning model (975B total, 41B active) for text, image, and audio", "Original GPT-5 workhorse for reasoning, coding, writing, and tool workflows", "gemini-flash", "xAI's Grok model for chat, coding, agentic tools, and lower hallucination risk", "Claude model for demanding reasoning and long-horizon agentic work", "Small GPT-5 for responsive agents, coding help, and everyday automation", "Largest Nemotron 3 model for maximum open-weight reasoning and agent accuracy", "deprecated", "merge_gateway", "Affordable GPT-4.1 lane for fast coding help and structured extraction", "Flagship Mistral model for advanced reasoning, coding, and multilingual work", "Tool-capable chat model for instruction following and agentic application workflows", "azure_cognitive_services", "New Gemini flash lane bringing frontier-style multimodal reasoning to cheaper runs", "Multimodal model for analyzing text, images, documents, and rich media", "Automatic model router for matching prompts to suitable backends and budgets", "Small omni GPT for cheap multimodal assistance and production-scale traffic", "Tiny GPT-4.1 option for classification, routing, and very high-volume tasks", "codecategories", "Sharper GPT-5 generation for coding, product work, and tool-assisted tasks", "Earlier Qwen multimodal workhorse for million-token agent and document tasks", "StepFun flash model for efficient multimodal reasoning, coding, and tool use", "Instruction following, chat", "2026-04-22", "openai/gpt-oss-120b", "toggle", "Tiny GPT-5 lane for routing, extraction, classification, and bulk jobs", "Fast Gemini workhorse for multimodal apps where latency and price matter", "gemini-flash-lite", "2026-06-13", "Omni-era GPT for multimodal chat, practical coding, and general assistants", "Highest-accuracy GPT-5.5 tier for slower, precision-heavy reasoning and coding", "High-speed MiniMax model for low-latency coding and agent workflows", "Hybrid-reasoning DeepSeek model with thinking and non-thinking modes, sparse attention, and tool-use", "Google's proven reasoning model for coding, math, and multimodal analysis", "Nemotron model for efficient reasoning, coding, and specialized AI agents", "Quality-first multi-agent model for hard research, analysis, and competitions", "minimal", "Deliberate o-series reasoner for hard math, coding, and multi-step analysis", "Newer StepFun flash model for faster agents, coding, and multimodal prompts", "@ai-sdk/openai-compatible", "Budget GLM lane for fast coding help, routing, and everyday automation", "merge_by_id", "models", "@ai-sdk/anthropic", "2026-02-16", "2025-01", "2026-04-02", "2025-04", "Experimental multimodal DeepSeek V4 Flash model for image understanding, coding, and agentic work", "Efficient Qwen model for fast chat, extraction, and high-volume workloads", "More exact GPT-5.4 tier for demanding professional reasoning and agent tasks", "Flagship Qwen3 model for coding agents, complex reasoning, and tool use", "Late GLM-4 workhorse for coding agents, reasoning, and structured tasks", "Fast o-series model for compact reasoning, coding, and tool use", "Flagship DeepSeek model for coding, reasoning, and agentic work", "2026-08-14", "2026-08-26", "Kimi reasoning model for long-horizon research, planning, and tool use", "Muse Spark is a natively multimodal reasoning model with support for tool-use, visual chain of thought, and multi-agent orchestration.", "Smaller o-series reasoner for economical coding, math, and planning tasks", "openai_images", "Fast DeepSeek model for efficient chat, coding help, and agent loops", "Fast Grok coding model tuned for agentic engineering and iterative edits", "Open multimodal Llama model for strong reasoning and fast responses", "Lean Gemini 2.5 lane for cheap multimodal traffic and quick agents", "2025-08-07", "Open Llama multimodal model for image understanding and text reasoning", "llmgateway", "@ai-sdk/openai", "Balanced Mistral model for enterprise assistants, multilingual work, and tools", "https://${AZURE_COGNITIVE_SERVICES_RESOURCE_NAME}.services.ai.azure.com/anthropic/v1", "openai_embeddings", "Small Nemotron 3 MoE for efficient coding, math, and long-context agents", "Classic open reasoning model for transparent math, coding, and deliberate problem solving", "2025-04-14", "2026-04-21", "2026-06-12", "uicomponent", "2026-07-16", "/images/generations", "2025-12-01", "2026-02-12", "Fast Mistral production model for chat, extraction, and cost-sensitive agents", "2026-03-18", "web_search", "2026-04-16", "Fast GLM vision model for screenshots, documents, and multimodal agent tasks", "A next-generation productivity model with significantly enhanced Agent and complex task execution capabilities.", "deepseek-ai/DeepSeek-V4-Flash-0731", "Qwen/Qwen3-235B-A22B-Instruct-2507", "Mistral's largest general model for enterprise agents, coding, and multilingual reasoning", "Nano Banana Pro for higher-fidelity image generation and design-heavy edits", "2026-05-28", "Qwen omni model for text, vision, audio, and multimodal agent tasks", "2025-01-01", "openai_completion", "Hybrid-reasoning GLM release that made the 4.5 line broadly useful", "deepseek-ai/DeepSeek-V4-Pro", "Earlier MiniMax agent model for practical coding and productivity tasks", "http_request", "Mistral's coding-agent model for repository work, terminal tasks, and software fixes", "Reranking model for improving retrieval quality in search and recommendation systems", "Lower-latency Kimi Code variant for interactive edits and coding-agent loops", "2026-01-31", "2026-08-12", "Lighter GLM-4.5 variant for fast coding assistance and cheaper agents", "Low-latency M2.7 variant for interactive coding plans and agent loops", "/responses", "2026-03-17", "2025-06-17", "Open multimodal Llama model for long-context analysis and efficient agents", "Updated large open Qwen3 MoE instruct model for multilingual chat, coding, and tool use", "Hosted Qwen coder for software agents, repo edits, and long-context code", "2026-02-23", "Smaller Qwen coder for efficient local agents and repo-level fixes", "Largest open Gemma 3 instruction model for multilingual text generation and visual understanding", "Efficient open MiniMax model built for coding agents and tool-heavy workflows", "https://${AZURE_RESOURCE_NAME}.services.ai.azure.com/anthropic/v1", "unsloth"];
|
|
@@ -1,3 +1,3 @@
|
|
|
1
1
|
import type { Manifest } from "../types.js";
|
|
2
|
-
export type KnownProviderId = "302ai" | "abacus" | "abliteration_ai" | "above" | "agentrouter" | "agnes" | "ai_router" | "aiand" | "aihubmix" | "aixy" | "aki_io" | "alibaba" | "alibaba_cn" | "alibaba_coding_plan" | "alibaba_coding_plan_cn" | "alibaba_token_plan" | "alibaba_token_plan_cn" | "amazon_bedrock" | "ambient" | "amd" | "anthropic" | "anyapi" | "arcee" | "atomic_chat" | "auriko" | "azure" | "azure_cognitive_services" | "bailing" | "baseten" | "berget" | "blueclaw" | "bothub" | "cerebras" | "chutes" | "clarifai" | "claudinio" | "cline_pass" | "cloudferro_sherlock" | "cloudflare_ai_gateway" | "cloudflare_workers_ai" | "cohere" | "coralbricks" | "cortecs" | "crof" | "crossmodel" | "crusoe" | "daoxe" | "databricks" | "deepinfra" | "deepseek" | "digitalocean" | "dinference" | "drun" | "ebcloud" | "echo" | "edenai" | "elevenlabs" | "empiriolabs" | "evroc" | "fastrouter" | "fireworks_ai" | "freemodel" | "friendli" | "frogbot" | "github_copilot" | "github_models" | "gitlab" | "gmicloud" | "google" | "google_vertex" | "google_vertex_anthropic" | "greenpt" | "groq" | "helicone" | "hetzner" | "hpc_ai" | "huggingface" | "hyper" | "iflowcn" | "impossibl" | "inception" | "inceptron" | "infer" | "inference" | "inferx" | "infomaniak" | "io_net" | "iteracompute" | "jalapeno" | "jiekou" | "kenari" | "kilo" | "kimi_for_coding" | "klokintegration" | "kosmik" | "kuae_cloud_coding_plan" | "lilac" | "llama" | "llmgateway" | "llmgateway_providers" | "llmtech" | "llmtr" | "lmstudio" | "longcat" | "lucidquery" | "lynkr" | "meganova" | "melious" | "merge_gateway" | "meta" | "minimax" | "minimax_cn" | "minimax_cn_coding_plan" | "minimax_coding_plan" | "mistral" | "mixlayer" | "moark" | "modal" | "model_oracle_ai" | "modelis" | "modelscope" | "moonshotai" | "moonshotai_cn" | "morph" | "nan" | "nano_gpt" | "nearai" | "nebius" | "neon" | "neosmith" | "neuralwatt" | "nova" | "novita_ai" | "nvidia" | "ofox" | "ollama_cloud" | "openai" | "opencode" | "opencode_go" | "openreason" | "openrouter" | "opper" | "orcarouter" | "ovhcloud" | "pendra" | "perplexity" | "perplexity_agent" | "pioneer" | "poe" | "poolside" | "privatemode_ai" | "qihang_ai" | "qiniu_ai" | "qvac" | "regolo_ai" | "requesty" | "routing_run" | "runinfra" | "sakana" | "salad_cloud" | "sap_ai_core" | "sarvam" | "scaleway" | "scnet_token_plan" | "scx_ai" | "sensenova" | "siliconflow" | "siliconflow_cn" | "snowflake_cortex" | "stackit" | "standardcompute" | "stepfun" | "stepfun_ai" | "stepfun_ai_step_plan" | "stepfun_step_plan" | "subconscious" | "submodel" | "synthetic" | "tencent_coding_plan" | "tencent_token_plan" | "tencent_tokenhub" | "tensorx" | "the_grid_ai" | "thinkingmachines" | "tinfoil" | "togetherai" | "tokengo" | "tokenrouter" | "trustedrouter" | "umans_ai" | "umans_ai_coding_plan" | "unorouter" | "upstage" | "v0" | "vancine" | "venice" | "vercel" | "vispark" | "vivgrid" | "volcengine" | "volcengine_coding_plan" | "vultr" | "wafer_ai" | "wallaby" | "wandb" | "watsonx" | "xai" | "xiaomi" | "xiaomi_token_plan_ams" | "xiaomi_token_plan_cn" | "xiaomi_token_plan_sgp" | "xpersona" | "zai" | "zai_coder" | "zai_coding_plan" | "zeldoc" | "zenifra" | "zenmux" | "zhipuai" | "zhipuai_coding_plan";
|
|
2
|
+
export type KnownProviderId = "302ai" | "abacus" | "abliteration_ai" | "above" | "agentrouter" | "agnes" | "ai21" | "ai_router" | "aiand" | "aihubmix" | "aixy" | "aki_io" | "alibaba" | "alibaba_cn" | "alibaba_coding_plan" | "alibaba_coding_plan_cn" | "alibaba_token_plan" | "alibaba_token_plan_cn" | "amazon_bedrock" | "ambient" | "amd" | "anthropic" | "anyapi" | "arcee" | "atomic_chat" | "auriko" | "azure" | "azure_cognitive_services" | "bailing" | "baseten" | "berget" | "blueclaw" | "bothub" | "cerebras" | "chutes" | "clarifai" | "claudinio" | "cline_pass" | "cloudferro_sherlock" | "cloudflare_ai_gateway" | "cloudflare_workers_ai" | "cohere" | "coralbricks" | "cortecs" | "crof" | "crossmodel" | "crusoe" | "daoxe" | "databricks" | "deepinfra" | "deepseek" | "digitalocean" | "dinference" | "drun" | "ebcloud" | "echo" | "edenai" | "elevenlabs" | "empiriolabs" | "evroc" | "fastrouter" | "fireworks_ai" | "freemodel" | "friendli" | "frogbot" | "github_copilot" | "github_models" | "gitlab" | "gmicloud" | "google" | "google_vertex" | "google_vertex_anthropic" | "greenpt" | "groq" | "helicone" | "hetzner" | "hpc_ai" | "huggingface" | "hyper" | "iflowcn" | "impossibl" | "inception" | "inceptron" | "inco" | "infer" | "inference" | "inferx" | "infomaniak" | "io_net" | "iteracompute" | "jalapeno" | "jiekou" | "kenari" | "kilo" | "kimi_for_coding" | "klokintegration" | "kosmik" | "kuae_cloud_coding_plan" | "lilac" | "llama" | "llmgateway" | "llmgateway_providers" | "llmtech" | "llmtr" | "lmstudio" | "longcat" | "lucidquery" | "lynkr" | "meganova" | "melious" | "merge_gateway" | "meta" | "minimax" | "minimax_cn" | "minimax_cn_coding_plan" | "minimax_coding_plan" | "mistral" | "mixlayer" | "moark" | "modal" | "model_oracle_ai" | "modelis" | "modelscope" | "moonshotai" | "moonshotai_cn" | "morph" | "nan" | "nano_gpt" | "nearai" | "nebius" | "neon" | "neosmith" | "neuralwatt" | "nova" | "novita_ai" | "nvidia" | "oci" | "ofox" | "ollama_cloud" | "openai" | "opencode" | "opencode_go" | "openreason" | "openrouter" | "opper" | "orcarouter" | "ovhcloud" | "pendra" | "perplexity" | "perplexity_agent" | "pioneer" | "poe" | "poolside" | "privatemode_ai" | "qihang_ai" | "qiniu_ai" | "qvac" | "regolo_ai" | "requesty" | "routing_run" | "runinfra" | "sakana" | "salad_cloud" | "sap_ai_core" | "sarvam" | "scaleway" | "scnet_token_plan" | "scx_ai" | "sensenova" | "siliconflow" | "siliconflow_cn" | "snowflake_cortex" | "stackit" | "standardcompute" | "stepfun" | "stepfun_ai" | "stepfun_ai_step_plan" | "stepfun_step_plan" | "subconscious" | "submodel" | "synthetic" | "tencent_coding_plan" | "tencent_token_plan" | "tencent_tokenhub" | "tensorx" | "the_grid_ai" | "thinkingmachines" | "tinfoil" | "togetherai" | "tokengo" | "tokenrouter" | "trustedrouter" | "typesafe" | "umans_ai" | "umans_ai_coding_plan" | "unorouter" | "upstage" | "v0" | "vancine" | "venice" | "vercel" | "vispark" | "vivgrid" | "volcengine" | "volcengine_coding_plan" | "vultr" | "wafer_ai" | "wallaby" | "wandb" | "watsonx" | "xai" | "xiaomi" | "xiaomi_token_plan_ams" | "xiaomi_token_plan_cn" | "xiaomi_token_plan_sgp" | "xpersona" | "zai" | "zai_coder" | "zai_coding_plan" | "zeldoc" | "zenifra" | "zenmux" | "zhipuai" | "zhipuai_coding_plan";
|
|
3
3
|
export declare const manifest: Manifest;
|