@homeflare/distilled-litellm 0.2.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/LICENSE +201 -0
- package/README.md +77 -0
- package/dist/credentials.d.ts +27 -0
- package/dist/credentials.d.ts.map +1 -0
- package/dist/credentials.js +67 -0
- package/dist/credentials.js.map +1 -0
- package/dist/errors.d.ts +92 -0
- package/dist/errors.d.ts.map +1 -0
- package/dist/errors.js +72 -0
- package/dist/errors.js.map +1 -0
- package/dist/index.d.ts +30 -0
- package/dist/index.d.ts.map +1 -0
- package/dist/index.js +30 -0
- package/dist/index.js.map +1 -0
- package/dist/protocol.d.ts +18 -0
- package/dist/protocol.d.ts.map +1 -0
- package/dist/protocol.js +135 -0
- package/dist/protocol.js.map +1 -0
- package/dist/retry.d.ts +50 -0
- package/dist/retry.d.ts.map +1 -0
- package/dist/retry.js +43 -0
- package/dist/retry.js.map +1 -0
- package/dist/services/a2a.d.ts +70 -0
- package/dist/services/a2a.d.ts.map +1 -0
- package/dist/services/a2a.js +128 -0
- package/dist/services/a2a.js.map +1 -0
- package/dist/services/a2a_registration.d.ts +43 -0
- package/dist/services/a2a_registration.d.ts.map +1 -0
- package/dist/services/a2a_registration.js +40 -0
- package/dist/services/a2a_registration.js.map +1 -0
- package/dist/services/access_groups.d.ts +192 -0
- package/dist/services/access_groups.d.ts.map +1 -0
- package/dist/services/access_groups.js +283 -0
- package/dist/services/access_groups.js.map +1 -0
- package/dist/services/adaptive_router.d.ts +15 -0
- package/dist/services/adaptive_router.d.ts.map +1 -0
- package/dist/services/adaptive_router.js +25 -0
- package/dist/services/adaptive_router.js.map +1 -0
- package/dist/services/agents.d.ts +537 -0
- package/dist/services/agents.d.ts.map +1 -0
- package/dist/services/agents.js +481 -0
- package/dist/services/agents.js.map +1 -0
- package/dist/services/alerting.d.ts +15 -0
- package/dist/services/alerting.d.ts.map +1 -0
- package/dist/services/alerting.js +25 -0
- package/dist/services/alerting.js.map +1 -0
- package/dist/services/anthropic_pass_through.d.ts +70 -0
- package/dist/services/anthropic_pass_through.d.ts.map +1 -0
- package/dist/services/anthropic_pass_through.js +114 -0
- package/dist/services/anthropic_pass_through.js.map +1 -0
- package/dist/services/anthropic_passthrough.d.ts +35 -0
- package/dist/services/anthropic_passthrough.d.ts.map +1 -0
- package/dist/services/anthropic_passthrough.js +59 -0
- package/dist/services/anthropic_passthrough.js.map +1 -0
- package/dist/services/anthropic_skills.d.ts +74 -0
- package/dist/services/anthropic_skills.d.ts.map +1 -0
- package/dist/services/anthropic_skills.js +94 -0
- package/dist/services/anthropic_skills.js.map +1 -0
- package/dist/services/assembly_ai_eu_pass_through.d.ts +70 -0
- package/dist/services/assembly_ai_eu_pass_through.d.ts.map +1 -0
- package/dist/services/assembly_ai_eu_pass_through.js +114 -0
- package/dist/services/assembly_ai_eu_pass_through.js.map +1 -0
- package/dist/services/assembly_ai_pass_through.d.ts +70 -0
- package/dist/services/assembly_ai_pass_through.d.ts.map +1 -0
- package/dist/services/assembly_ai_pass_through.js +114 -0
- package/dist/services/assembly_ai_pass_through.js.map +1 -0
- package/dist/services/assistants.d.ts +185 -0
- package/dist/services/assistants.d.ts.map +1 -0
- package/dist/services/assistants.js +331 -0
- package/dist/services/assistants.js.map +1 -0
- package/dist/services/audio.d.ts +57 -0
- package/dist/services/audio.d.ts.map +1 -0
- package/dist/services/audio.js +96 -0
- package/dist/services/audio.js.map +1 -0
- package/dist/services/audit_logging.d.ts +98 -0
- package/dist/services/audit_logging.d.ts.map +1 -0
- package/dist/services/audit_logging.js +83 -0
- package/dist/services/audit_logging.js.map +1 -0
- package/dist/services/auto_router.d.ts +286 -0
- package/dist/services/auto_router.d.ts.map +1 -0
- package/dist/services/auto_router.js +248 -0
- package/dist/services/auto_router.js.map +1 -0
- package/dist/services/aws_comprehend_medical_pass_through.d.ts +36 -0
- package/dist/services/aws_comprehend_medical_pass_through.d.ts.map +1 -0
- package/dist/services/aws_comprehend_medical_pass_through.js +56 -0
- package/dist/services/aws_comprehend_medical_pass_through.js.map +1 -0
- package/dist/services/azure_ai_pass_through.d.ts +70 -0
- package/dist/services/azure_ai_pass_through.d.ts.map +1 -0
- package/dist/services/azure_ai_pass_through.js +112 -0
- package/dist/services/azure_ai_pass_through.js.map +1 -0
- package/dist/services/azure_pass_through.d.ts +70 -0
- package/dist/services/azure_pass_through.d.ts.map +1 -0
- package/dist/services/azure_pass_through.js +107 -0
- package/dist/services/azure_pass_through.js.map +1 -0
- package/dist/services/batch.d.ts +165 -0
- package/dist/services/batch.d.ts.map +1 -0
- package/dist/services/batch.js +265 -0
- package/dist/services/batch.js.map +1 -0
- package/dist/services/bedrock_pass_through.d.ts +70 -0
- package/dist/services/bedrock_pass_through.d.ts.map +1 -0
- package/dist/services/bedrock_pass_through.js +114 -0
- package/dist/services/bedrock_pass_through.js.map +1 -0
- package/dist/services/beta_agents.d.ts +199 -0
- package/dist/services/beta_agents.d.ts.map +1 -0
- package/dist/services/beta_agents.js +151 -0
- package/dist/services/beta_agents.js.map +1 -0
- package/dist/services/beta_mcp.d.ts +37 -0
- package/dist/services/beta_mcp.d.ts.map +1 -0
- package/dist/services/beta_mcp.js +40 -0
- package/dist/services/beta_mcp.js.map +1 -0
- package/dist/services/budget_management.d.ts +185 -0
- package/dist/services/budget_management.d.ts.map +1 -0
- package/dist/services/budget_management.js +196 -0
- package/dist/services/budget_management.js.map +1 -0
- package/dist/services/budget_spend_tracking.d.ts +967 -0
- package/dist/services/budget_spend_tracking.d.ts.map +1 -0
- package/dist/services/budget_spend_tracking.js +1103 -0
- package/dist/services/budget_spend_tracking.js.map +1 -0
- package/dist/services/cache_settings.d.ts +96 -0
- package/dist/services/cache_settings.d.ts.map +1 -0
- package/dist/services/cache_settings.js +95 -0
- package/dist/services/cache_settings.js.map +1 -0
- package/dist/services/caching.d.ts +54 -0
- package/dist/services/caching.d.ts.map +1 -0
- package/dist/services/caching.js +78 -0
- package/dist/services/caching.js.map +1 -0
- package/dist/services/chat_completions.d.ts +656 -0
- package/dist/services/chat_completions.d.ts.map +1 -0
- package/dist/services/chat_completions.js +590 -0
- package/dist/services/chat_completions.js.map +1 -0
- package/dist/services/claude_code_marketplace.d.ts +212 -0
- package/dist/services/claude_code_marketplace.d.ts.map +1 -0
- package/dist/services/claude_code_marketplace.js +248 -0
- package/dist/services/claude_code_marketplace.js.map +1 -0
- package/dist/services/cloudzero.d.ts +117 -0
- package/dist/services/cloudzero.d.ts.map +1 -0
- package/dist/services/cloudzero.js +130 -0
- package/dist/services/cloudzero.js.map +1 -0
- package/dist/services/cohere_pass_through.d.ts +70 -0
- package/dist/services/cohere_pass_through.d.ts.map +1 -0
- package/dist/services/cohere_pass_through.js +112 -0
- package/dist/services/cohere_pass_through.js.map +1 -0
- package/dist/services/completions.d.ts +59 -0
- package/dist/services/completions.d.ts.map +1 -0
- package/dist/services/completions.js +98 -0
- package/dist/services/completions.js.map +1 -0
- package/dist/services/compliance.d.ts +66 -0
- package/dist/services/compliance.d.ts.map +1 -0
- package/dist/services/compliance.js +74 -0
- package/dist/services/compliance.js.map +1 -0
- package/dist/services/config_overrides.d.ts +157 -0
- package/dist/services/config_overrides.d.ts.map +1 -0
- package/dist/services/config_overrides.js +207 -0
- package/dist/services/config_overrides.js.map +1 -0
- package/dist/services/config_yaml.d.ts +647 -0
- package/dist/services/config_yaml.d.ts.map +1 -0
- package/dist/services/config_yaml.js +504 -0
- package/dist/services/config_yaml.js.map +1 -0
- package/dist/services/containers.d.ts +215 -0
- package/dist/services/containers.d.ts.map +1 -0
- package/dist/services/containers.js +416 -0
- package/dist/services/containers.js.map +1 -0
- package/dist/services/coordination_redis_settings.d.ts +93 -0
- package/dist/services/coordination_redis_settings.d.ts.map +1 -0
- package/dist/services/coordination_redis_settings.js +104 -0
- package/dist/services/coordination_redis_settings.js.map +1 -0
- package/dist/services/cost_tracking.d.ts +140 -0
- package/dist/services/cost_tracking.d.ts.map +1 -0
- package/dist/services/cost_tracking.js +181 -0
- package/dist/services/cost_tracking.js.map +1 -0
- package/dist/services/credential_management.d.ts +133 -0
- package/dist/services/credential_management.d.ts.map +1 -0
- package/dist/services/credential_management.js +198 -0
- package/dist/services/credential_management.js.map +1 -0
- package/dist/services/cursor_pass_through.d.ts +70 -0
- package/dist/services/cursor_pass_through.d.ts.map +1 -0
- package/dist/services/cursor_pass_through.js +112 -0
- package/dist/services/cursor_pass_through.js.map +1 -0
- package/dist/services/customer_management.d.ts +545 -0
- package/dist/services/customer_management.d.ts.map +1 -0
- package/dist/services/customer_management.js +567 -0
- package/dist/services/customer_management.js.map +1 -0
- package/dist/services/email_management.d.ts +57 -0
- package/dist/services/email_management.d.ts.map +1 -0
- package/dist/services/email_management.js +79 -0
- package/dist/services/email_management.js.map +1 -0
- package/dist/services/embeddings.d.ts +155 -0
- package/dist/services/embeddings.d.ts.map +1 -0
- package/dist/services/embeddings.js +168 -0
- package/dist/services/embeddings.js.map +1 -0
- package/dist/services/evals.d.ts +245 -0
- package/dist/services/evals.d.ts.map +1 -0
- package/dist/services/evals.js +292 -0
- package/dist/services/evals.js.map +1 -0
- package/dist/services/experimental.d.ts +166 -0
- package/dist/services/experimental.d.ts.map +1 -0
- package/dist/services/experimental.js +265 -0
- package/dist/services/experimental.js.map +1 -0
- package/dist/services/fallback_management.d.ts +91 -0
- package/dist/services/fallback_management.d.ts.map +1 -0
- package/dist/services/fallback_management.js +86 -0
- package/dist/services/fallback_management.js.map +1 -0
- package/dist/services/files.d.ts +219 -0
- package/dist/services/files.d.ts.map +1 -0
- package/dist/services/files.js +358 -0
- package/dist/services/files.js.map +1 -0
- package/dist/services/fine_tuning.d.ts +155 -0
- package/dist/services/fine_tuning.d.ts.map +1 -0
- package/dist/services/fine_tuning.js +232 -0
- package/dist/services/fine_tuning.js.map +1 -0
- package/dist/services/gemini_agents.d.ts +68 -0
- package/dist/services/gemini_agents.d.ts.map +1 -0
- package/dist/services/gemini_agents.js +110 -0
- package/dist/services/gemini_agents.js.map +1 -0
- package/dist/services/google_ai_studio_pass_through.d.ts +70 -0
- package/dist/services/google_ai_studio_pass_through.d.ts.map +1 -0
- package/dist/services/google_ai_studio_pass_through.js +112 -0
- package/dist/services/google_ai_studio_pass_through.js.map +1 -0
- package/dist/services/google_genai_endpoints.d.ts +174 -0
- package/dist/services/google_genai_endpoints.d.ts.map +1 -0
- package/dist/services/google_genai_endpoints.js +340 -0
- package/dist/services/google_genai_endpoints.js.map +1 -0
- package/dist/services/guardrails.d.ts +1221 -0
- package/dist/services/guardrails.d.ts.map +1 -0
- package/dist/services/guardrails.js +1087 -0
- package/dist/services/guardrails.js.map +1 -0
- package/dist/services/health.d.ts +212 -0
- package/dist/services/health.d.ts.map +1 -0
- package/dist/services/health.js +309 -0
- package/dist/services/health.js.map +1 -0
- package/dist/services/images.d.ts +117 -0
- package/dist/services/images.d.ts.map +1 -0
- package/dist/services/images.js +184 -0
- package/dist/services/images.js.map +1 -0
- package/dist/services/index.d.ts +106 -0
- package/dist/services/index.d.ts.map +1 -0
- package/dist/services/index.js +107 -0
- package/dist/services/index.js.map +1 -0
- package/dist/services/internal_user_management.d.ts +1062 -0
- package/dist/services/internal_user_management.d.ts.map +1 -0
- package/dist/services/internal_user_management.js +850 -0
- package/dist/services/internal_user_management.js.map +1 -0
- package/dist/services/invite_links.d.ts +56 -0
- package/dist/services/invite_links.d.ts.map +1 -0
- package/dist/services/invite_links.js +82 -0
- package/dist/services/invite_links.js.map +1 -0
- package/dist/services/jwt_mappings.d.ts +79 -0
- package/dist/services/jwt_mappings.d.ts.map +1 -0
- package/dist/services/jwt_mappings.js +116 -0
- package/dist/services/jwt_mappings.js.map +1 -0
- package/dist/services/key_management.d.ts +1640 -0
- package/dist/services/key_management.d.ts.map +1 -0
- package/dist/services/key_management.js +1344 -0
- package/dist/services/key_management.js.map +1 -0
- package/dist/services/langfuse_passthrough.d.ts +70 -0
- package/dist/services/langfuse_passthrough.d.ts.map +1 -0
- package/dist/services/langfuse_passthrough.js +114 -0
- package/dist/services/langfuse_passthrough.js.map +1 -0
- package/dist/services/llm_utils.d.ts +101 -0
- package/dist/services/llm_utils.d.ts.map +1 -0
- package/dist/services/llm_utils.js +110 -0
- package/dist/services/llm_utils.js.map +1 -0
- package/dist/services/logging_callbacks.d.ts +33 -0
- package/dist/services/logging_callbacks.d.ts.map +1 -0
- package/dist/services/logging_callbacks.js +46 -0
- package/dist/services/logging_callbacks.js.map +1 -0
- package/dist/services/mcp_byok_oauth.d.ts +112 -0
- package/dist/services/mcp_byok_oauth.d.ts.map +1 -0
- package/dist/services/mcp_byok_oauth.js +226 -0
- package/dist/services/mcp_byok_oauth.js.map +1 -0
- package/dist/services/mcp_discoverable.d.ts +197 -0
- package/dist/services/mcp_discoverable.d.ts.map +1 -0
- package/dist/services/mcp_discoverable.js +304 -0
- package/dist/services/mcp_discoverable.js.map +1 -0
- package/dist/services/mcp_management.d.ts +956 -0
- package/dist/services/mcp_management.d.ts.map +1 -0
- package/dist/services/mcp_management.js +1139 -0
- package/dist/services/mcp_management.js.map +1 -0
- package/dist/services/mcp_rest.d.ts +281 -0
- package/dist/services/mcp_rest.d.ts.map +1 -0
- package/dist/services/mcp_rest.js +269 -0
- package/dist/services/mcp_rest.js.map +1 -0
- package/dist/services/memory_management.d.ts +93 -0
- package/dist/services/memory_management.d.ts.map +1 -0
- package/dist/services/memory_management.js +117 -0
- package/dist/services/memory_management.js.map +1 -0
- package/dist/services/milvus_pass_through.d.ts +70 -0
- package/dist/services/milvus_pass_through.d.ts.map +1 -0
- package/dist/services/milvus_pass_through.js +112 -0
- package/dist/services/milvus_pass_through.js.map +1 -0
- package/dist/services/misc.d.ts +676 -0
- package/dist/services/misc.d.ts.map +1 -0
- package/dist/services/misc.js +954 -0
- package/dist/services/misc.js.map +1 -0
- package/dist/services/mistral_pass_through.d.ts +70 -0
- package/dist/services/mistral_pass_through.d.ts.map +1 -0
- package/dist/services/mistral_pass_through.js +114 -0
- package/dist/services/mistral_pass_through.js.map +1 -0
- package/dist/services/model_management.d.ts +1649 -0
- package/dist/services/model_management.d.ts.map +1 -0
- package/dist/services/model_management.js +1805 -0
- package/dist/services/model_management.js.map +1 -0
- package/dist/services/moderations.d.ts +25 -0
- package/dist/services/moderations.d.ts.map +1 -0
- package/dist/services/moderations.js +39 -0
- package/dist/services/moderations.js.map +1 -0
- package/dist/services/ocr.d.ts +25 -0
- package/dist/services/ocr.d.ts.map +1 -0
- package/dist/services/ocr.js +39 -0
- package/dist/services/ocr.js.map +1 -0
- package/dist/services/open_ai_pass_through.d.ts +125 -0
- package/dist/services/open_ai_pass_through.d.ts.map +1 -0
- package/dist/services/open_ai_pass_through.js +232 -0
- package/dist/services/open_ai_pass_through.js.map +1 -0
- package/dist/services/organization_management.d.ts +664 -0
- package/dist/services/organization_management.d.ts.map +1 -0
- package/dist/services/organization_management.js +644 -0
- package/dist/services/organization_management.js.map +1 -0
- package/dist/services/plugins.d.ts +106 -0
- package/dist/services/plugins.d.ts.map +1 -0
- package/dist/services/plugins.js +180 -0
- package/dist/services/plugins.js.map +1 -0
- package/dist/services/policies.d.ts +618 -0
- package/dist/services/policies.d.ts.map +1 -0
- package/dist/services/policies.js +616 -0
- package/dist/services/policies.js.map +1 -0
- package/dist/services/policy_engine.d.ts +483 -0
- package/dist/services/policy_engine.d.ts.map +1 -0
- package/dist/services/policy_engine.js +501 -0
- package/dist/services/policy_engine.js.map +1 -0
- package/dist/services/project_management.d.ts +356 -0
- package/dist/services/project_management.d.ts.map +1 -0
- package/dist/services/project_management.js +318 -0
- package/dist/services/project_management.js.map +1 -0
- package/dist/services/prompts.d.ts +193 -0
- package/dist/services/prompts.d.ts.map +1 -0
- package/dist/services/prompts.js +260 -0
- package/dist/services/prompts.js.map +1 -0
- package/dist/services/public.d.ts +300 -0
- package/dist/services/public.d.ts.map +1 -0
- package/dist/services/public.js +362 -0
- package/dist/services/public.js.map +1 -0
- package/dist/services/rag.d.ts +45 -0
- package/dist/services/rag.d.ts.map +1 -0
- package/dist/services/rag.js +71 -0
- package/dist/services/rag.js.map +1 -0
- package/dist/services/realtime.d.ts +91 -0
- package/dist/services/realtime.d.ts.map +1 -0
- package/dist/services/realtime.js +164 -0
- package/dist/services/realtime.js.map +1 -0
- package/dist/services/rerank.d.ts +35 -0
- package/dist/services/rerank.d.ts.map +1 -0
- package/dist/services/rerank.js +55 -0
- package/dist/services/rerank.js.map +1 -0
- package/dist/services/responses.d.ts +237 -0
- package/dist/services/responses.d.ts.map +1 -0
- package/dist/services/responses.js +445 -0
- package/dist/services/responses.js.map +1 -0
- package/dist/services/router_settings.d.ts +67 -0
- package/dist/services/router_settings.d.ts.map +1 -0
- package/dist/services/router_settings.js +63 -0
- package/dist/services/router_settings.js.map +1 -0
- package/dist/services/rust_control_plane.d.ts +52 -0
- package/dist/services/rust_control_plane.d.ts.map +1 -0
- package/dist/services/rust_control_plane.js +54 -0
- package/dist/services/rust_control_plane.js.map +1 -0
- package/dist/services/scim.d.ts +412 -0
- package/dist/services/scim.d.ts.map +1 -0
- package/dist/services/scim.js +519 -0
- package/dist/services/scim.js.map +1 -0
- package/dist/services/search.d.ts +79 -0
- package/dist/services/search.d.ts.map +1 -0
- package/dist/services/search.js +122 -0
- package/dist/services/search.js.map +1 -0
- package/dist/services/search_tools.d.ts +140 -0
- package/dist/services/search_tools.d.ts.map +1 -0
- package/dist/services/search_tools.js +201 -0
- package/dist/services/search_tools.js.map +1 -0
- package/dist/services/settings.d.ts +53 -0
- package/dist/services/settings.d.ts.map +1 -0
- package/dist/services/settings.js +67 -0
- package/dist/services/settings.js.map +1 -0
- package/dist/services/sso_settings.d.ts +239 -0
- package/dist/services/sso_settings.d.ts.map +1 -0
- package/dist/services/sso_settings.js +214 -0
- package/dist/services/sso_settings.js.map +1 -0
- package/dist/services/tag_management.d.ts +390 -0
- package/dist/services/tag_management.d.ts.map +1 -0
- package/dist/services/tag_management.js +407 -0
- package/dist/services/tag_management.js.map +1 -0
- package/dist/services/team_management.d.ts +1502 -0
- package/dist/services/team_management.d.ts.map +1 -0
- package/dist/services/team_management.js +1429 -0
- package/dist/services/team_management.js.map +1 -0
- package/dist/services/tools.d.ts +227 -0
- package/dist/services/tools.d.ts.map +1 -0
- package/dist/services/tools.js +266 -0
- package/dist/services/tools.js.map +1 -0
- package/dist/services/ui_settings.d.ts +90 -0
- package/dist/services/ui_settings.d.ts.map +1 -0
- package/dist/services/ui_settings.js +96 -0
- package/dist/services/ui_settings.js.map +1 -0
- package/dist/services/ui_theme_settings.d.ts +62 -0
- package/dist/services/ui_theme_settings.d.ts.map +1 -0
- package/dist/services/ui_theme_settings.js +79 -0
- package/dist/services/ui_theme_settings.js.map +1 -0
- package/dist/services/usage_ai.d.ts +39 -0
- package/dist/services/usage_ai.d.ts.map +1 -0
- package/dist/services/usage_ai.js +40 -0
- package/dist/services/usage_ai.js.map +1 -0
- package/dist/services/vantage.d.ts +108 -0
- package/dist/services/vantage.d.ts.map +1 -0
- package/dist/services/vantage.js +124 -0
- package/dist/services/vantage.js.map +1 -0
- package/dist/services/vector_store_management.d.ts +161 -0
- package/dist/services/vector_store_management.d.ts.map +1 -0
- package/dist/services/vector_store_management.js +191 -0
- package/dist/services/vector_store_management.js.map +1 -0
- package/dist/services/vector_stores.d.ts +342 -0
- package/dist/services/vector_stores.d.ts.map +1 -0
- package/dist/services/vector_stores.js +639 -0
- package/dist/services/vector_stores.js.map +1 -0
- package/dist/services/vertex_ai_pass_through.d.ts +180 -0
- package/dist/services/vertex_ai_pass_through.d.ts.map +1 -0
- package/dist/services/vertex_ai_pass_through.js +334 -0
- package/dist/services/vertex_ai_pass_through.js.map +1 -0
- package/dist/services/videos.d.ts +207 -0
- package/dist/services/videos.d.ts.map +1 -0
- package/dist/services/videos.js +373 -0
- package/dist/services/videos.js.map +1 -0
- package/dist/services/vllm_pass_through.d.ts +70 -0
- package/dist/services/vllm_pass_through.d.ts.map +1 -0
- package/dist/services/vllm_pass_through.js +104 -0
- package/dist/services/vllm_pass_through.js.map +1 -0
- package/dist/services/watsonx_pass_through.d.ts +70 -0
- package/dist/services/watsonx_pass_through.d.ts.map +1 -0
- package/dist/services/watsonx_pass_through.js +114 -0
- package/dist/services/watsonx_pass_through.js.map +1 -0
- package/dist/services/web_socket.d.ts +77 -0
- package/dist/services/web_socket.d.ts.map +1 -0
- package/dist/services/web_socket.js +135 -0
- package/dist/services/web_socket.js.map +1 -0
- package/dist/services/workflow_management.d.ts +140 -0
- package/dist/services/workflow_management.d.ts.map +1 -0
- package/dist/services/workflow_management.js +220 -0
- package/dist/services/workflow_management.js.map +1 -0
- package/dist/traits.d.ts +14 -0
- package/dist/traits.d.ts.map +1 -0
- package/dist/traits.js +15 -0
- package/dist/traits.js.map +1 -0
- package/package.json +75 -0
- package/src/credentials.ts +97 -0
- package/src/errors.ts +108 -0
- package/src/index.ts +33 -0
- package/src/protocol.ts +160 -0
- package/src/retry.ts +74 -0
- package/src/services/a2a.ts +250 -0
- package/src/services/a2a_registration.ts +94 -0
- package/src/services/access_groups.ts +716 -0
- package/src/services/adaptive_router.ts +49 -0
- package/src/services/agents.ts +1314 -0
- package/src/services/alerting.ts +49 -0
- package/src/services/anthropic_pass_through.ts +233 -0
- package/src/services/anthropic_passthrough.ts +123 -0
- package/src/services/anthropic_skills.ts +200 -0
- package/src/services/assembly_ai_eu_pass_through.ts +237 -0
- package/src/services/assembly_ai_pass_through.ts +237 -0
- package/src/services/assistants.ts +685 -0
- package/src/services/audio.ts +189 -0
- package/src/services/audit_logging.ts +194 -0
- package/src/services/auto_router.ts +625 -0
- package/src/services/aws_comprehend_medical_pass_through.ts +109 -0
- package/src/services/azure_ai_pass_through.ts +231 -0
- package/src/services/azure_pass_through.ts +227 -0
- package/src/services/batch.ts +553 -0
- package/src/services/bedrock_pass_through.ts +229 -0
- package/src/services/beta_agents.ts +444 -0
- package/src/services/beta_mcp.ts +105 -0
- package/src/services/budget_management.ts +465 -0
- package/src/services/budget_spend_tracking.ts +2670 -0
- package/src/services/cache_settings.ts +234 -0
- package/src/services/caching.ts +175 -0
- package/src/services/chat_completions.ts +1788 -0
- package/src/services/claude_code_marketplace.ts +569 -0
- package/src/services/cloudzero.ts +297 -0
- package/src/services/cohere_pass_through.ts +227 -0
- package/src/services/completions.ts +194 -0
- package/src/services/compliance.ts +175 -0
- package/src/services/config_overrides.ts +466 -0
- package/src/services/config_yaml.ts +1537 -0
- package/src/services/containers.ts +848 -0
- package/src/services/coordination_redis_settings.ts +251 -0
- package/src/services/cost_tracking.ts +407 -0
- package/src/services/credential_management.ts +436 -0
- package/src/services/cursor_pass_through.ts +227 -0
- package/src/services/customer_management.ts +1475 -0
- package/src/services/email_management.ts +171 -0
- package/src/services/embeddings.ts +424 -0
- package/src/services/evals.ts +672 -0
- package/src/services/experimental.ts +572 -0
- package/src/services/fallback_management.ts +224 -0
- package/src/services/files.ts +732 -0
- package/src/services/fine_tuning.ts +547 -0
- package/src/services/gemini_agents.ts +227 -0
- package/src/services/google_ai_studio_pass_through.ts +227 -0
- package/src/services/google_genai_endpoints.ts +688 -0
- package/src/services/guardrails.ts +3006 -0
- package/src/services/health.ts +713 -0
- package/src/services/images.ts +430 -0
- package/src/services/index.ts +106 -0
- package/src/services/internal_user_management.ts +2574 -0
- package/src/services/invite_links.ts +167 -0
- package/src/services/jwt_mappings.ts +238 -0
- package/src/services/key_management.ts +4162 -0
- package/src/services/langfuse_passthrough.ts +231 -0
- package/src/services/llm_utils.ts +429 -0
- package/src/services/logging_callbacks.ts +104 -0
- package/src/services/mcp_byok_oauth.ts +468 -0
- package/src/services/mcp_discoverable.ts +626 -0
- package/src/services/mcp_management.ts +2915 -0
- package/src/services/mcp_rest.ts +779 -0
- package/src/services/memory_management.ts +248 -0
- package/src/services/milvus_pass_through.ts +227 -0
- package/src/services/misc.ts +2128 -0
- package/src/services/mistral_pass_through.ts +229 -0
- package/src/services/model_management.ts +4501 -0
- package/src/services/moderations.ts +80 -0
- package/src/services/ocr.ts +78 -0
- package/src/services/open_ai_pass_through.ts +462 -0
- package/src/services/organization_management.ts +1749 -0
- package/src/services/plugins.ts +367 -0
- package/src/services/policies.ts +1676 -0
- package/src/services/policy_engine.ts +1318 -0
- package/src/services/project_management.ts +938 -0
- package/src/services/prompts.ts +592 -0
- package/src/services/public.ts +880 -0
- package/src/services/rag.ts +148 -0
- package/src/services/realtime.ts +350 -0
- package/src/services/rerank.ts +111 -0
- package/src/services/responses.ts +909 -0
- package/src/services/router_settings.ts +168 -0
- package/src/services/rust_control_plane.ts +123 -0
- package/src/services/scim.ts +1243 -0
- package/src/services/search.ts +260 -0
- package/src/services/search_tools.ts +439 -0
- package/src/services/settings.ts +146 -0
- package/src/services/sso_settings.ts +626 -0
- package/src/services/tag_management.ts +1004 -0
- package/src/services/team_management.ts +3990 -0
- package/src/services/tools.ts +628 -0
- package/src/services/ui_settings.ts +237 -0
- package/src/services/ui_theme_settings.ts +173 -0
- package/src/services/usage_ai.ts +86 -0
- package/src/services/vantage.ts +283 -0
- package/src/services/vector_store_management.ts +454 -0
- package/src/services/vector_stores.ts +1321 -0
- package/src/services/vertex_ai_pass_through.ts +684 -0
- package/src/services/videos.ts +762 -0
- package/src/services/vllm_pass_through.ts +227 -0
- package/src/services/watsonx_pass_through.ts +229 -0
- package/src/services/web_socket.ts +254 -0
- package/src/services/workflow_management.ts +482 -0
- package/src/traits.ts +44 -0
|
@@ -0,0 +1,4501 @@
|
|
|
1
|
+
// AUTO-GENERATED by scripts/generate.ts from .generated-specs (pinned litellm-openapi @ 1.100.0 — see README). Do not edit.
|
|
2
|
+
import * as S from "@distilled.cloud/core/schema";
|
|
3
|
+
import * as API from "@distilled.cloud/core/api";
|
|
4
|
+
import * as C from "@distilled.cloud/core/category";
|
|
5
|
+
import * as T from "../traits.ts";
|
|
6
|
+
import {
|
|
7
|
+
LitellmProtocol,
|
|
8
|
+
type LitellmOpError,
|
|
9
|
+
type LitellmOpContext,
|
|
10
|
+
} from "../protocol.ts";
|
|
11
|
+
import * as Retry from "../retry.ts";
|
|
12
|
+
|
|
13
|
+
export type { LitellmOpError, LitellmOpContext };
|
|
14
|
+
|
|
15
|
+
export class BadRequest
|
|
16
|
+
extends /*@__PURE__*/ T.applyErrorMatchers(
|
|
17
|
+
/*@__PURE__*/ S.TaggedError<BadRequest>()("BadRequest", {
|
|
18
|
+
code: S.Number,
|
|
19
|
+
message: S.String,
|
|
20
|
+
}).pipe(C.withBadRequestError),
|
|
21
|
+
[{ status: 400 }],
|
|
22
|
+
) {}
|
|
23
|
+
|
|
24
|
+
export class UnprocessableEntity
|
|
25
|
+
extends /*@__PURE__*/ T.applyErrorMatchers(
|
|
26
|
+
/*@__PURE__*/ S.TaggedError<UnprocessableEntity>()("UnprocessableEntity", {
|
|
27
|
+
code: S.Number,
|
|
28
|
+
message: S.String,
|
|
29
|
+
}).pipe(C.withBadRequestError),
|
|
30
|
+
[{ status: 422 }],
|
|
31
|
+
) {}
|
|
32
|
+
|
|
33
|
+
export type LiteLLMParamsAdaptiveRouterConfigMap = {
|
|
34
|
+
[key: string]: unknown | undefined;
|
|
35
|
+
};
|
|
36
|
+
export const LiteLLMParamsAdaptiveRouterConfigMap = /*@__PURE__*/ S.Record(
|
|
37
|
+
S.String,
|
|
38
|
+
S.Unknown,
|
|
39
|
+
) as any as S.Schema<LiteLLMParamsAdaptiveRouterConfigMap>;
|
|
40
|
+
|
|
41
|
+
export type LiteLLMParamsBedrockTagsList = Array<unknown>;
|
|
42
|
+
export const LiteLLMParamsBedrockTagsList = /*@__PURE__*/ S.Array(
|
|
43
|
+
S.Unknown,
|
|
44
|
+
) as any as S.Schema<LiteLLMParamsBedrockTagsList>;
|
|
45
|
+
|
|
46
|
+
export type LiteLLMParamsComplexityRouterConfigMap = {
|
|
47
|
+
[key: string]: unknown | undefined;
|
|
48
|
+
};
|
|
49
|
+
export const LiteLLMParamsComplexityRouterConfigMap = /*@__PURE__*/ S.Record(
|
|
50
|
+
S.String,
|
|
51
|
+
S.Unknown,
|
|
52
|
+
) as any as S.Schema<LiteLLMParamsComplexityRouterConfigMap>;
|
|
53
|
+
|
|
54
|
+
export interface ConfigurableClientsideParamsCustomAuthInput {
|
|
55
|
+
api_base: string;
|
|
56
|
+
}
|
|
57
|
+
export const ConfigurableClientsideParamsCustomAuthInput =
|
|
58
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
59
|
+
S.Struct({
|
|
60
|
+
api_base: S.String,
|
|
61
|
+
}),
|
|
62
|
+
).annotate({
|
|
63
|
+
identifier: "ConfigurableClientsideParamsCustomAuthInput",
|
|
64
|
+
}) as any as S.Schema<ConfigurableClientsideParamsCustomAuthInput>;
|
|
65
|
+
|
|
66
|
+
export type LiteLLMParamsConfigurableClientsideAuthParamsItem =
|
|
67
|
+
| string
|
|
68
|
+
| ConfigurableClientsideParamsCustomAuthInput;
|
|
69
|
+
export const LiteLLMParamsConfigurableClientsideAuthParamsItem =
|
|
70
|
+
S.Unknown as any as S.Schema<LiteLLMParamsConfigurableClientsideAuthParamsItem>;
|
|
71
|
+
|
|
72
|
+
export type LiteLLMParamsConfigurableClientsideAuthParamsList =
|
|
73
|
+
Array<LiteLLMParamsConfigurableClientsideAuthParamsItem>;
|
|
74
|
+
export const LiteLLMParamsConfigurableClientsideAuthParamsList =
|
|
75
|
+
/*@__PURE__*/ S.Array(
|
|
76
|
+
LiteLLMParamsConfigurableClientsideAuthParamsItem,
|
|
77
|
+
) as any as S.Schema<LiteLLMParamsConfigurableClientsideAuthParamsList>;
|
|
78
|
+
|
|
79
|
+
export type LiteLLMParamsMilvusPartitionNamesList = Array<string>;
|
|
80
|
+
export const LiteLLMParamsMilvusPartitionNamesList = /*@__PURE__*/ S.Array(
|
|
81
|
+
S.String,
|
|
82
|
+
) as any as S.Schema<LiteLLMParamsMilvusPartitionNamesList>;
|
|
83
|
+
|
|
84
|
+
export type ChoicesFinishReason =
|
|
85
|
+
| "stop"
|
|
86
|
+
| "content_filter"
|
|
87
|
+
| "function_call"
|
|
88
|
+
| "tool_calls"
|
|
89
|
+
| "length"
|
|
90
|
+
| "guardrail_intervened"
|
|
91
|
+
| "eos"
|
|
92
|
+
| "finish_reason_unspecified"
|
|
93
|
+
| "malformed_function_call";
|
|
94
|
+
export const ChoicesFinishReason = S.String;
|
|
95
|
+
|
|
96
|
+
export type ChatCompletionTokenLogprobBytesList = Array<number>;
|
|
97
|
+
export const ChatCompletionTokenLogprobBytesList = /*@__PURE__*/ S.Array(
|
|
98
|
+
S.Number,
|
|
99
|
+
) as any as S.Schema<ChatCompletionTokenLogprobBytesList>;
|
|
100
|
+
|
|
101
|
+
export type TopLogprobBytesList = Array<number>;
|
|
102
|
+
export const TopLogprobBytesList = /*@__PURE__*/ S.Array(
|
|
103
|
+
S.Number,
|
|
104
|
+
) as any as S.Schema<TopLogprobBytesList>;
|
|
105
|
+
|
|
106
|
+
export interface TopLogprob {
|
|
107
|
+
bytes?: TopLogprobBytesList | null;
|
|
108
|
+
logprob: number;
|
|
109
|
+
token: string;
|
|
110
|
+
}
|
|
111
|
+
export const TopLogprob = /*@__PURE__*/ S.suspend(() =>
|
|
112
|
+
S.Struct({
|
|
113
|
+
bytes: S.optional(S.NullOr(TopLogprobBytesList)),
|
|
114
|
+
logprob: S.Number,
|
|
115
|
+
token: S.String,
|
|
116
|
+
}),
|
|
117
|
+
).annotate({ identifier: "TopLogprob" }) as any as S.Schema<TopLogprob>;
|
|
118
|
+
|
|
119
|
+
export type ChatCompletionTokenLogprobTopLogprobsList = Array<TopLogprob>;
|
|
120
|
+
export const ChatCompletionTokenLogprobTopLogprobsList = /*@__PURE__*/ S.Array(
|
|
121
|
+
TopLogprob,
|
|
122
|
+
) as any as S.Schema<ChatCompletionTokenLogprobTopLogprobsList>;
|
|
123
|
+
|
|
124
|
+
export interface ChatCompletionTokenLogprob {
|
|
125
|
+
bytes?: ChatCompletionTokenLogprobBytesList | null;
|
|
126
|
+
logprob: number;
|
|
127
|
+
token: string;
|
|
128
|
+
top_logprobs: ChatCompletionTokenLogprobTopLogprobsList;
|
|
129
|
+
}
|
|
130
|
+
export const ChatCompletionTokenLogprob = /*@__PURE__*/ S.suspend(() =>
|
|
131
|
+
S.Struct({
|
|
132
|
+
bytes: S.optional(S.NullOr(ChatCompletionTokenLogprobBytesList)),
|
|
133
|
+
logprob: S.Number,
|
|
134
|
+
token: S.String,
|
|
135
|
+
top_logprobs: ChatCompletionTokenLogprobTopLogprobsList,
|
|
136
|
+
}),
|
|
137
|
+
).annotate({
|
|
138
|
+
identifier: "ChatCompletionTokenLogprob",
|
|
139
|
+
}) as any as S.Schema<ChatCompletionTokenLogprob>;
|
|
140
|
+
|
|
141
|
+
export type ChoiceLogprobsContentList = Array<ChatCompletionTokenLogprob>;
|
|
142
|
+
export const ChoiceLogprobsContentList = /*@__PURE__*/ S.Array(
|
|
143
|
+
ChatCompletionTokenLogprob,
|
|
144
|
+
) as any as S.Schema<ChoiceLogprobsContentList>;
|
|
145
|
+
|
|
146
|
+
export interface ChoiceLogprobs {
|
|
147
|
+
content?: ChoiceLogprobsContentList | null;
|
|
148
|
+
}
|
|
149
|
+
export const ChoiceLogprobs = /*@__PURE__*/ S.suspend(() =>
|
|
150
|
+
S.Struct({
|
|
151
|
+
content: S.optional(S.NullOr(ChoiceLogprobsContentList)),
|
|
152
|
+
}),
|
|
153
|
+
).annotate({ identifier: "ChoiceLogprobs" }) as any as S.Schema<ChoiceLogprobs>;
|
|
154
|
+
|
|
155
|
+
export type ChoicesLogprobs = ChoiceLogprobs | unknown;
|
|
156
|
+
export const ChoicesLogprobs = S.Unknown as any as S.Schema<ChoicesLogprobs>;
|
|
157
|
+
|
|
158
|
+
export interface ChatCompletionAnnotationURLCitation {
|
|
159
|
+
end_index?: number;
|
|
160
|
+
start_index?: number;
|
|
161
|
+
title?: string;
|
|
162
|
+
url?: string;
|
|
163
|
+
}
|
|
164
|
+
export const ChatCompletionAnnotationURLCitation = /*@__PURE__*/ S.suspend(() =>
|
|
165
|
+
S.Struct({
|
|
166
|
+
end_index: S.optional(S.Number),
|
|
167
|
+
start_index: S.optional(S.Number),
|
|
168
|
+
title: S.optional(S.String),
|
|
169
|
+
url: S.optional(S.String),
|
|
170
|
+
}),
|
|
171
|
+
).annotate({
|
|
172
|
+
identifier: "ChatCompletionAnnotationURLCitation",
|
|
173
|
+
}) as any as S.Schema<ChatCompletionAnnotationURLCitation>;
|
|
174
|
+
|
|
175
|
+
export interface ChatCompletionAnnotation {
|
|
176
|
+
type?: string;
|
|
177
|
+
url_citation?: ChatCompletionAnnotationURLCitation;
|
|
178
|
+
}
|
|
179
|
+
export const ChatCompletionAnnotation = /*@__PURE__*/ S.suspend(() =>
|
|
180
|
+
S.Struct({
|
|
181
|
+
type: S.optional(S.String),
|
|
182
|
+
url_citation: S.optional(ChatCompletionAnnotationURLCitation),
|
|
183
|
+
}),
|
|
184
|
+
).annotate({
|
|
185
|
+
identifier: "ChatCompletionAnnotation",
|
|
186
|
+
}) as any as S.Schema<ChatCompletionAnnotation>;
|
|
187
|
+
|
|
188
|
+
export type MessageAnnotationsList = Array<ChatCompletionAnnotation>;
|
|
189
|
+
export const MessageAnnotationsList = /*@__PURE__*/ S.Array(
|
|
190
|
+
ChatCompletionAnnotation,
|
|
191
|
+
) as any as S.Schema<MessageAnnotationsList>;
|
|
192
|
+
|
|
193
|
+
export interface ChatCompletionAudioResponse {
|
|
194
|
+
data: string;
|
|
195
|
+
expires_at: number;
|
|
196
|
+
id: string;
|
|
197
|
+
transcript: string;
|
|
198
|
+
}
|
|
199
|
+
export const ChatCompletionAudioResponse = /*@__PURE__*/ S.suspend(() =>
|
|
200
|
+
S.Struct({
|
|
201
|
+
data: S.String,
|
|
202
|
+
expires_at: S.Number,
|
|
203
|
+
id: S.String,
|
|
204
|
+
transcript: S.String,
|
|
205
|
+
}),
|
|
206
|
+
).annotate({
|
|
207
|
+
identifier: "ChatCompletionAudioResponse",
|
|
208
|
+
}) as any as S.Schema<ChatCompletionAudioResponse>;
|
|
209
|
+
|
|
210
|
+
export interface FunctionCall {
|
|
211
|
+
arguments: string;
|
|
212
|
+
name?: string | null;
|
|
213
|
+
}
|
|
214
|
+
export const FunctionCall = /*@__PURE__*/ S.suspend(() =>
|
|
215
|
+
S.Struct({
|
|
216
|
+
arguments: S.String,
|
|
217
|
+
name: S.optional(S.NullOr(S.String)),
|
|
218
|
+
}),
|
|
219
|
+
).annotate({ identifier: "FunctionCall" }) as any as S.Schema<FunctionCall>;
|
|
220
|
+
|
|
221
|
+
export interface ImageURLObject {
|
|
222
|
+
detail?: string | null;
|
|
223
|
+
url: string;
|
|
224
|
+
}
|
|
225
|
+
export const ImageURLObject = /*@__PURE__*/ S.suspend(() =>
|
|
226
|
+
S.Struct({
|
|
227
|
+
detail: S.optional(S.NullOr(S.String)),
|
|
228
|
+
url: S.String,
|
|
229
|
+
}),
|
|
230
|
+
).annotate({ identifier: "ImageURLObject" }) as any as S.Schema<ImageURLObject>;
|
|
231
|
+
|
|
232
|
+
export interface ImageURLListItem {
|
|
233
|
+
image_url: ImageURLObject;
|
|
234
|
+
index: number;
|
|
235
|
+
type: string;
|
|
236
|
+
}
|
|
237
|
+
export const ImageURLListItem = /*@__PURE__*/ S.suspend(() =>
|
|
238
|
+
S.Struct({
|
|
239
|
+
image_url: ImageURLObject,
|
|
240
|
+
index: S.Number,
|
|
241
|
+
type: S.String,
|
|
242
|
+
}),
|
|
243
|
+
).annotate({
|
|
244
|
+
identifier: "ImageURLListItem",
|
|
245
|
+
}) as any as S.Schema<ImageURLListItem>;
|
|
246
|
+
|
|
247
|
+
export type MessageImagesList = Array<ImageURLListItem>;
|
|
248
|
+
export const MessageImagesList = /*@__PURE__*/ S.Array(
|
|
249
|
+
ImageURLListItem,
|
|
250
|
+
) as any as S.Schema<MessageImagesList>;
|
|
251
|
+
|
|
252
|
+
export type MessageProviderSpecificFieldsMap = {
|
|
253
|
+
[key: string]: unknown | undefined;
|
|
254
|
+
};
|
|
255
|
+
export const MessageProviderSpecificFieldsMap = /*@__PURE__*/ S.Record(
|
|
256
|
+
S.String,
|
|
257
|
+
S.Unknown,
|
|
258
|
+
) as any as S.Schema<MessageProviderSpecificFieldsMap>;
|
|
259
|
+
|
|
260
|
+
export interface ChatCompletionReasoningSummaryTextBlock {
|
|
261
|
+
text?: string;
|
|
262
|
+
type: string;
|
|
263
|
+
}
|
|
264
|
+
export const ChatCompletionReasoningSummaryTextBlock = /*@__PURE__*/ S.suspend(
|
|
265
|
+
() =>
|
|
266
|
+
S.Struct({
|
|
267
|
+
text: S.optional(S.String),
|
|
268
|
+
type: S.String,
|
|
269
|
+
}),
|
|
270
|
+
).annotate({
|
|
271
|
+
identifier: "ChatCompletionReasoningSummaryTextBlock",
|
|
272
|
+
}) as any as S.Schema<ChatCompletionReasoningSummaryTextBlock>;
|
|
273
|
+
|
|
274
|
+
export type ChatCompletionReasoningItemSummaryList =
|
|
275
|
+
Array<ChatCompletionReasoningSummaryTextBlock>;
|
|
276
|
+
export const ChatCompletionReasoningItemSummaryList = /*@__PURE__*/ S.Array(
|
|
277
|
+
ChatCompletionReasoningSummaryTextBlock,
|
|
278
|
+
) as any as S.Schema<ChatCompletionReasoningItemSummaryList>;
|
|
279
|
+
|
|
280
|
+
/** Represents an OpenAI Responses API reasoning item for round-tripping in conversation history. */
|
|
281
|
+
export interface ChatCompletionReasoningItem {
|
|
282
|
+
encrypted_content?: string | null;
|
|
283
|
+
id?: string;
|
|
284
|
+
summary?: ChatCompletionReasoningItemSummaryList;
|
|
285
|
+
type: string;
|
|
286
|
+
}
|
|
287
|
+
export const ChatCompletionReasoningItem = /*@__PURE__*/ S.suspend(() =>
|
|
288
|
+
S.Struct({
|
|
289
|
+
encrypted_content: S.optional(S.NullOr(S.String)),
|
|
290
|
+
id: S.optional(S.String),
|
|
291
|
+
summary: S.optional(ChatCompletionReasoningItemSummaryList),
|
|
292
|
+
type: S.String,
|
|
293
|
+
}),
|
|
294
|
+
).annotate({
|
|
295
|
+
identifier: "ChatCompletionReasoningItem",
|
|
296
|
+
}) as any as S.Schema<ChatCompletionReasoningItem>;
|
|
297
|
+
|
|
298
|
+
export type MessageReasoningItemsList = Array<ChatCompletionReasoningItem>;
|
|
299
|
+
export const MessageReasoningItemsList = /*@__PURE__*/ S.Array(
|
|
300
|
+
ChatCompletionReasoningItem,
|
|
301
|
+
) as any as S.Schema<MessageReasoningItemsList>;
|
|
302
|
+
|
|
303
|
+
export type MessageRole = "assistant" | "user" | "system" | "tool" | "function";
|
|
304
|
+
export const MessageRole = S.String;
|
|
305
|
+
|
|
306
|
+
export type ChatCompletionThinkingBlockCacheControlCase0Map = {
|
|
307
|
+
[key: string]: unknown | undefined;
|
|
308
|
+
};
|
|
309
|
+
export const ChatCompletionThinkingBlockCacheControlCase0Map =
|
|
310
|
+
/*@__PURE__*/ S.Record(
|
|
311
|
+
S.String,
|
|
312
|
+
S.Unknown,
|
|
313
|
+
) as any as S.Schema<ChatCompletionThinkingBlockCacheControlCase0Map>;
|
|
314
|
+
|
|
315
|
+
export type ChatCompletionCachedContentTtl = "5m" | "1h";
|
|
316
|
+
export const ChatCompletionCachedContentTtl = S.String;
|
|
317
|
+
|
|
318
|
+
export interface ChatCompletionCachedContent {
|
|
319
|
+
ttl?: ChatCompletionCachedContentTtl | (string & {});
|
|
320
|
+
type: string;
|
|
321
|
+
}
|
|
322
|
+
export const ChatCompletionCachedContent = /*@__PURE__*/ S.suspend(() =>
|
|
323
|
+
S.Struct({
|
|
324
|
+
ttl: S.optional(ChatCompletionCachedContentTtl),
|
|
325
|
+
type: S.String,
|
|
326
|
+
}),
|
|
327
|
+
).annotate({
|
|
328
|
+
identifier: "ChatCompletionCachedContent",
|
|
329
|
+
}) as any as S.Schema<ChatCompletionCachedContent>;
|
|
330
|
+
|
|
331
|
+
export type ChatCompletionThinkingBlockCacheControl =
|
|
332
|
+
| ChatCompletionThinkingBlockCacheControlCase0Map
|
|
333
|
+
| ChatCompletionCachedContent;
|
|
334
|
+
export const ChatCompletionThinkingBlockCacheControl =
|
|
335
|
+
S.Unknown as any as S.Schema<ChatCompletionThinkingBlockCacheControl>;
|
|
336
|
+
|
|
337
|
+
export interface ChatCompletionThinkingBlock {
|
|
338
|
+
cache_control?: ChatCompletionThinkingBlockCacheControl | null;
|
|
339
|
+
signature?: string | null;
|
|
340
|
+
thinking?: string;
|
|
341
|
+
type: string;
|
|
342
|
+
}
|
|
343
|
+
export const ChatCompletionThinkingBlock = /*@__PURE__*/ S.suspend(() =>
|
|
344
|
+
S.Struct({
|
|
345
|
+
cache_control: S.optional(
|
|
346
|
+
S.NullOr(ChatCompletionThinkingBlockCacheControl),
|
|
347
|
+
),
|
|
348
|
+
signature: S.optional(S.NullOr(S.String)),
|
|
349
|
+
thinking: S.optional(S.String),
|
|
350
|
+
type: S.String,
|
|
351
|
+
}),
|
|
352
|
+
).annotate({
|
|
353
|
+
identifier: "ChatCompletionThinkingBlock",
|
|
354
|
+
}) as any as S.Schema<ChatCompletionThinkingBlock>;
|
|
355
|
+
|
|
356
|
+
export type ChatCompletionRedactedThinkingBlockCacheControlCase0Map = {
|
|
357
|
+
[key: string]: unknown | undefined;
|
|
358
|
+
};
|
|
359
|
+
export const ChatCompletionRedactedThinkingBlockCacheControlCase0Map =
|
|
360
|
+
/*@__PURE__*/ S.Record(
|
|
361
|
+
S.String,
|
|
362
|
+
S.Unknown,
|
|
363
|
+
) as any as S.Schema<ChatCompletionRedactedThinkingBlockCacheControlCase0Map>;
|
|
364
|
+
|
|
365
|
+
export type ChatCompletionRedactedThinkingBlockCacheControl =
|
|
366
|
+
| ChatCompletionRedactedThinkingBlockCacheControlCase0Map
|
|
367
|
+
| ChatCompletionCachedContent;
|
|
368
|
+
export const ChatCompletionRedactedThinkingBlockCacheControl =
|
|
369
|
+
S.Unknown as any as S.Schema<ChatCompletionRedactedThinkingBlockCacheControl>;
|
|
370
|
+
|
|
371
|
+
export interface ChatCompletionRedactedThinkingBlock {
|
|
372
|
+
cache_control?: ChatCompletionRedactedThinkingBlockCacheControl | null;
|
|
373
|
+
data?: string;
|
|
374
|
+
type: string;
|
|
375
|
+
}
|
|
376
|
+
export const ChatCompletionRedactedThinkingBlock = /*@__PURE__*/ S.suspend(() =>
|
|
377
|
+
S.Struct({
|
|
378
|
+
cache_control: S.optional(
|
|
379
|
+
S.NullOr(ChatCompletionRedactedThinkingBlockCacheControl),
|
|
380
|
+
),
|
|
381
|
+
data: S.optional(S.String),
|
|
382
|
+
type: S.String,
|
|
383
|
+
}),
|
|
384
|
+
).annotate({
|
|
385
|
+
identifier: "ChatCompletionRedactedThinkingBlock",
|
|
386
|
+
}) as any as S.Schema<ChatCompletionRedactedThinkingBlock>;
|
|
387
|
+
|
|
388
|
+
export type MessageThinkingBlocksItem =
|
|
389
|
+
| ChatCompletionThinkingBlock
|
|
390
|
+
| ChatCompletionRedactedThinkingBlock;
|
|
391
|
+
export const MessageThinkingBlocksItem =
|
|
392
|
+
S.Unknown as any as S.Schema<MessageThinkingBlocksItem>;
|
|
393
|
+
|
|
394
|
+
export type MessageThinkingBlocksList = Array<MessageThinkingBlocksItem>;
|
|
395
|
+
export const MessageThinkingBlocksList = /*@__PURE__*/ S.Array(
|
|
396
|
+
MessageThinkingBlocksItem,
|
|
397
|
+
) as any as S.Schema<MessageThinkingBlocksList>;
|
|
398
|
+
|
|
399
|
+
export type ChatCompletionMessageToolCall = {
|
|
400
|
+
[key: string]: unknown | undefined;
|
|
401
|
+
};
|
|
402
|
+
export const ChatCompletionMessageToolCall = /*@__PURE__*/ S.Record(
|
|
403
|
+
S.String,
|
|
404
|
+
S.Unknown,
|
|
405
|
+
) as any as S.Schema<ChatCompletionMessageToolCall>;
|
|
406
|
+
|
|
407
|
+
export interface ChatCompletionCustomToolCallPayload {
|
|
408
|
+
input: string;
|
|
409
|
+
name: string;
|
|
410
|
+
}
|
|
411
|
+
export const ChatCompletionCustomToolCallPayload = /*@__PURE__*/ S.suspend(() =>
|
|
412
|
+
S.Struct({
|
|
413
|
+
input: S.String,
|
|
414
|
+
name: S.String,
|
|
415
|
+
}),
|
|
416
|
+
).annotate({
|
|
417
|
+
identifier: "ChatCompletionCustomToolCallPayload",
|
|
418
|
+
}) as any as S.Schema<ChatCompletionCustomToolCallPayload>;
|
|
419
|
+
|
|
420
|
+
export interface ChatCompletionMessageCustomToolCall {
|
|
421
|
+
custom: ChatCompletionCustomToolCallPayload;
|
|
422
|
+
id: string;
|
|
423
|
+
type?: string;
|
|
424
|
+
}
|
|
425
|
+
export const ChatCompletionMessageCustomToolCall = /*@__PURE__*/ S.suspend(() =>
|
|
426
|
+
S.Struct({
|
|
427
|
+
custom: ChatCompletionCustomToolCallPayload,
|
|
428
|
+
id: S.String,
|
|
429
|
+
type: S.optional(S.String),
|
|
430
|
+
}),
|
|
431
|
+
).annotate({
|
|
432
|
+
identifier: "ChatCompletionMessageCustomToolCall",
|
|
433
|
+
}) as any as S.Schema<ChatCompletionMessageCustomToolCall>;
|
|
434
|
+
|
|
435
|
+
export type MessageToolCallsItem =
|
|
436
|
+
| ChatCompletionMessageToolCall
|
|
437
|
+
| ChatCompletionMessageCustomToolCall;
|
|
438
|
+
export const MessageToolCallsItem =
|
|
439
|
+
S.Unknown as any as S.Schema<MessageToolCallsItem>;
|
|
440
|
+
|
|
441
|
+
export type MessageToolCallsList = Array<MessageToolCallsItem>;
|
|
442
|
+
export const MessageToolCallsList = /*@__PURE__*/ S.Array(
|
|
443
|
+
MessageToolCallsItem,
|
|
444
|
+
) as any as S.Schema<MessageToolCallsList>;
|
|
445
|
+
|
|
446
|
+
export interface Message {
|
|
447
|
+
annotations?: MessageAnnotationsList | null;
|
|
448
|
+
audio?: ChatCompletionAudioResponse | null;
|
|
449
|
+
content: string | null;
|
|
450
|
+
function_call: FunctionCall | null;
|
|
451
|
+
images?: MessageImagesList | null;
|
|
452
|
+
provider_specific_fields?: MessageProviderSpecificFieldsMap | null;
|
|
453
|
+
reasoning_content?: string | null;
|
|
454
|
+
reasoning_items?: MessageReasoningItemsList | null;
|
|
455
|
+
role: MessageRole | (string & {});
|
|
456
|
+
thinking_blocks?: MessageThinkingBlocksList | null;
|
|
457
|
+
tool_calls: MessageToolCallsList | null;
|
|
458
|
+
}
|
|
459
|
+
export const Message = /*@__PURE__*/ S.suspend(() =>
|
|
460
|
+
S.Struct({
|
|
461
|
+
annotations: S.optional(S.NullOr(MessageAnnotationsList)),
|
|
462
|
+
audio: S.optional(S.NullOr(ChatCompletionAudioResponse)),
|
|
463
|
+
content: S.NullOr(S.String),
|
|
464
|
+
function_call: S.NullOr(FunctionCall),
|
|
465
|
+
images: S.optional(S.NullOr(MessageImagesList)),
|
|
466
|
+
provider_specific_fields: S.optional(
|
|
467
|
+
S.NullOr(MessageProviderSpecificFieldsMap),
|
|
468
|
+
),
|
|
469
|
+
reasoning_content: S.optional(S.NullOr(S.String)),
|
|
470
|
+
reasoning_items: S.optional(S.NullOr(MessageReasoningItemsList)),
|
|
471
|
+
role: MessageRole,
|
|
472
|
+
thinking_blocks: S.optional(S.NullOr(MessageThinkingBlocksList)),
|
|
473
|
+
tool_calls: S.NullOr(MessageToolCallsList),
|
|
474
|
+
}),
|
|
475
|
+
).annotate({ identifier: "Message" }) as any as S.Schema<Message>;
|
|
476
|
+
|
|
477
|
+
export type ChoicesProviderSpecificFieldsMap = {
|
|
478
|
+
[key: string]: unknown | undefined;
|
|
479
|
+
};
|
|
480
|
+
export const ChoicesProviderSpecificFieldsMap = /*@__PURE__*/ S.Record(
|
|
481
|
+
S.String,
|
|
482
|
+
S.Unknown,
|
|
483
|
+
) as any as S.Schema<ChoicesProviderSpecificFieldsMap>;
|
|
484
|
+
|
|
485
|
+
export interface Choices {
|
|
486
|
+
finish_reason: ChoicesFinishReason | (string & {});
|
|
487
|
+
index: number;
|
|
488
|
+
logprobs?: ChoicesLogprobs | null;
|
|
489
|
+
message: Message;
|
|
490
|
+
provider_specific_fields?: ChoicesProviderSpecificFieldsMap | null;
|
|
491
|
+
}
|
|
492
|
+
export const Choices = /*@__PURE__*/ S.suspend(() =>
|
|
493
|
+
S.Struct({
|
|
494
|
+
finish_reason: ChoicesFinishReason,
|
|
495
|
+
index: S.Number,
|
|
496
|
+
logprobs: S.optional(S.NullOr(ChoicesLogprobs)),
|
|
497
|
+
message: Message,
|
|
498
|
+
provider_specific_fields: S.optional(
|
|
499
|
+
S.NullOr(ChoicesProviderSpecificFieldsMap),
|
|
500
|
+
),
|
|
501
|
+
}),
|
|
502
|
+
).annotate({ identifier: "Choices" }) as any as S.Schema<Choices>;
|
|
503
|
+
|
|
504
|
+
export type ModelResponseChoicesList = Array<Choices>;
|
|
505
|
+
export const ModelResponseChoicesList = /*@__PURE__*/ S.Array(
|
|
506
|
+
Choices,
|
|
507
|
+
) as any as S.Schema<ModelResponseChoicesList>;
|
|
508
|
+
|
|
509
|
+
export interface ModelResponse {
|
|
510
|
+
choices: ModelResponseChoicesList;
|
|
511
|
+
created: number;
|
|
512
|
+
id: string;
|
|
513
|
+
model?: string | null;
|
|
514
|
+
object: string;
|
|
515
|
+
system_fingerprint?: string | null;
|
|
516
|
+
}
|
|
517
|
+
export const ModelResponse = /*@__PURE__*/ S.suspend(() =>
|
|
518
|
+
S.Struct({
|
|
519
|
+
choices: ModelResponseChoicesList,
|
|
520
|
+
created: S.Number,
|
|
521
|
+
id: S.String,
|
|
522
|
+
model: S.optional(S.NullOr(S.String)),
|
|
523
|
+
object: S.String,
|
|
524
|
+
system_fingerprint: S.optional(S.NullOr(S.String)),
|
|
525
|
+
}),
|
|
526
|
+
).annotate({ identifier: "ModelResponse" }) as any as S.Schema<ModelResponse>;
|
|
527
|
+
|
|
528
|
+
export type LiteLLMParamsMockResponse = string | ModelResponse | unknown;
|
|
529
|
+
export const LiteLLMParamsMockResponse =
|
|
530
|
+
S.Unknown as any as S.Schema<LiteLLMParamsMockResponse>;
|
|
531
|
+
|
|
532
|
+
export type LiteLLMParamsModelInfoMap = { [key: string]: unknown | undefined };
|
|
533
|
+
export const LiteLLMParamsModelInfoMap = /*@__PURE__*/ S.Record(
|
|
534
|
+
S.String,
|
|
535
|
+
S.Unknown,
|
|
536
|
+
) as any as S.Schema<LiteLLMParamsModelInfoMap>;
|
|
537
|
+
|
|
538
|
+
export type LiteLLMParamsQualityRouterConfigMap = {
|
|
539
|
+
[key: string]: unknown | undefined;
|
|
540
|
+
};
|
|
541
|
+
export const LiteLLMParamsQualityRouterConfigMap = /*@__PURE__*/ S.Record(
|
|
542
|
+
S.String,
|
|
543
|
+
S.Unknown,
|
|
544
|
+
) as any as S.Schema<LiteLLMParamsQualityRouterConfigMap>;
|
|
545
|
+
|
|
546
|
+
export type LiteLLMParamsSearchContextCostPerQueryMap = {
|
|
547
|
+
[key: string]: unknown | undefined;
|
|
548
|
+
};
|
|
549
|
+
export const LiteLLMParamsSearchContextCostPerQueryMap = /*@__PURE__*/ S.Record(
|
|
550
|
+
S.String,
|
|
551
|
+
S.Unknown,
|
|
552
|
+
) as any as S.Schema<LiteLLMParamsSearchContextCostPerQueryMap>;
|
|
553
|
+
|
|
554
|
+
export type LiteLLMParamsStreamTimeout = number | string;
|
|
555
|
+
export const LiteLLMParamsStreamTimeout =
|
|
556
|
+
S.Unknown as any as S.Schema<LiteLLMParamsStreamTimeout>;
|
|
557
|
+
|
|
558
|
+
export type LiteLLMParamsTagRegexList = Array<string>;
|
|
559
|
+
export const LiteLLMParamsTagRegexList = /*@__PURE__*/ S.Array(
|
|
560
|
+
S.String,
|
|
561
|
+
) as any as S.Schema<LiteLLMParamsTagRegexList>;
|
|
562
|
+
|
|
563
|
+
export type LiteLLMParamsTagsList = Array<string>;
|
|
564
|
+
export const LiteLLMParamsTagsList = /*@__PURE__*/ S.Array(
|
|
565
|
+
S.String,
|
|
566
|
+
) as any as S.Schema<LiteLLMParamsTagsList>;
|
|
567
|
+
|
|
568
|
+
export type LiteLLMParamsTieredPricingItemMap = {
|
|
569
|
+
[key: string]: unknown | undefined;
|
|
570
|
+
};
|
|
571
|
+
export const LiteLLMParamsTieredPricingItemMap = /*@__PURE__*/ S.Record(
|
|
572
|
+
S.String,
|
|
573
|
+
S.Unknown,
|
|
574
|
+
) as any as S.Schema<LiteLLMParamsTieredPricingItemMap>;
|
|
575
|
+
|
|
576
|
+
export type LiteLLMParamsTieredPricingList =
|
|
577
|
+
Array<LiteLLMParamsTieredPricingItemMap>;
|
|
578
|
+
export const LiteLLMParamsTieredPricingList = /*@__PURE__*/ S.Array(
|
|
579
|
+
LiteLLMParamsTieredPricingItemMap,
|
|
580
|
+
) as any as S.Schema<LiteLLMParamsTieredPricingList>;
|
|
581
|
+
|
|
582
|
+
export type LiteLLMParamsTimeout = number | string;
|
|
583
|
+
export const LiteLLMParamsTimeout =
|
|
584
|
+
S.Unknown as any as S.Schema<LiteLLMParamsTimeout>;
|
|
585
|
+
|
|
586
|
+
export type LiteLLMParamsVertexCredentialsCase1Map = {
|
|
587
|
+
[key: string]: unknown | undefined;
|
|
588
|
+
};
|
|
589
|
+
export const LiteLLMParamsVertexCredentialsCase1Map = /*@__PURE__*/ S.Record(
|
|
590
|
+
S.String,
|
|
591
|
+
S.Unknown,
|
|
592
|
+
) as any as S.Schema<LiteLLMParamsVertexCredentialsCase1Map>;
|
|
593
|
+
|
|
594
|
+
export type LiteLLMParamsVertexCredentials =
|
|
595
|
+
| string
|
|
596
|
+
| LiteLLMParamsVertexCredentialsCase1Map;
|
|
597
|
+
export const LiteLLMParamsVertexCredentials =
|
|
598
|
+
S.Unknown as any as S.Schema<LiteLLMParamsVertexCredentials>;
|
|
599
|
+
|
|
600
|
+
/** LiteLLM Params with 'model' requirement - used for completions */
|
|
601
|
+
export interface LiteLLMParams {
|
|
602
|
+
adaptive_router_config?: LiteLLMParamsAdaptiveRouterConfigMap | null;
|
|
603
|
+
adaptive_router_default_model?: string | null;
|
|
604
|
+
allow_client_keepalive_override?: boolean | null;
|
|
605
|
+
annotation_cost_per_page?: number | null;
|
|
606
|
+
api_base?: string | null;
|
|
607
|
+
api_key?: string | null;
|
|
608
|
+
api_version?: string | null;
|
|
609
|
+
auto_router_config?: string | null;
|
|
610
|
+
auto_router_config_path?: string | null;
|
|
611
|
+
auto_router_default_model?: string | null;
|
|
612
|
+
auto_router_embedding_model?: string | null;
|
|
613
|
+
auto_router_max_input_chars?: number | null;
|
|
614
|
+
aws_access_key_id?: string | null;
|
|
615
|
+
aws_batch_role_arn?: string | null;
|
|
616
|
+
aws_bedrock_project_id?: string | null;
|
|
617
|
+
aws_bedrock_runtime_endpoint?: string | null;
|
|
618
|
+
aws_external_id?: string | null;
|
|
619
|
+
aws_profile_name?: string | null;
|
|
620
|
+
aws_region_name?: string | null;
|
|
621
|
+
aws_role_name?: string | null;
|
|
622
|
+
aws_secret_access_key?: string | null;
|
|
623
|
+
aws_session_name?: string | null;
|
|
624
|
+
aws_session_token?: string | null;
|
|
625
|
+
aws_sts_endpoint?: string | null;
|
|
626
|
+
aws_web_identity_token?: string | null;
|
|
627
|
+
azure_ad_token?: string | null;
|
|
628
|
+
bedrock_tags?: LiteLLMParamsBedrockTagsList | null;
|
|
629
|
+
budget_duration?: string | null;
|
|
630
|
+
cache_creation_input_audio_token_cost?: number | null;
|
|
631
|
+
cache_creation_input_token_cost?: number | null;
|
|
632
|
+
cache_creation_input_token_cost_above_1hr?: number | null;
|
|
633
|
+
cache_creation_input_token_cost_above_200k_tokens?: number | null;
|
|
634
|
+
cache_creation_input_token_cost_above_272k_tokens?: number | null;
|
|
635
|
+
cache_creation_input_token_cost_above_272k_tokens_flex?: number | null;
|
|
636
|
+
cache_creation_input_token_cost_above_272k_tokens_priority?: number | null;
|
|
637
|
+
cache_creation_input_token_cost_flex?: number | null;
|
|
638
|
+
cache_creation_input_token_cost_priority?: number | null;
|
|
639
|
+
cache_creation_input_token_cost_ultrafast?: number | null;
|
|
640
|
+
cache_read_input_audio_token_cost?: number | null;
|
|
641
|
+
cache_read_input_token_cost?: number | null;
|
|
642
|
+
cache_read_input_token_cost_above_200k_tokens?: number | null;
|
|
643
|
+
cache_read_input_token_cost_above_200k_tokens_priority?: number | null;
|
|
644
|
+
cache_read_input_token_cost_above_272k_tokens?: number | null;
|
|
645
|
+
cache_read_input_token_cost_above_272k_tokens_flex?: number | null;
|
|
646
|
+
cache_read_input_token_cost_above_272k_tokens_priority?: number | null;
|
|
647
|
+
cache_read_input_token_cost_above_512k_tokens?: number | null;
|
|
648
|
+
cache_read_input_token_cost_flex?: number | null;
|
|
649
|
+
cache_read_input_token_cost_priority?: number | null;
|
|
650
|
+
cache_read_input_token_cost_ultrafast?: number | null;
|
|
651
|
+
citation_cost_per_token?: number | null;
|
|
652
|
+
complexity_router_config?: LiteLLMParamsComplexityRouterConfigMap | null;
|
|
653
|
+
complexity_router_default_model?: string | null;
|
|
654
|
+
configurable_clientside_auth_params?: LiteLLMParamsConfigurableClientsideAuthParamsList | null;
|
|
655
|
+
custom_llm_provider?: string | null;
|
|
656
|
+
default_api_key_rpm_limit?: number | null;
|
|
657
|
+
default_api_key_tpm_limit?: number | null;
|
|
658
|
+
gcs_bucket_name?: string | null;
|
|
659
|
+
google_maps_grounding_cost_per_query?: number | null;
|
|
660
|
+
input_cost_per_audio_per_second?: number | null;
|
|
661
|
+
input_cost_per_audio_per_second_above_128k_tokens?: number | null;
|
|
662
|
+
input_cost_per_audio_token?: number | null;
|
|
663
|
+
input_cost_per_character?: number | null;
|
|
664
|
+
input_cost_per_character_above_128k_tokens?: number | null;
|
|
665
|
+
input_cost_per_image?: number | null;
|
|
666
|
+
input_cost_per_image_above_128k_tokens?: number | null;
|
|
667
|
+
input_cost_per_image_token?: number | null;
|
|
668
|
+
input_cost_per_pixel?: number | null;
|
|
669
|
+
input_cost_per_query?: number | null;
|
|
670
|
+
input_cost_per_second?: number | null;
|
|
671
|
+
input_cost_per_token?: number | null;
|
|
672
|
+
input_cost_per_token_above_128k_tokens?: number | null;
|
|
673
|
+
input_cost_per_token_above_200k_tokens?: number | null;
|
|
674
|
+
input_cost_per_token_above_200k_tokens_priority?: number | null;
|
|
675
|
+
input_cost_per_token_above_272k_tokens?: number | null;
|
|
676
|
+
input_cost_per_token_above_272k_tokens_flex?: number | null;
|
|
677
|
+
input_cost_per_token_above_272k_tokens_priority?: number | null;
|
|
678
|
+
input_cost_per_token_above_512k_tokens?: number | null;
|
|
679
|
+
input_cost_per_token_batches?: number | null;
|
|
680
|
+
input_cost_per_token_cache_hit?: number | null;
|
|
681
|
+
input_cost_per_token_flex?: number | null;
|
|
682
|
+
input_cost_per_token_priority?: number | null;
|
|
683
|
+
input_cost_per_token_ultrafast?: number | null;
|
|
684
|
+
input_cost_per_video_per_second?: number | null;
|
|
685
|
+
input_cost_per_video_per_second_above_128k_tokens?: number | null;
|
|
686
|
+
input_cost_per_video_per_second_above_15s_interval?: number | null;
|
|
687
|
+
input_cost_per_video_per_second_above_8s_interval?: number | null;
|
|
688
|
+
input_cost_per_video_token?: number | null;
|
|
689
|
+
itpm?: number | null;
|
|
690
|
+
keepalive_seconds?: number | null;
|
|
691
|
+
litellm_credential_name?: string | null;
|
|
692
|
+
litellm_trace_id?: string | null;
|
|
693
|
+
max_budget?: number | null;
|
|
694
|
+
max_file_size_mb?: number | null;
|
|
695
|
+
max_retries?: number | null;
|
|
696
|
+
merge_reasoning_content_in_choices?: boolean | null;
|
|
697
|
+
milvus_db_name?: string | null;
|
|
698
|
+
milvus_partition_names?: LiteLLMParamsMilvusPartitionNamesList | null;
|
|
699
|
+
milvus_text_field?: string | null;
|
|
700
|
+
mock_response?: LiteLLMParamsMockResponse | null;
|
|
701
|
+
model: string;
|
|
702
|
+
model_info?: LiteLLMParamsModelInfoMap | null;
|
|
703
|
+
ocr_cost_per_credit?: number | null;
|
|
704
|
+
ocr_cost_per_page?: number | null;
|
|
705
|
+
organization?: string | null;
|
|
706
|
+
otpm?: number | null;
|
|
707
|
+
output_cost_per_audio_per_second?: number | null;
|
|
708
|
+
output_cost_per_audio_token?: number | null;
|
|
709
|
+
output_cost_per_character?: number | null;
|
|
710
|
+
output_cost_per_character_above_128k_tokens?: number | null;
|
|
711
|
+
output_cost_per_image?: number | null;
|
|
712
|
+
output_cost_per_image_token?: number | null;
|
|
713
|
+
output_cost_per_pixel?: number | null;
|
|
714
|
+
output_cost_per_reasoning_token?: number | null;
|
|
715
|
+
output_cost_per_reasoning_token_flex?: number | null;
|
|
716
|
+
output_cost_per_reasoning_token_priority?: number | null;
|
|
717
|
+
output_cost_per_second?: number | null;
|
|
718
|
+
output_cost_per_second_1080p?: number | null;
|
|
719
|
+
output_cost_per_second_480p?: number | null;
|
|
720
|
+
output_cost_per_second_4k?: number | null;
|
|
721
|
+
output_cost_per_token?: number | null;
|
|
722
|
+
output_cost_per_token_above_128k_tokens?: number | null;
|
|
723
|
+
output_cost_per_token_above_200k_tokens?: number | null;
|
|
724
|
+
output_cost_per_token_above_200k_tokens_priority?: number | null;
|
|
725
|
+
output_cost_per_token_above_272k_tokens?: number | null;
|
|
726
|
+
output_cost_per_token_above_272k_tokens_flex?: number | null;
|
|
727
|
+
output_cost_per_token_above_272k_tokens_priority?: number | null;
|
|
728
|
+
output_cost_per_token_above_512k_tokens?: number | null;
|
|
729
|
+
output_cost_per_token_batches?: number | null;
|
|
730
|
+
output_cost_per_token_flex?: number | null;
|
|
731
|
+
output_cost_per_token_priority?: number | null;
|
|
732
|
+
output_cost_per_token_ultrafast?: number | null;
|
|
733
|
+
output_cost_per_video_per_second?: number | null;
|
|
734
|
+
output_cost_per_video_token?: number | null;
|
|
735
|
+
output_vector_size?: number | null;
|
|
736
|
+
quality_router_config?: LiteLLMParamsQualityRouterConfigMap | null;
|
|
737
|
+
quality_router_default_model?: string | null;
|
|
738
|
+
region_name?: string | null;
|
|
739
|
+
regional_endpoint_uplift_multiplier?: number | null;
|
|
740
|
+
regional_processing_uplift_multiplier_eu?: number | null;
|
|
741
|
+
regional_processing_uplift_multiplier_us?: number | null;
|
|
742
|
+
rpm?: number | null;
|
|
743
|
+
s3_bucket_name?: string | null;
|
|
744
|
+
s3_encryption_key_id?: string | null;
|
|
745
|
+
s3_output_bucket_name?: string | null;
|
|
746
|
+
s3_region_name?: string | null;
|
|
747
|
+
search_context_cost_per_query?: LiteLLMParamsSearchContextCostPerQueryMap | null;
|
|
748
|
+
stream_timeout?: LiteLLMParamsStreamTimeout | null;
|
|
749
|
+
tag_regex?: LiteLLMParamsTagRegexList | null;
|
|
750
|
+
tags?: LiteLLMParamsTagsList | null;
|
|
751
|
+
tiered_pricing?: LiteLLMParamsTieredPricingList | null;
|
|
752
|
+
timeout?: LiteLLMParamsTimeout | null;
|
|
753
|
+
tpm?: number | null;
|
|
754
|
+
use_chat_completions_api?: boolean | null;
|
|
755
|
+
use_in_pass_through?: boolean | null;
|
|
756
|
+
use_litellm_proxy?: boolean | null;
|
|
757
|
+
/** Use stored xAI OAuth credentials when no xAI API key is configured. */
|
|
758
|
+
use_xai_oauth?: boolean | null;
|
|
759
|
+
valkey_embedding_field?: string | null;
|
|
760
|
+
valkey_host?: string | null;
|
|
761
|
+
valkey_password?: string | null;
|
|
762
|
+
valkey_port?: number | null;
|
|
763
|
+
valkey_ssl?: boolean | null;
|
|
764
|
+
valkey_text_field?: string | null;
|
|
765
|
+
vector_store_id?: string | null;
|
|
766
|
+
vertex_credentials?: LiteLLMParamsVertexCredentials | null;
|
|
767
|
+
vertex_location?: string | null;
|
|
768
|
+
vertex_project?: string | null;
|
|
769
|
+
watsonx_region_name?: string | null;
|
|
770
|
+
}
|
|
771
|
+
export const LiteLLMParams = /*@__PURE__*/ S.suspend(() =>
|
|
772
|
+
S.Struct({
|
|
773
|
+
adaptive_router_config: S.optional(
|
|
774
|
+
S.NullOr(LiteLLMParamsAdaptiveRouterConfigMap),
|
|
775
|
+
),
|
|
776
|
+
adaptive_router_default_model: S.optional(S.NullOr(S.String)),
|
|
777
|
+
allow_client_keepalive_override: S.optional(S.NullOr(S.Boolean)),
|
|
778
|
+
annotation_cost_per_page: S.optional(S.NullOr(S.Number)),
|
|
779
|
+
api_base: S.optional(S.NullOr(S.String)),
|
|
780
|
+
api_key: S.optional(S.NullOr(S.String)),
|
|
781
|
+
api_version: S.optional(S.NullOr(S.String)),
|
|
782
|
+
auto_router_config: S.optional(S.NullOr(S.String)),
|
|
783
|
+
auto_router_config_path: S.optional(S.NullOr(S.String)),
|
|
784
|
+
auto_router_default_model: S.optional(S.NullOr(S.String)),
|
|
785
|
+
auto_router_embedding_model: S.optional(S.NullOr(S.String)),
|
|
786
|
+
auto_router_max_input_chars: S.optional(S.NullOr(S.Number)),
|
|
787
|
+
aws_access_key_id: S.optional(S.NullOr(S.String)),
|
|
788
|
+
aws_batch_role_arn: S.optional(S.NullOr(S.String)),
|
|
789
|
+
aws_bedrock_project_id: S.optional(S.NullOr(S.String)),
|
|
790
|
+
aws_bedrock_runtime_endpoint: S.optional(S.NullOr(S.String)),
|
|
791
|
+
aws_external_id: S.optional(S.NullOr(S.String)),
|
|
792
|
+
aws_profile_name: S.optional(S.NullOr(S.String)),
|
|
793
|
+
aws_region_name: S.optional(S.NullOr(S.String)),
|
|
794
|
+
aws_role_name: S.optional(S.NullOr(S.String)),
|
|
795
|
+
aws_secret_access_key: S.optional(S.NullOr(S.String)),
|
|
796
|
+
aws_session_name: S.optional(S.NullOr(S.String)),
|
|
797
|
+
aws_session_token: S.optional(S.NullOr(S.String)),
|
|
798
|
+
aws_sts_endpoint: S.optional(S.NullOr(S.String)),
|
|
799
|
+
aws_web_identity_token: S.optional(S.NullOr(S.String)),
|
|
800
|
+
azure_ad_token: S.optional(S.NullOr(S.String)),
|
|
801
|
+
bedrock_tags: S.optional(S.NullOr(LiteLLMParamsBedrockTagsList)),
|
|
802
|
+
budget_duration: S.optional(S.NullOr(S.String)),
|
|
803
|
+
cache_creation_input_audio_token_cost: S.optional(S.NullOr(S.Number)),
|
|
804
|
+
cache_creation_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
805
|
+
cache_creation_input_token_cost_above_1hr: S.optional(S.NullOr(S.Number)),
|
|
806
|
+
cache_creation_input_token_cost_above_200k_tokens: S.optional(
|
|
807
|
+
S.NullOr(S.Number),
|
|
808
|
+
),
|
|
809
|
+
cache_creation_input_token_cost_above_272k_tokens: S.optional(
|
|
810
|
+
S.NullOr(S.Number),
|
|
811
|
+
),
|
|
812
|
+
cache_creation_input_token_cost_above_272k_tokens_flex: S.optional(
|
|
813
|
+
S.NullOr(S.Number),
|
|
814
|
+
),
|
|
815
|
+
cache_creation_input_token_cost_above_272k_tokens_priority: S.optional(
|
|
816
|
+
S.NullOr(S.Number),
|
|
817
|
+
),
|
|
818
|
+
cache_creation_input_token_cost_flex: S.optional(S.NullOr(S.Number)),
|
|
819
|
+
cache_creation_input_token_cost_priority: S.optional(S.NullOr(S.Number)),
|
|
820
|
+
cache_creation_input_token_cost_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
821
|
+
cache_read_input_audio_token_cost: S.optional(S.NullOr(S.Number)),
|
|
822
|
+
cache_read_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
823
|
+
cache_read_input_token_cost_above_200k_tokens: S.optional(
|
|
824
|
+
S.NullOr(S.Number),
|
|
825
|
+
),
|
|
826
|
+
cache_read_input_token_cost_above_200k_tokens_priority: S.optional(
|
|
827
|
+
S.NullOr(S.Number),
|
|
828
|
+
),
|
|
829
|
+
cache_read_input_token_cost_above_272k_tokens: S.optional(
|
|
830
|
+
S.NullOr(S.Number),
|
|
831
|
+
),
|
|
832
|
+
cache_read_input_token_cost_above_272k_tokens_flex: S.optional(
|
|
833
|
+
S.NullOr(S.Number),
|
|
834
|
+
),
|
|
835
|
+
cache_read_input_token_cost_above_272k_tokens_priority: S.optional(
|
|
836
|
+
S.NullOr(S.Number),
|
|
837
|
+
),
|
|
838
|
+
cache_read_input_token_cost_above_512k_tokens: S.optional(
|
|
839
|
+
S.NullOr(S.Number),
|
|
840
|
+
),
|
|
841
|
+
cache_read_input_token_cost_flex: S.optional(S.NullOr(S.Number)),
|
|
842
|
+
cache_read_input_token_cost_priority: S.optional(S.NullOr(S.Number)),
|
|
843
|
+
cache_read_input_token_cost_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
844
|
+
citation_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
845
|
+
complexity_router_config: S.optional(
|
|
846
|
+
S.NullOr(LiteLLMParamsComplexityRouterConfigMap),
|
|
847
|
+
),
|
|
848
|
+
complexity_router_default_model: S.optional(S.NullOr(S.String)),
|
|
849
|
+
configurable_clientside_auth_params: S.optional(
|
|
850
|
+
S.NullOr(LiteLLMParamsConfigurableClientsideAuthParamsList),
|
|
851
|
+
),
|
|
852
|
+
custom_llm_provider: S.optional(S.NullOr(S.String)),
|
|
853
|
+
default_api_key_rpm_limit: S.optional(S.NullOr(S.Number)),
|
|
854
|
+
default_api_key_tpm_limit: S.optional(S.NullOr(S.Number)),
|
|
855
|
+
gcs_bucket_name: S.optional(S.NullOr(S.String)),
|
|
856
|
+
google_maps_grounding_cost_per_query: S.optional(S.NullOr(S.Number)),
|
|
857
|
+
input_cost_per_audio_per_second: S.optional(S.NullOr(S.Number)),
|
|
858
|
+
input_cost_per_audio_per_second_above_128k_tokens: S.optional(
|
|
859
|
+
S.NullOr(S.Number),
|
|
860
|
+
),
|
|
861
|
+
input_cost_per_audio_token: S.optional(S.NullOr(S.Number)),
|
|
862
|
+
input_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
863
|
+
input_cost_per_character_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
864
|
+
input_cost_per_image: S.optional(S.NullOr(S.Number)),
|
|
865
|
+
input_cost_per_image_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
866
|
+
input_cost_per_image_token: S.optional(S.NullOr(S.Number)),
|
|
867
|
+
input_cost_per_pixel: S.optional(S.NullOr(S.Number)),
|
|
868
|
+
input_cost_per_query: S.optional(S.NullOr(S.Number)),
|
|
869
|
+
input_cost_per_second: S.optional(S.NullOr(S.Number)),
|
|
870
|
+
input_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
871
|
+
input_cost_per_token_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
872
|
+
input_cost_per_token_above_200k_tokens: S.optional(S.NullOr(S.Number)),
|
|
873
|
+
input_cost_per_token_above_200k_tokens_priority: S.optional(
|
|
874
|
+
S.NullOr(S.Number),
|
|
875
|
+
),
|
|
876
|
+
input_cost_per_token_above_272k_tokens: S.optional(S.NullOr(S.Number)),
|
|
877
|
+
input_cost_per_token_above_272k_tokens_flex: S.optional(S.NullOr(S.Number)),
|
|
878
|
+
input_cost_per_token_above_272k_tokens_priority: S.optional(
|
|
879
|
+
S.NullOr(S.Number),
|
|
880
|
+
),
|
|
881
|
+
input_cost_per_token_above_512k_tokens: S.optional(S.NullOr(S.Number)),
|
|
882
|
+
input_cost_per_token_batches: S.optional(S.NullOr(S.Number)),
|
|
883
|
+
input_cost_per_token_cache_hit: S.optional(S.NullOr(S.Number)),
|
|
884
|
+
input_cost_per_token_flex: S.optional(S.NullOr(S.Number)),
|
|
885
|
+
input_cost_per_token_priority: S.optional(S.NullOr(S.Number)),
|
|
886
|
+
input_cost_per_token_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
887
|
+
input_cost_per_video_per_second: S.optional(S.NullOr(S.Number)),
|
|
888
|
+
input_cost_per_video_per_second_above_128k_tokens: S.optional(
|
|
889
|
+
S.NullOr(S.Number),
|
|
890
|
+
),
|
|
891
|
+
input_cost_per_video_per_second_above_15s_interval: S.optional(
|
|
892
|
+
S.NullOr(S.Number),
|
|
893
|
+
),
|
|
894
|
+
input_cost_per_video_per_second_above_8s_interval: S.optional(
|
|
895
|
+
S.NullOr(S.Number),
|
|
896
|
+
),
|
|
897
|
+
input_cost_per_video_token: S.optional(S.NullOr(S.Number)),
|
|
898
|
+
itpm: S.optional(S.NullOr(S.Number)),
|
|
899
|
+
keepalive_seconds: S.optional(S.NullOr(S.Number)),
|
|
900
|
+
litellm_credential_name: S.optional(S.NullOr(S.String)),
|
|
901
|
+
litellm_trace_id: S.optional(S.NullOr(S.String)),
|
|
902
|
+
max_budget: S.optional(S.NullOr(S.Number)),
|
|
903
|
+
max_file_size_mb: S.optional(S.NullOr(S.Number)),
|
|
904
|
+
max_retries: S.optional(S.NullOr(S.Number)),
|
|
905
|
+
merge_reasoning_content_in_choices: S.optional(S.NullOr(S.Boolean)),
|
|
906
|
+
milvus_db_name: S.optional(S.NullOr(S.String)),
|
|
907
|
+
milvus_partition_names: S.optional(
|
|
908
|
+
S.NullOr(LiteLLMParamsMilvusPartitionNamesList),
|
|
909
|
+
),
|
|
910
|
+
milvus_text_field: S.optional(S.NullOr(S.String)),
|
|
911
|
+
mock_response: S.optional(S.NullOr(LiteLLMParamsMockResponse)),
|
|
912
|
+
model: S.String,
|
|
913
|
+
model_info: S.optional(S.NullOr(LiteLLMParamsModelInfoMap)),
|
|
914
|
+
ocr_cost_per_credit: S.optional(S.NullOr(S.Number)),
|
|
915
|
+
ocr_cost_per_page: S.optional(S.NullOr(S.Number)),
|
|
916
|
+
organization: S.optional(S.NullOr(S.String)),
|
|
917
|
+
otpm: S.optional(S.NullOr(S.Number)),
|
|
918
|
+
output_cost_per_audio_per_second: S.optional(S.NullOr(S.Number)),
|
|
919
|
+
output_cost_per_audio_token: S.optional(S.NullOr(S.Number)),
|
|
920
|
+
output_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
921
|
+
output_cost_per_character_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
922
|
+
output_cost_per_image: S.optional(S.NullOr(S.Number)),
|
|
923
|
+
output_cost_per_image_token: S.optional(S.NullOr(S.Number)),
|
|
924
|
+
output_cost_per_pixel: S.optional(S.NullOr(S.Number)),
|
|
925
|
+
output_cost_per_reasoning_token: S.optional(S.NullOr(S.Number)),
|
|
926
|
+
output_cost_per_reasoning_token_flex: S.optional(S.NullOr(S.Number)),
|
|
927
|
+
output_cost_per_reasoning_token_priority: S.optional(S.NullOr(S.Number)),
|
|
928
|
+
output_cost_per_second: S.optional(S.NullOr(S.Number)),
|
|
929
|
+
output_cost_per_second_1080p: S.optional(S.NullOr(S.Number)),
|
|
930
|
+
output_cost_per_second_480p: S.optional(S.NullOr(S.Number)),
|
|
931
|
+
output_cost_per_second_4k: S.optional(S.NullOr(S.Number)),
|
|
932
|
+
output_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
933
|
+
output_cost_per_token_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
934
|
+
output_cost_per_token_above_200k_tokens: S.optional(S.NullOr(S.Number)),
|
|
935
|
+
output_cost_per_token_above_200k_tokens_priority: S.optional(
|
|
936
|
+
S.NullOr(S.Number),
|
|
937
|
+
),
|
|
938
|
+
output_cost_per_token_above_272k_tokens: S.optional(S.NullOr(S.Number)),
|
|
939
|
+
output_cost_per_token_above_272k_tokens_flex: S.optional(
|
|
940
|
+
S.NullOr(S.Number),
|
|
941
|
+
),
|
|
942
|
+
output_cost_per_token_above_272k_tokens_priority: S.optional(
|
|
943
|
+
S.NullOr(S.Number),
|
|
944
|
+
),
|
|
945
|
+
output_cost_per_token_above_512k_tokens: S.optional(S.NullOr(S.Number)),
|
|
946
|
+
output_cost_per_token_batches: S.optional(S.NullOr(S.Number)),
|
|
947
|
+
output_cost_per_token_flex: S.optional(S.NullOr(S.Number)),
|
|
948
|
+
output_cost_per_token_priority: S.optional(S.NullOr(S.Number)),
|
|
949
|
+
output_cost_per_token_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
950
|
+
output_cost_per_video_per_second: S.optional(S.NullOr(S.Number)),
|
|
951
|
+
output_cost_per_video_token: S.optional(S.NullOr(S.Number)),
|
|
952
|
+
output_vector_size: S.optional(S.NullOr(S.Number)),
|
|
953
|
+
quality_router_config: S.optional(
|
|
954
|
+
S.NullOr(LiteLLMParamsQualityRouterConfigMap),
|
|
955
|
+
),
|
|
956
|
+
quality_router_default_model: S.optional(S.NullOr(S.String)),
|
|
957
|
+
region_name: S.optional(S.NullOr(S.String)),
|
|
958
|
+
regional_endpoint_uplift_multiplier: S.optional(S.NullOr(S.Number)),
|
|
959
|
+
regional_processing_uplift_multiplier_eu: S.optional(S.NullOr(S.Number)),
|
|
960
|
+
regional_processing_uplift_multiplier_us: S.optional(S.NullOr(S.Number)),
|
|
961
|
+
rpm: S.optional(S.NullOr(S.Number)),
|
|
962
|
+
s3_bucket_name: S.optional(S.NullOr(S.String)),
|
|
963
|
+
s3_encryption_key_id: S.optional(S.NullOr(S.String)),
|
|
964
|
+
s3_output_bucket_name: S.optional(S.NullOr(S.String)),
|
|
965
|
+
s3_region_name: S.optional(S.NullOr(S.String)),
|
|
966
|
+
search_context_cost_per_query: S.optional(
|
|
967
|
+
S.NullOr(LiteLLMParamsSearchContextCostPerQueryMap),
|
|
968
|
+
),
|
|
969
|
+
stream_timeout: S.optional(S.NullOr(LiteLLMParamsStreamTimeout)),
|
|
970
|
+
tag_regex: S.optional(S.NullOr(LiteLLMParamsTagRegexList)),
|
|
971
|
+
tags: S.optional(S.NullOr(LiteLLMParamsTagsList)),
|
|
972
|
+
tiered_pricing: S.optional(S.NullOr(LiteLLMParamsTieredPricingList)),
|
|
973
|
+
timeout: S.optional(S.NullOr(LiteLLMParamsTimeout)),
|
|
974
|
+
tpm: S.optional(S.NullOr(S.Number)),
|
|
975
|
+
use_chat_completions_api: S.optional(S.NullOr(S.Boolean)),
|
|
976
|
+
use_in_pass_through: S.optional(S.NullOr(S.Boolean)),
|
|
977
|
+
use_litellm_proxy: S.optional(S.NullOr(S.Boolean)),
|
|
978
|
+
use_xai_oauth: S.optional(S.NullOr(S.Boolean)),
|
|
979
|
+
valkey_embedding_field: S.optional(S.NullOr(S.String)),
|
|
980
|
+
valkey_host: S.optional(S.NullOr(S.String)),
|
|
981
|
+
valkey_password: S.optional(S.NullOr(S.String)),
|
|
982
|
+
valkey_port: S.optional(S.NullOr(S.Number)),
|
|
983
|
+
valkey_ssl: S.optional(S.NullOr(S.Boolean)),
|
|
984
|
+
valkey_text_field: S.optional(S.NullOr(S.String)),
|
|
985
|
+
vector_store_id: S.optional(S.NullOr(S.String)),
|
|
986
|
+
vertex_credentials: S.optional(S.NullOr(LiteLLMParamsVertexCredentials)),
|
|
987
|
+
vertex_location: S.optional(S.NullOr(S.String)),
|
|
988
|
+
vertex_project: S.optional(S.NullOr(S.String)),
|
|
989
|
+
watsonx_region_name: S.optional(S.NullOr(S.String)),
|
|
990
|
+
}),
|
|
991
|
+
).annotate({ identifier: "LiteLLMParams" }) as any as S.Schema<LiteLLMParams>;
|
|
992
|
+
|
|
993
|
+
export type LitellmTypesRouterModelInfoTier = "free" | "paid";
|
|
994
|
+
export const LitellmTypesRouterModelInfoTier = S.String;
|
|
995
|
+
|
|
996
|
+
export type LitellmTypesRouterModelInfoTieredPricingItemMap = {
|
|
997
|
+
[key: string]: unknown | undefined;
|
|
998
|
+
};
|
|
999
|
+
export const LitellmTypesRouterModelInfoTieredPricingItemMap =
|
|
1000
|
+
/*@__PURE__*/ S.Record(
|
|
1001
|
+
S.String,
|
|
1002
|
+
S.Unknown,
|
|
1003
|
+
) as any as S.Schema<LitellmTypesRouterModelInfoTieredPricingItemMap>;
|
|
1004
|
+
|
|
1005
|
+
export type LitellmTypesRouterModelInfoTieredPricingList =
|
|
1006
|
+
Array<LitellmTypesRouterModelInfoTieredPricingItemMap>;
|
|
1007
|
+
export const LitellmTypesRouterModelInfoTieredPricingList =
|
|
1008
|
+
/*@__PURE__*/ S.Array(
|
|
1009
|
+
LitellmTypesRouterModelInfoTieredPricingItemMap,
|
|
1010
|
+
) as any as S.Schema<LitellmTypesRouterModelInfoTieredPricingList>;
|
|
1011
|
+
|
|
1012
|
+
export interface LitellmTypesRouterModelInfo {
|
|
1013
|
+
allow_fail_open?: boolean | null;
|
|
1014
|
+
base_model?: string | null;
|
|
1015
|
+
blocked?: boolean | null;
|
|
1016
|
+
cache_creation_input_token_cost?: number | null;
|
|
1017
|
+
cache_read_input_token_cost?: number | null;
|
|
1018
|
+
cost_per_ptu_per_hour?: number | null;
|
|
1019
|
+
created_at?: string | null;
|
|
1020
|
+
created_by?: string | null;
|
|
1021
|
+
db_model?: boolean;
|
|
1022
|
+
enable_tag_filtering?: boolean | null;
|
|
1023
|
+
id: string | null;
|
|
1024
|
+
input_cost_per_character?: number | null;
|
|
1025
|
+
input_cost_per_token?: number | null;
|
|
1026
|
+
output_cost_per_character?: number | null;
|
|
1027
|
+
output_cost_per_token?: number | null;
|
|
1028
|
+
ptu_count?: number | null;
|
|
1029
|
+
ptu_effective_from?: string | null;
|
|
1030
|
+
ptu_effective_to?: string | null;
|
|
1031
|
+
team_id?: string | null;
|
|
1032
|
+
team_public_model_name?: string | null;
|
|
1033
|
+
tier?: LitellmTypesRouterModelInfoTier | (string & {}) | null;
|
|
1034
|
+
tiered_pricing?: LitellmTypesRouterModelInfoTieredPricingList | null;
|
|
1035
|
+
updated_at?: string | null;
|
|
1036
|
+
updated_by?: string | null;
|
|
1037
|
+
}
|
|
1038
|
+
export const LitellmTypesRouterModelInfo = /*@__PURE__*/ S.suspend(() =>
|
|
1039
|
+
S.Struct({
|
|
1040
|
+
allow_fail_open: S.optional(S.NullOr(S.Boolean)),
|
|
1041
|
+
base_model: S.optional(S.NullOr(S.String)),
|
|
1042
|
+
blocked: S.optional(S.NullOr(S.Boolean)),
|
|
1043
|
+
cache_creation_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
1044
|
+
cache_read_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
1045
|
+
cost_per_ptu_per_hour: S.optional(S.NullOr(S.Number)),
|
|
1046
|
+
created_at: S.optional(S.NullOr(S.String)),
|
|
1047
|
+
created_by: S.optional(S.NullOr(S.String)),
|
|
1048
|
+
db_model: S.optional(S.Boolean),
|
|
1049
|
+
enable_tag_filtering: S.optional(S.NullOr(S.Boolean)),
|
|
1050
|
+
id: S.NullOr(S.String),
|
|
1051
|
+
input_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
1052
|
+
input_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
1053
|
+
output_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
1054
|
+
output_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
1055
|
+
ptu_count: S.optional(S.NullOr(S.Number)),
|
|
1056
|
+
ptu_effective_from: S.optional(S.NullOr(S.String)),
|
|
1057
|
+
ptu_effective_to: S.optional(S.NullOr(S.String)),
|
|
1058
|
+
team_id: S.optional(S.NullOr(S.String)),
|
|
1059
|
+
team_public_model_name: S.optional(S.NullOr(S.String)),
|
|
1060
|
+
tier: S.optional(S.NullOr(LitellmTypesRouterModelInfoTier)),
|
|
1061
|
+
tiered_pricing: S.optional(
|
|
1062
|
+
S.NullOr(LitellmTypesRouterModelInfoTieredPricingList),
|
|
1063
|
+
),
|
|
1064
|
+
updated_at: S.optional(S.NullOr(S.String)),
|
|
1065
|
+
updated_by: S.optional(S.NullOr(S.String)),
|
|
1066
|
+
}),
|
|
1067
|
+
).annotate({
|
|
1068
|
+
identifier: "LitellmTypesRouterModelInfo",
|
|
1069
|
+
}) as any as S.Schema<LitellmTypesRouterModelInfo>;
|
|
1070
|
+
|
|
1071
|
+
export interface AddNewModelModelNewPostRequest {
|
|
1072
|
+
litellm_params: LiteLLMParams;
|
|
1073
|
+
model_info: LitellmTypesRouterModelInfo;
|
|
1074
|
+
model_name: string;
|
|
1075
|
+
}
|
|
1076
|
+
export const AddNewModelModelNewPostRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1077
|
+
S.Struct({
|
|
1078
|
+
litellm_params: LiteLLMParams,
|
|
1079
|
+
model_info: LitellmTypesRouterModelInfo,
|
|
1080
|
+
model_name: S.String,
|
|
1081
|
+
}).pipe(T.Http({ method: "POST", uri: "/model/new", code: 200 })),
|
|
1082
|
+
).annotate({
|
|
1083
|
+
identifier: "AddNewModelModelNewPostRequest",
|
|
1084
|
+
}) as any as S.Schema<AddNewModelModelNewPostRequest>;
|
|
1085
|
+
|
|
1086
|
+
export interface AddNewModelModelNewPostResponse {
|
|
1087
|
+
body: unknown;
|
|
1088
|
+
}
|
|
1089
|
+
export const AddNewModelModelNewPostResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1090
|
+
S.Struct({
|
|
1091
|
+
body: S.Unknown,
|
|
1092
|
+
}),
|
|
1093
|
+
).annotate({
|
|
1094
|
+
identifier: "AddNewModelModelNewPostResponse",
|
|
1095
|
+
}) as any as S.Schema<AddNewModelModelNewPostResponse>;
|
|
1096
|
+
|
|
1097
|
+
export interface CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest {}
|
|
1098
|
+
export const CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest =
|
|
1099
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1100
|
+
S.Struct({}).pipe(
|
|
1101
|
+
T.Http({
|
|
1102
|
+
method: "DELETE",
|
|
1103
|
+
uri: "/schedule/anthropic_beta_headers_reload",
|
|
1104
|
+
code: 200,
|
|
1105
|
+
}),
|
|
1106
|
+
),
|
|
1107
|
+
).annotate({
|
|
1108
|
+
identifier:
|
|
1109
|
+
"CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest",
|
|
1110
|
+
}) as any as S.Schema<CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest>;
|
|
1111
|
+
|
|
1112
|
+
export interface CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse {
|
|
1113
|
+
body: unknown;
|
|
1114
|
+
}
|
|
1115
|
+
export const CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse =
|
|
1116
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1117
|
+
S.Struct({
|
|
1118
|
+
body: S.Unknown,
|
|
1119
|
+
}),
|
|
1120
|
+
).annotate({
|
|
1121
|
+
identifier:
|
|
1122
|
+
"CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse",
|
|
1123
|
+
}) as any as S.Schema<CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse>;
|
|
1124
|
+
|
|
1125
|
+
export interface CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest {}
|
|
1126
|
+
export const CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest =
|
|
1127
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1128
|
+
S.Struct({}).pipe(
|
|
1129
|
+
T.Http({
|
|
1130
|
+
method: "DELETE",
|
|
1131
|
+
uri: "/schedule/model_cost_map_reload",
|
|
1132
|
+
code: 200,
|
|
1133
|
+
}),
|
|
1134
|
+
),
|
|
1135
|
+
).annotate({
|
|
1136
|
+
identifier:
|
|
1137
|
+
"CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest",
|
|
1138
|
+
}) as any as S.Schema<CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest>;
|
|
1139
|
+
|
|
1140
|
+
export interface CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse {
|
|
1141
|
+
body: unknown;
|
|
1142
|
+
}
|
|
1143
|
+
export const CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse =
|
|
1144
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1145
|
+
S.Struct({
|
|
1146
|
+
body: S.Unknown,
|
|
1147
|
+
}),
|
|
1148
|
+
).annotate({
|
|
1149
|
+
identifier:
|
|
1150
|
+
"CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse",
|
|
1151
|
+
}) as any as S.Schema<CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse>;
|
|
1152
|
+
|
|
1153
|
+
export type CreateModelGroupAccessGroupNewPostRequestModelIdsList =
|
|
1154
|
+
Array<string>;
|
|
1155
|
+
export const CreateModelGroupAccessGroupNewPostRequestModelIdsList =
|
|
1156
|
+
/*@__PURE__*/ S.Array(
|
|
1157
|
+
S.String,
|
|
1158
|
+
) as any as S.Schema<CreateModelGroupAccessGroupNewPostRequestModelIdsList>;
|
|
1159
|
+
|
|
1160
|
+
export type CreateModelGroupAccessGroupNewPostRequestModelNamesList =
|
|
1161
|
+
Array<string>;
|
|
1162
|
+
export const CreateModelGroupAccessGroupNewPostRequestModelNamesList =
|
|
1163
|
+
/*@__PURE__*/ S.Array(
|
|
1164
|
+
S.String,
|
|
1165
|
+
) as any as S.Schema<CreateModelGroupAccessGroupNewPostRequestModelNamesList>;
|
|
1166
|
+
|
|
1167
|
+
export interface CreateModelGroupAccessGroupNewPostRequest {
|
|
1168
|
+
access_group: string;
|
|
1169
|
+
model_ids?: CreateModelGroupAccessGroupNewPostRequestModelIdsList | null;
|
|
1170
|
+
model_names?: CreateModelGroupAccessGroupNewPostRequestModelNamesList | null;
|
|
1171
|
+
}
|
|
1172
|
+
export const CreateModelGroupAccessGroupNewPostRequest =
|
|
1173
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1174
|
+
S.Struct({
|
|
1175
|
+
access_group: S.String,
|
|
1176
|
+
model_ids: S.optional(
|
|
1177
|
+
S.NullOr(CreateModelGroupAccessGroupNewPostRequestModelIdsList),
|
|
1178
|
+
),
|
|
1179
|
+
model_names: S.optional(
|
|
1180
|
+
S.NullOr(CreateModelGroupAccessGroupNewPostRequestModelNamesList),
|
|
1181
|
+
),
|
|
1182
|
+
}).pipe(T.Http({ method: "POST", uri: "/access_group/new", code: 200 })),
|
|
1183
|
+
).annotate({
|
|
1184
|
+
identifier: "CreateModelGroupAccessGroupNewPostRequest",
|
|
1185
|
+
}) as any as S.Schema<CreateModelGroupAccessGroupNewPostRequest>;
|
|
1186
|
+
|
|
1187
|
+
export type NewModelGroupResponseModelIdsList = Array<string>;
|
|
1188
|
+
export const NewModelGroupResponseModelIdsList = /*@__PURE__*/ S.Array(
|
|
1189
|
+
S.String,
|
|
1190
|
+
) as any as S.Schema<NewModelGroupResponseModelIdsList>;
|
|
1191
|
+
|
|
1192
|
+
export type NewModelGroupResponseModelNamesList = Array<string>;
|
|
1193
|
+
export const NewModelGroupResponseModelNamesList = /*@__PURE__*/ S.Array(
|
|
1194
|
+
S.String,
|
|
1195
|
+
) as any as S.Schema<NewModelGroupResponseModelNamesList>;
|
|
1196
|
+
|
|
1197
|
+
export interface NewModelGroupResponse {
|
|
1198
|
+
access_group: string;
|
|
1199
|
+
model_ids?: NewModelGroupResponseModelIdsList | null;
|
|
1200
|
+
model_names?: NewModelGroupResponseModelNamesList | null;
|
|
1201
|
+
models_updated: number;
|
|
1202
|
+
}
|
|
1203
|
+
export const NewModelGroupResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1204
|
+
S.Struct({
|
|
1205
|
+
access_group: S.String,
|
|
1206
|
+
model_ids: S.optional(S.NullOr(NewModelGroupResponseModelIdsList)),
|
|
1207
|
+
model_names: S.optional(S.NullOr(NewModelGroupResponseModelNamesList)),
|
|
1208
|
+
models_updated: S.Number,
|
|
1209
|
+
}),
|
|
1210
|
+
).annotate({
|
|
1211
|
+
identifier: "NewModelGroupResponse",
|
|
1212
|
+
}) as any as S.Schema<NewModelGroupResponse>;
|
|
1213
|
+
|
|
1214
|
+
export interface DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest {
|
|
1215
|
+
access_group: string;
|
|
1216
|
+
}
|
|
1217
|
+
export const DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest =
|
|
1218
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1219
|
+
S.Struct({
|
|
1220
|
+
access_group: S.String.pipe(T.Label()),
|
|
1221
|
+
}).pipe(
|
|
1222
|
+
T.Http({
|
|
1223
|
+
method: "DELETE",
|
|
1224
|
+
uri: "/access_group/{access_group}/delete",
|
|
1225
|
+
code: 200,
|
|
1226
|
+
}),
|
|
1227
|
+
),
|
|
1228
|
+
).annotate({
|
|
1229
|
+
identifier: "DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest",
|
|
1230
|
+
}) as any as S.Schema<DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest>;
|
|
1231
|
+
|
|
1232
|
+
export interface DeleteModelGroupResponse {
|
|
1233
|
+
access_group: string;
|
|
1234
|
+
message: string;
|
|
1235
|
+
models_updated: number;
|
|
1236
|
+
}
|
|
1237
|
+
export const DeleteModelGroupResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1238
|
+
S.Struct({
|
|
1239
|
+
access_group: S.String,
|
|
1240
|
+
message: S.String,
|
|
1241
|
+
models_updated: S.Number,
|
|
1242
|
+
}),
|
|
1243
|
+
).annotate({
|
|
1244
|
+
identifier: "DeleteModelGroupResponse",
|
|
1245
|
+
}) as any as S.Schema<DeleteModelGroupResponse>;
|
|
1246
|
+
|
|
1247
|
+
export interface DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest {
|
|
1248
|
+
access_group: string;
|
|
1249
|
+
}
|
|
1250
|
+
export const DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest =
|
|
1251
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1252
|
+
S.Struct({
|
|
1253
|
+
access_group: S.String.pipe(T.Label()),
|
|
1254
|
+
}).pipe(
|
|
1255
|
+
T.Http({
|
|
1256
|
+
method: "DELETE",
|
|
1257
|
+
uri: "/access_group/{access_group}/budget",
|
|
1258
|
+
code: 200,
|
|
1259
|
+
}),
|
|
1260
|
+
),
|
|
1261
|
+
).annotate({
|
|
1262
|
+
identifier:
|
|
1263
|
+
"DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest",
|
|
1264
|
+
}) as any as S.Schema<DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest>;
|
|
1265
|
+
|
|
1266
|
+
export interface DeleteAccessGroupBudgetResponse {
|
|
1267
|
+
access_group: string;
|
|
1268
|
+
budget_deleted: boolean;
|
|
1269
|
+
message: string;
|
|
1270
|
+
}
|
|
1271
|
+
export const DeleteAccessGroupBudgetResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1272
|
+
S.Struct({
|
|
1273
|
+
access_group: S.String,
|
|
1274
|
+
budget_deleted: S.Boolean,
|
|
1275
|
+
message: S.String,
|
|
1276
|
+
}),
|
|
1277
|
+
).annotate({
|
|
1278
|
+
identifier: "DeleteAccessGroupBudgetResponse",
|
|
1279
|
+
}) as any as S.Schema<DeleteAccessGroupBudgetResponse>;
|
|
1280
|
+
|
|
1281
|
+
export interface DeleteModelModelDeletePostRequest {
|
|
1282
|
+
id: string;
|
|
1283
|
+
}
|
|
1284
|
+
export const DeleteModelModelDeletePostRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1285
|
+
S.Struct({
|
|
1286
|
+
id: S.String,
|
|
1287
|
+
}).pipe(T.Http({ method: "POST", uri: "/model/delete", code: 200 })),
|
|
1288
|
+
).annotate({
|
|
1289
|
+
identifier: "DeleteModelModelDeletePostRequest",
|
|
1290
|
+
}) as any as S.Schema<DeleteModelModelDeletePostRequest>;
|
|
1291
|
+
|
|
1292
|
+
export interface DeleteModelModelDeletePostResponse {
|
|
1293
|
+
body: unknown;
|
|
1294
|
+
}
|
|
1295
|
+
export const DeleteModelModelDeletePostResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1296
|
+
S.Struct({
|
|
1297
|
+
body: S.Unknown,
|
|
1298
|
+
}),
|
|
1299
|
+
).annotate({
|
|
1300
|
+
identifier: "DeleteModelModelDeletePostResponse",
|
|
1301
|
+
}) as any as S.Schema<DeleteModelModelDeletePostResponse>;
|
|
1302
|
+
|
|
1303
|
+
export interface GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest {
|
|
1304
|
+
access_group: string;
|
|
1305
|
+
}
|
|
1306
|
+
export const GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest =
|
|
1307
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1308
|
+
S.Struct({
|
|
1309
|
+
access_group: S.String.pipe(T.Label()),
|
|
1310
|
+
}).pipe(
|
|
1311
|
+
T.Http({
|
|
1312
|
+
method: "GET",
|
|
1313
|
+
uri: "/access_group/{access_group}/budget",
|
|
1314
|
+
code: 200,
|
|
1315
|
+
}),
|
|
1316
|
+
),
|
|
1317
|
+
).annotate({
|
|
1318
|
+
identifier: "GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest",
|
|
1319
|
+
}) as any as S.Schema<GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest>;
|
|
1320
|
+
|
|
1321
|
+
export interface AccessGroupBudget {
|
|
1322
|
+
budget_duration?: string | null;
|
|
1323
|
+
budget_id: string;
|
|
1324
|
+
budget_reset_at?: string | null;
|
|
1325
|
+
max_budget?: number | null;
|
|
1326
|
+
soft_budget?: number | null;
|
|
1327
|
+
}
|
|
1328
|
+
export const AccessGroupBudget = /*@__PURE__*/ S.suspend(() =>
|
|
1329
|
+
S.Struct({
|
|
1330
|
+
budget_duration: S.optional(S.NullOr(S.String)),
|
|
1331
|
+
budget_id: S.String,
|
|
1332
|
+
budget_reset_at: S.optional(S.NullOr(S.String)),
|
|
1333
|
+
max_budget: S.optional(S.NullOr(S.Number)),
|
|
1334
|
+
soft_budget: S.optional(S.NullOr(S.Number)),
|
|
1335
|
+
}),
|
|
1336
|
+
).annotate({
|
|
1337
|
+
identifier: "AccessGroupBudget",
|
|
1338
|
+
}) as any as S.Schema<AccessGroupBudget>;
|
|
1339
|
+
|
|
1340
|
+
export interface AccessGroupBudgetResponse {
|
|
1341
|
+
access_group: string;
|
|
1342
|
+
budget?: AccessGroupBudget | null;
|
|
1343
|
+
spend: number;
|
|
1344
|
+
}
|
|
1345
|
+
export const AccessGroupBudgetResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1346
|
+
S.Struct({
|
|
1347
|
+
access_group: S.String,
|
|
1348
|
+
budget: S.optional(S.NullOr(AccessGroupBudget)),
|
|
1349
|
+
spend: S.Number,
|
|
1350
|
+
}),
|
|
1351
|
+
).annotate({
|
|
1352
|
+
identifier: "AccessGroupBudgetResponse",
|
|
1353
|
+
}) as any as S.Schema<AccessGroupBudgetResponse>;
|
|
1354
|
+
|
|
1355
|
+
export interface GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest {
|
|
1356
|
+
access_group: string;
|
|
1357
|
+
}
|
|
1358
|
+
export const GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest =
|
|
1359
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1360
|
+
S.Struct({
|
|
1361
|
+
access_group: S.String.pipe(T.Label()),
|
|
1362
|
+
}).pipe(
|
|
1363
|
+
T.Http({
|
|
1364
|
+
method: "GET",
|
|
1365
|
+
uri: "/access_group/{access_group}/info",
|
|
1366
|
+
code: 200,
|
|
1367
|
+
}),
|
|
1368
|
+
),
|
|
1369
|
+
).annotate({
|
|
1370
|
+
identifier: "GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest",
|
|
1371
|
+
}) as any as S.Schema<GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest>;
|
|
1372
|
+
|
|
1373
|
+
export type AccessGroupInfoModelNamesList = Array<string>;
|
|
1374
|
+
export const AccessGroupInfoModelNamesList = /*@__PURE__*/ S.Array(
|
|
1375
|
+
S.String,
|
|
1376
|
+
) as any as S.Schema<AccessGroupInfoModelNamesList>;
|
|
1377
|
+
|
|
1378
|
+
export interface AccessGroupInfo {
|
|
1379
|
+
access_group: string;
|
|
1380
|
+
budget?: AccessGroupBudget | null;
|
|
1381
|
+
deployment_count: number;
|
|
1382
|
+
model_names: AccessGroupInfoModelNamesList;
|
|
1383
|
+
spend?: number | null;
|
|
1384
|
+
}
|
|
1385
|
+
export const AccessGroupInfo = /*@__PURE__*/ S.suspend(() =>
|
|
1386
|
+
S.Struct({
|
|
1387
|
+
access_group: S.String,
|
|
1388
|
+
budget: S.optional(S.NullOr(AccessGroupBudget)),
|
|
1389
|
+
deployment_count: S.Number,
|
|
1390
|
+
model_names: AccessGroupInfoModelNamesList,
|
|
1391
|
+
spend: S.optional(S.NullOr(S.Number)),
|
|
1392
|
+
}),
|
|
1393
|
+
).annotate({
|
|
1394
|
+
identifier: "AccessGroupInfo",
|
|
1395
|
+
}) as any as S.Schema<AccessGroupInfo>;
|
|
1396
|
+
|
|
1397
|
+
export interface GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest {}
|
|
1398
|
+
export const GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest =
|
|
1399
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1400
|
+
S.Struct({}).pipe(
|
|
1401
|
+
T.Http({
|
|
1402
|
+
method: "GET",
|
|
1403
|
+
uri: "/schedule/anthropic_beta_headers_reload/status",
|
|
1404
|
+
code: 200,
|
|
1405
|
+
}),
|
|
1406
|
+
),
|
|
1407
|
+
).annotate({
|
|
1408
|
+
identifier:
|
|
1409
|
+
"GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest",
|
|
1410
|
+
}) as any as S.Schema<GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest>;
|
|
1411
|
+
|
|
1412
|
+
export interface GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse {
|
|
1413
|
+
body: unknown;
|
|
1414
|
+
}
|
|
1415
|
+
export const GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse =
|
|
1416
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1417
|
+
S.Struct({
|
|
1418
|
+
body: S.Unknown,
|
|
1419
|
+
}),
|
|
1420
|
+
).annotate({
|
|
1421
|
+
identifier:
|
|
1422
|
+
"GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse",
|
|
1423
|
+
}) as any as S.Schema<GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse>;
|
|
1424
|
+
|
|
1425
|
+
/** Which calibration examples, and for BUSINESS which tier criteria, the built-in classifier rubric carries. */
|
|
1426
|
+
export type ClassificationRubric = "legacy" | "agentic" | "chat" | "business";
|
|
1427
|
+
export const ClassificationRubric = S.String;
|
|
1428
|
+
|
|
1429
|
+
export interface GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest {
|
|
1430
|
+
context_window_size?: number;
|
|
1431
|
+
tier_labels?: string;
|
|
1432
|
+
classification_rubric?: ClassificationRubric | (string & {});
|
|
1433
|
+
}
|
|
1434
|
+
export const GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest =
|
|
1435
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1436
|
+
S.Struct({
|
|
1437
|
+
context_window_size: S.optional(S.Number.pipe(T.Query())),
|
|
1438
|
+
tier_labels: S.optional(S.String.pipe(T.Query())),
|
|
1439
|
+
classification_rubric: S.optional(ClassificationRubric.pipe(T.Query())),
|
|
1440
|
+
}).pipe(
|
|
1441
|
+
T.Http({
|
|
1442
|
+
method: "GET",
|
|
1443
|
+
uri: "/auto_router/classifier/default_prompt",
|
|
1444
|
+
code: 200,
|
|
1445
|
+
}),
|
|
1446
|
+
),
|
|
1447
|
+
).annotate({
|
|
1448
|
+
identifier:
|
|
1449
|
+
"GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest",
|
|
1450
|
+
}) as any as S.Schema<GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest>;
|
|
1451
|
+
|
|
1452
|
+
/** The built-in system prompt an auto-router's LLM classifier uses when none is configured. Served so the dashboard's prompt editor prefills the rubric the proxy actually sends, rather than a copy in the frontend that drifts the moment the rubric is edited. */
|
|
1453
|
+
export interface AutoRouterClassifierDefaultPromptResponse {
|
|
1454
|
+
system_prompt: string;
|
|
1455
|
+
}
|
|
1456
|
+
export const AutoRouterClassifierDefaultPromptResponse =
|
|
1457
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1458
|
+
S.Struct({
|
|
1459
|
+
system_prompt: S.String,
|
|
1460
|
+
}),
|
|
1461
|
+
).annotate({
|
|
1462
|
+
identifier: "AutoRouterClassifierDefaultPromptResponse",
|
|
1463
|
+
}) as any as S.Schema<AutoRouterClassifierDefaultPromptResponse>;
|
|
1464
|
+
|
|
1465
|
+
export interface GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest {}
|
|
1466
|
+
export const GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest =
|
|
1467
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1468
|
+
S.Struct({}).pipe(
|
|
1469
|
+
T.Http({
|
|
1470
|
+
method: "GET",
|
|
1471
|
+
uri: "/schedule/model_cost_map_reload/status",
|
|
1472
|
+
code: 200,
|
|
1473
|
+
}),
|
|
1474
|
+
),
|
|
1475
|
+
).annotate({
|
|
1476
|
+
identifier:
|
|
1477
|
+
"GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest",
|
|
1478
|
+
}) as any as S.Schema<GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest>;
|
|
1479
|
+
|
|
1480
|
+
export interface GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse {
|
|
1481
|
+
body: unknown;
|
|
1482
|
+
}
|
|
1483
|
+
export const GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse =
|
|
1484
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1485
|
+
S.Struct({
|
|
1486
|
+
body: S.Unknown,
|
|
1487
|
+
}),
|
|
1488
|
+
).annotate({
|
|
1489
|
+
identifier:
|
|
1490
|
+
"GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse",
|
|
1491
|
+
}) as any as S.Schema<GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse>;
|
|
1492
|
+
|
|
1493
|
+
export interface GetModelCostMapSourceModelCostMapSourceGetRequest {}
|
|
1494
|
+
export const GetModelCostMapSourceModelCostMapSourceGetRequest =
|
|
1495
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1496
|
+
S.Struct({}).pipe(
|
|
1497
|
+
T.Http({ method: "GET", uri: "/model/cost_map/source", code: 200 }),
|
|
1498
|
+
),
|
|
1499
|
+
).annotate({
|
|
1500
|
+
identifier: "GetModelCostMapSourceModelCostMapSourceGetRequest",
|
|
1501
|
+
}) as any as S.Schema<GetModelCostMapSourceModelCostMapSourceGetRequest>;
|
|
1502
|
+
|
|
1503
|
+
export interface GetModelCostMapSourceModelCostMapSourceGetResponse {
|
|
1504
|
+
body: unknown;
|
|
1505
|
+
}
|
|
1506
|
+
export const GetModelCostMapSourceModelCostMapSourceGetResponse =
|
|
1507
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1508
|
+
S.Struct({
|
|
1509
|
+
body: S.Unknown,
|
|
1510
|
+
}),
|
|
1511
|
+
).annotate({
|
|
1512
|
+
identifier: "GetModelCostMapSourceModelCostMapSourceGetResponse",
|
|
1513
|
+
}) as any as S.Schema<GetModelCostMapSourceModelCostMapSourceGetResponse>;
|
|
1514
|
+
|
|
1515
|
+
export interface GetModelDeprecationsModelDeprecationRequest {
|
|
1516
|
+
warn_within_days?: number;
|
|
1517
|
+
}
|
|
1518
|
+
export const GetModelDeprecationsModelDeprecationRequest =
|
|
1519
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1520
|
+
S.Struct({
|
|
1521
|
+
warn_within_days: S.optional(S.Number.pipe(T.Query())),
|
|
1522
|
+
}).pipe(T.Http({ method: "GET", uri: "/model/deprecations", code: 200 })),
|
|
1523
|
+
).annotate({
|
|
1524
|
+
identifier: "GetModelDeprecationsModelDeprecationRequest",
|
|
1525
|
+
}) as any as S.Schema<GetModelDeprecationsModelDeprecationRequest>;
|
|
1526
|
+
|
|
1527
|
+
/** 'deprecated' if the date has passed, 'imminent' if it falls within warn_within_days, 'upcoming' otherwise. */
|
|
1528
|
+
export type ModelDeprecationInfoStatus = "upcoming" | "imminent" | "deprecated";
|
|
1529
|
+
export const ModelDeprecationInfoStatus = S.String;
|
|
1530
|
+
|
|
1531
|
+
export interface ModelDeprecationInfo {
|
|
1532
|
+
/** Days remaining until the deprecation date. Negative if the model is already deprecated. */
|
|
1533
|
+
days_until_deprecation: number;
|
|
1534
|
+
/** The date (UTC) when the model becomes deprecated. */
|
|
1535
|
+
deprecation_date: string;
|
|
1536
|
+
/** The underlying litellm model string the deprecation date is sourced from. */
|
|
1537
|
+
litellm_model?: string | null;
|
|
1538
|
+
/** The provider this model belongs to. */
|
|
1539
|
+
litellm_provider?: string | null;
|
|
1540
|
+
/** The public name of the model on the proxy (model_group). */
|
|
1541
|
+
model_name: string;
|
|
1542
|
+
/** 'deprecated' if the date has passed, 'imminent' if it falls within warn_within_days, 'upcoming' otherwise. */
|
|
1543
|
+
status: ModelDeprecationInfoStatus;
|
|
1544
|
+
}
|
|
1545
|
+
export const ModelDeprecationInfo = /*@__PURE__*/ S.suspend(() =>
|
|
1546
|
+
S.Struct({
|
|
1547
|
+
days_until_deprecation: S.Number,
|
|
1548
|
+
deprecation_date: S.String,
|
|
1549
|
+
litellm_model: S.optional(S.NullOr(S.String)),
|
|
1550
|
+
litellm_provider: S.optional(S.NullOr(S.String)),
|
|
1551
|
+
model_name: S.String,
|
|
1552
|
+
status: ModelDeprecationInfoStatus,
|
|
1553
|
+
}),
|
|
1554
|
+
).annotate({
|
|
1555
|
+
identifier: "ModelDeprecationInfo",
|
|
1556
|
+
}) as any as S.Schema<ModelDeprecationInfo>;
|
|
1557
|
+
|
|
1558
|
+
/** Models whose deprecation date has already passed. */
|
|
1559
|
+
export type ModelDeprecationResponseDeprecatedList =
|
|
1560
|
+
Array<ModelDeprecationInfo>;
|
|
1561
|
+
export const ModelDeprecationResponseDeprecatedList = /*@__PURE__*/ S.Array(
|
|
1562
|
+
ModelDeprecationInfo,
|
|
1563
|
+
) as any as S.Schema<ModelDeprecationResponseDeprecatedList>;
|
|
1564
|
+
|
|
1565
|
+
/** Models whose deprecation date is within warn_within_days from today and require immediate migration planning. */
|
|
1566
|
+
export type ModelDeprecationResponseImminentList = Array<ModelDeprecationInfo>;
|
|
1567
|
+
export const ModelDeprecationResponseImminentList = /*@__PURE__*/ S.Array(
|
|
1568
|
+
ModelDeprecationInfo,
|
|
1569
|
+
) as any as S.Schema<ModelDeprecationResponseImminentList>;
|
|
1570
|
+
|
|
1571
|
+
/** Models with a future deprecation date outside the warn window. */
|
|
1572
|
+
export type ModelDeprecationResponseUpcomingList = Array<ModelDeprecationInfo>;
|
|
1573
|
+
export const ModelDeprecationResponseUpcomingList = /*@__PURE__*/ S.Array(
|
|
1574
|
+
ModelDeprecationInfo,
|
|
1575
|
+
) as any as S.Schema<ModelDeprecationResponseUpcomingList>;
|
|
1576
|
+
|
|
1577
|
+
export interface ModelDeprecationResponse {
|
|
1578
|
+
/** UTC timestamp when the deprecation snapshot was generated. */
|
|
1579
|
+
checked_at: string;
|
|
1580
|
+
/** Models whose deprecation date has already passed. */
|
|
1581
|
+
deprecated?: ModelDeprecationResponseDeprecatedList;
|
|
1582
|
+
/** Models whose deprecation date is within warn_within_days from today and require immediate migration planning. */
|
|
1583
|
+
imminent?: ModelDeprecationResponseImminentList;
|
|
1584
|
+
/** Models with a future deprecation date outside the warn window. */
|
|
1585
|
+
upcoming?: ModelDeprecationResponseUpcomingList;
|
|
1586
|
+
/** The window (in days) used to bucket 'imminent' models. */
|
|
1587
|
+
warn_within_days: number;
|
|
1588
|
+
}
|
|
1589
|
+
export const ModelDeprecationResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1590
|
+
S.Struct({
|
|
1591
|
+
checked_at: S.String,
|
|
1592
|
+
deprecated: S.optional(ModelDeprecationResponseDeprecatedList),
|
|
1593
|
+
imminent: S.optional(ModelDeprecationResponseImminentList),
|
|
1594
|
+
upcoming: S.optional(ModelDeprecationResponseUpcomingList),
|
|
1595
|
+
warn_within_days: S.Number,
|
|
1596
|
+
}),
|
|
1597
|
+
).annotate({
|
|
1598
|
+
identifier: "ModelDeprecationResponse",
|
|
1599
|
+
}) as any as S.Schema<ModelDeprecationResponse>;
|
|
1600
|
+
|
|
1601
|
+
export interface GetModelDeprecationsV1ModelDeprecationRequest {
|
|
1602
|
+
warn_within_days?: number;
|
|
1603
|
+
}
|
|
1604
|
+
export const GetModelDeprecationsV1ModelDeprecationRequest =
|
|
1605
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1606
|
+
S.Struct({
|
|
1607
|
+
warn_within_days: S.optional(S.Number.pipe(T.Query())),
|
|
1608
|
+
}).pipe(
|
|
1609
|
+
T.Http({ method: "GET", uri: "/v1/model/deprecations", code: 200 }),
|
|
1610
|
+
),
|
|
1611
|
+
).annotate({
|
|
1612
|
+
identifier: "GetModelDeprecationsV1ModelDeprecationRequest",
|
|
1613
|
+
}) as any as S.Schema<GetModelDeprecationsV1ModelDeprecationRequest>;
|
|
1614
|
+
|
|
1615
|
+
export interface GetModelGroupInfoModelGroupInfoRequest {
|
|
1616
|
+
model_group?: string;
|
|
1617
|
+
}
|
|
1618
|
+
export const GetModelGroupInfoModelGroupInfoRequest = /*@__PURE__*/ S.suspend(
|
|
1619
|
+
() =>
|
|
1620
|
+
S.Struct({
|
|
1621
|
+
model_group: S.optional(S.String.pipe(T.Query())),
|
|
1622
|
+
}).pipe(T.Http({ method: "GET", uri: "/model_group/info", code: 200 })),
|
|
1623
|
+
).annotate({
|
|
1624
|
+
identifier: "GetModelGroupInfoModelGroupInfoRequest",
|
|
1625
|
+
}) as any as S.Schema<GetModelGroupInfoModelGroupInfoRequest>;
|
|
1626
|
+
|
|
1627
|
+
export interface GetModelGroupInfoModelGroupInfoResponse {
|
|
1628
|
+
body: unknown;
|
|
1629
|
+
}
|
|
1630
|
+
export const GetModelGroupInfoModelGroupInfoResponse = /*@__PURE__*/ S.suspend(
|
|
1631
|
+
() =>
|
|
1632
|
+
S.Struct({
|
|
1633
|
+
body: S.Unknown,
|
|
1634
|
+
}),
|
|
1635
|
+
).annotate({
|
|
1636
|
+
identifier: "GetModelGroupInfoModelGroupInfoResponse",
|
|
1637
|
+
}) as any as S.Schema<GetModelGroupInfoModelGroupInfoResponse>;
|
|
1638
|
+
|
|
1639
|
+
export interface GetModelInfoModelsModelIdRequest {
|
|
1640
|
+
model_id: string;
|
|
1641
|
+
team_id?: string;
|
|
1642
|
+
healthy_only?: boolean;
|
|
1643
|
+
}
|
|
1644
|
+
export const GetModelInfoModelsModelIdRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1645
|
+
S.Struct({
|
|
1646
|
+
model_id: S.String.pipe(T.Label()),
|
|
1647
|
+
team_id: S.optional(S.String.pipe(T.Query())),
|
|
1648
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
1649
|
+
}).pipe(T.Http({ method: "GET", uri: "/models/{model_id}", code: 200 })),
|
|
1650
|
+
).annotate({
|
|
1651
|
+
identifier: "GetModelInfoModelsModelIdRequest",
|
|
1652
|
+
}) as any as S.Schema<GetModelInfoModelsModelIdRequest>;
|
|
1653
|
+
|
|
1654
|
+
export interface GetModelInfoModelsModelIdResponse {
|
|
1655
|
+
body: unknown;
|
|
1656
|
+
}
|
|
1657
|
+
export const GetModelInfoModelsModelIdResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1658
|
+
S.Struct({
|
|
1659
|
+
body: S.Unknown,
|
|
1660
|
+
}),
|
|
1661
|
+
).annotate({
|
|
1662
|
+
identifier: "GetModelInfoModelsModelIdResponse",
|
|
1663
|
+
}) as any as S.Schema<GetModelInfoModelsModelIdResponse>;
|
|
1664
|
+
|
|
1665
|
+
export interface GetModelInfoV1ModelInfoRequest {
|
|
1666
|
+
litellm_model_id?: string;
|
|
1667
|
+
/** When true, filter to deployments the caller can use via direct access or team membership. */
|
|
1668
|
+
include_team_models?: boolean;
|
|
1669
|
+
/** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
|
|
1670
|
+
teamId?: string;
|
|
1671
|
+
healthy_only?: boolean;
|
|
1672
|
+
}
|
|
1673
|
+
export const GetModelInfoV1ModelInfoRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1674
|
+
S.Struct({
|
|
1675
|
+
litellm_model_id: S.optional(S.String.pipe(T.Query())),
|
|
1676
|
+
include_team_models: S.optional(S.Boolean.pipe(T.Query())),
|
|
1677
|
+
teamId: S.optional(S.String.pipe(T.Query())),
|
|
1678
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
1679
|
+
}).pipe(T.Http({ method: "GET", uri: "/model/info", code: 200 })),
|
|
1680
|
+
).annotate({
|
|
1681
|
+
identifier: "GetModelInfoV1ModelInfoRequest",
|
|
1682
|
+
}) as any as S.Schema<GetModelInfoV1ModelInfoRequest>;
|
|
1683
|
+
|
|
1684
|
+
export interface GetModelInfoV1ModelInfoResponse {
|
|
1685
|
+
body: unknown;
|
|
1686
|
+
}
|
|
1687
|
+
export const GetModelInfoV1ModelInfoResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1688
|
+
S.Struct({
|
|
1689
|
+
body: S.Unknown,
|
|
1690
|
+
}),
|
|
1691
|
+
).annotate({
|
|
1692
|
+
identifier: "GetModelInfoV1ModelInfoResponse",
|
|
1693
|
+
}) as any as S.Schema<GetModelInfoV1ModelInfoResponse>;
|
|
1694
|
+
|
|
1695
|
+
export interface GetModelInfoV1ModelsModelIdRequest {
|
|
1696
|
+
model_id: string;
|
|
1697
|
+
team_id?: string;
|
|
1698
|
+
healthy_only?: boolean;
|
|
1699
|
+
}
|
|
1700
|
+
export const GetModelInfoV1ModelsModelIdRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1701
|
+
S.Struct({
|
|
1702
|
+
model_id: S.String.pipe(T.Label()),
|
|
1703
|
+
team_id: S.optional(S.String.pipe(T.Query())),
|
|
1704
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
1705
|
+
}).pipe(T.Http({ method: "GET", uri: "/v1/models/{model_id}", code: 200 })),
|
|
1706
|
+
).annotate({
|
|
1707
|
+
identifier: "GetModelInfoV1ModelsModelIdRequest",
|
|
1708
|
+
}) as any as S.Schema<GetModelInfoV1ModelsModelIdRequest>;
|
|
1709
|
+
|
|
1710
|
+
export interface GetModelInfoV1ModelsModelIdResponse {
|
|
1711
|
+
body: unknown;
|
|
1712
|
+
}
|
|
1713
|
+
export const GetModelInfoV1ModelsModelIdResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1714
|
+
S.Struct({
|
|
1715
|
+
body: S.Unknown,
|
|
1716
|
+
}),
|
|
1717
|
+
).annotate({
|
|
1718
|
+
identifier: "GetModelInfoV1ModelsModelIdResponse",
|
|
1719
|
+
}) as any as S.Schema<GetModelInfoV1ModelsModelIdResponse>;
|
|
1720
|
+
|
|
1721
|
+
export interface GetModelInfoV1V1ModelInfoRequest {
|
|
1722
|
+
litellm_model_id?: string;
|
|
1723
|
+
/** When true, filter to deployments the caller can use via direct access or team membership. */
|
|
1724
|
+
include_team_models?: boolean;
|
|
1725
|
+
/** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
|
|
1726
|
+
teamId?: string;
|
|
1727
|
+
healthy_only?: boolean;
|
|
1728
|
+
}
|
|
1729
|
+
export const GetModelInfoV1V1ModelInfoRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1730
|
+
S.Struct({
|
|
1731
|
+
litellm_model_id: S.optional(S.String.pipe(T.Query())),
|
|
1732
|
+
include_team_models: S.optional(S.Boolean.pipe(T.Query())),
|
|
1733
|
+
teamId: S.optional(S.String.pipe(T.Query())),
|
|
1734
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
1735
|
+
}).pipe(T.Http({ method: "GET", uri: "/v1/model/info", code: 200 })),
|
|
1736
|
+
).annotate({
|
|
1737
|
+
identifier: "GetModelInfoV1V1ModelInfoRequest",
|
|
1738
|
+
}) as any as S.Schema<GetModelInfoV1V1ModelInfoRequest>;
|
|
1739
|
+
|
|
1740
|
+
export interface GetModelInfoV1V1ModelInfoResponse {
|
|
1741
|
+
body: unknown;
|
|
1742
|
+
}
|
|
1743
|
+
export const GetModelInfoV1V1ModelInfoResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1744
|
+
S.Struct({
|
|
1745
|
+
body: S.Unknown,
|
|
1746
|
+
}),
|
|
1747
|
+
).annotate({
|
|
1748
|
+
identifier: "GetModelInfoV1V1ModelInfoResponse",
|
|
1749
|
+
}) as any as S.Schema<GetModelInfoV1V1ModelInfoResponse>;
|
|
1750
|
+
|
|
1751
|
+
export interface GetModelInfoV2V2ModelInfoRequest {
|
|
1752
|
+
/** Specify the model name (optional) */
|
|
1753
|
+
model?: string;
|
|
1754
|
+
/** Only return models added by this user */
|
|
1755
|
+
user_models_only?: boolean;
|
|
1756
|
+
/** Return all models across all teams user is in. */
|
|
1757
|
+
include_team_models?: boolean;
|
|
1758
|
+
debug?: boolean;
|
|
1759
|
+
/** Page number */
|
|
1760
|
+
page?: number;
|
|
1761
|
+
/** Page size */
|
|
1762
|
+
size?: number;
|
|
1763
|
+
/** Search model names (case-insensitive partial match) */
|
|
1764
|
+
search?: string;
|
|
1765
|
+
/** Search for a specific model by its unique ID */
|
|
1766
|
+
modelId?: string;
|
|
1767
|
+
/** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
|
|
1768
|
+
teamId?: string;
|
|
1769
|
+
/** Field to sort by. Options: model_name, created_at, updated_at, costs, status */
|
|
1770
|
+
sortBy?: string;
|
|
1771
|
+
/** Sort order. Options: asc, desc */
|
|
1772
|
+
sortOrder?: string;
|
|
1773
|
+
/** Omit auto-router deployments (litellm model prefixed `auto_router/`). They select among deployments rather than being deployments themselves, so a caller rendering a deployment list can leave them out. Defaults to false, so existing callers are unaffected */
|
|
1774
|
+
exclude_auto_routers?: boolean;
|
|
1775
|
+
}
|
|
1776
|
+
export const GetModelInfoV2V2ModelInfoRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1777
|
+
S.Struct({
|
|
1778
|
+
model: S.optional(S.String.pipe(T.Query())),
|
|
1779
|
+
user_models_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
1780
|
+
include_team_models: S.optional(S.Boolean.pipe(T.Query())),
|
|
1781
|
+
debug: S.optional(S.Boolean.pipe(T.Query())),
|
|
1782
|
+
page: S.optional(S.Number.pipe(T.Query())),
|
|
1783
|
+
size: S.optional(S.Number.pipe(T.Query())),
|
|
1784
|
+
search: S.optional(S.String.pipe(T.Query())),
|
|
1785
|
+
modelId: S.optional(S.String.pipe(T.Query())),
|
|
1786
|
+
teamId: S.optional(S.String.pipe(T.Query())),
|
|
1787
|
+
sortBy: S.optional(S.String.pipe(T.Query())),
|
|
1788
|
+
sortOrder: S.optional(S.String.pipe(T.Query())),
|
|
1789
|
+
exclude_auto_routers: S.optional(S.Boolean.pipe(T.Query())),
|
|
1790
|
+
}).pipe(T.Http({ method: "GET", uri: "/v2/model/info", code: 200 })),
|
|
1791
|
+
).annotate({
|
|
1792
|
+
identifier: "GetModelInfoV2V2ModelInfoRequest",
|
|
1793
|
+
}) as any as S.Schema<GetModelInfoV2V2ModelInfoRequest>;
|
|
1794
|
+
|
|
1795
|
+
export interface GetModelInfoV2V2ModelInfoResponse {
|
|
1796
|
+
body: unknown;
|
|
1797
|
+
}
|
|
1798
|
+
export const GetModelInfoV2V2ModelInfoResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1799
|
+
S.Struct({
|
|
1800
|
+
body: S.Unknown,
|
|
1801
|
+
}),
|
|
1802
|
+
).annotate({
|
|
1803
|
+
identifier: "GetModelInfoV2V2ModelInfoResponse",
|
|
1804
|
+
}) as any as S.Schema<GetModelInfoV2V2ModelInfoResponse>;
|
|
1805
|
+
|
|
1806
|
+
export interface GetModelMetricsExceptionsModelMetricsExceptionRequest {
|
|
1807
|
+
_selected_model_group?: string;
|
|
1808
|
+
startTime?: string;
|
|
1809
|
+
endTime?: string;
|
|
1810
|
+
api_key?: string;
|
|
1811
|
+
customer?: string;
|
|
1812
|
+
}
|
|
1813
|
+
export const GetModelMetricsExceptionsModelMetricsExceptionRequest =
|
|
1814
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1815
|
+
S.Struct({
|
|
1816
|
+
_selected_model_group: S.optional(S.String.pipe(T.Query())),
|
|
1817
|
+
startTime: S.optional(S.String.pipe(T.Query())),
|
|
1818
|
+
endTime: S.optional(S.String.pipe(T.Query())),
|
|
1819
|
+
api_key: S.optional(S.String.pipe(T.Query())),
|
|
1820
|
+
customer: S.optional(S.String.pipe(T.Query())),
|
|
1821
|
+
}).pipe(
|
|
1822
|
+
T.Http({ method: "GET", uri: "/model/metrics/exceptions", code: 200 }),
|
|
1823
|
+
),
|
|
1824
|
+
).annotate({
|
|
1825
|
+
identifier: "GetModelMetricsExceptionsModelMetricsExceptionRequest",
|
|
1826
|
+
}) as any as S.Schema<GetModelMetricsExceptionsModelMetricsExceptionRequest>;
|
|
1827
|
+
|
|
1828
|
+
export interface GetModelMetricsExceptionsModelMetricsExceptionResponse {
|
|
1829
|
+
body: unknown;
|
|
1830
|
+
}
|
|
1831
|
+
export const GetModelMetricsExceptionsModelMetricsExceptionResponse =
|
|
1832
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1833
|
+
S.Struct({
|
|
1834
|
+
body: S.Unknown,
|
|
1835
|
+
}),
|
|
1836
|
+
).annotate({
|
|
1837
|
+
identifier: "GetModelMetricsExceptionsModelMetricsExceptionResponse",
|
|
1838
|
+
}) as any as S.Schema<GetModelMetricsExceptionsModelMetricsExceptionResponse>;
|
|
1839
|
+
|
|
1840
|
+
export interface GetModelMetricsModelMetricsRequest {
|
|
1841
|
+
_selected_model_group?: string;
|
|
1842
|
+
startTime?: string;
|
|
1843
|
+
endTime?: string;
|
|
1844
|
+
api_key?: string;
|
|
1845
|
+
customer?: string;
|
|
1846
|
+
}
|
|
1847
|
+
export const GetModelMetricsModelMetricsRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1848
|
+
S.Struct({
|
|
1849
|
+
_selected_model_group: S.optional(S.String.pipe(T.Query())),
|
|
1850
|
+
startTime: S.optional(S.String.pipe(T.Query())),
|
|
1851
|
+
endTime: S.optional(S.String.pipe(T.Query())),
|
|
1852
|
+
api_key: S.optional(S.String.pipe(T.Query())),
|
|
1853
|
+
customer: S.optional(S.String.pipe(T.Query())),
|
|
1854
|
+
}).pipe(T.Http({ method: "GET", uri: "/model/metrics", code: 200 })),
|
|
1855
|
+
).annotate({
|
|
1856
|
+
identifier: "GetModelMetricsModelMetricsRequest",
|
|
1857
|
+
}) as any as S.Schema<GetModelMetricsModelMetricsRequest>;
|
|
1858
|
+
|
|
1859
|
+
export interface GetModelMetricsModelMetricsResponse {
|
|
1860
|
+
body: unknown;
|
|
1861
|
+
}
|
|
1862
|
+
export const GetModelMetricsModelMetricsResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1863
|
+
S.Struct({
|
|
1864
|
+
body: S.Unknown,
|
|
1865
|
+
}),
|
|
1866
|
+
).annotate({
|
|
1867
|
+
identifier: "GetModelMetricsModelMetricsResponse",
|
|
1868
|
+
}) as any as S.Schema<GetModelMetricsModelMetricsResponse>;
|
|
1869
|
+
|
|
1870
|
+
export interface GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest {
|
|
1871
|
+
_selected_model_group?: string;
|
|
1872
|
+
startTime?: string;
|
|
1873
|
+
endTime?: string;
|
|
1874
|
+
api_key?: string;
|
|
1875
|
+
customer?: string;
|
|
1876
|
+
}
|
|
1877
|
+
export const GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest =
|
|
1878
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1879
|
+
S.Struct({
|
|
1880
|
+
_selected_model_group: S.optional(S.String.pipe(T.Query())),
|
|
1881
|
+
startTime: S.optional(S.String.pipe(T.Query())),
|
|
1882
|
+
endTime: S.optional(S.String.pipe(T.Query())),
|
|
1883
|
+
api_key: S.optional(S.String.pipe(T.Query())),
|
|
1884
|
+
customer: S.optional(S.String.pipe(T.Query())),
|
|
1885
|
+
}).pipe(
|
|
1886
|
+
T.Http({
|
|
1887
|
+
method: "GET",
|
|
1888
|
+
uri: "/model/metrics/slow_responses",
|
|
1889
|
+
code: 200,
|
|
1890
|
+
}),
|
|
1891
|
+
),
|
|
1892
|
+
).annotate({
|
|
1893
|
+
identifier: "GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest",
|
|
1894
|
+
}) as any as S.Schema<GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest>;
|
|
1895
|
+
|
|
1896
|
+
export interface GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse {
|
|
1897
|
+
body: unknown;
|
|
1898
|
+
}
|
|
1899
|
+
export const GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse =
|
|
1900
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1901
|
+
S.Struct({
|
|
1902
|
+
body: S.Unknown,
|
|
1903
|
+
}),
|
|
1904
|
+
).annotate({
|
|
1905
|
+
identifier: "GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse",
|
|
1906
|
+
}) as any as S.Schema<GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse>;
|
|
1907
|
+
|
|
1908
|
+
export interface GetModelSettingsModelSettingsRequest {}
|
|
1909
|
+
export const GetModelSettingsModelSettingsRequest = /*@__PURE__*/ S.suspend(
|
|
1910
|
+
() =>
|
|
1911
|
+
S.Struct({}).pipe(
|
|
1912
|
+
T.Http({ method: "GET", uri: "/model/settings", code: 200 }),
|
|
1913
|
+
),
|
|
1914
|
+
).annotate({
|
|
1915
|
+
identifier: "GetModelSettingsModelSettingsRequest",
|
|
1916
|
+
}) as any as S.Schema<GetModelSettingsModelSettingsRequest>;
|
|
1917
|
+
|
|
1918
|
+
export interface GetModelSettingsModelSettingsResponse {
|
|
1919
|
+
body: unknown;
|
|
1920
|
+
}
|
|
1921
|
+
export const GetModelSettingsModelSettingsResponse = /*@__PURE__*/ S.suspend(
|
|
1922
|
+
() =>
|
|
1923
|
+
S.Struct({
|
|
1924
|
+
body: S.Unknown,
|
|
1925
|
+
}),
|
|
1926
|
+
).annotate({
|
|
1927
|
+
identifier: "GetModelSettingsModelSettingsResponse",
|
|
1928
|
+
}) as any as S.Schema<GetModelSettingsModelSettingsResponse>;
|
|
1929
|
+
|
|
1930
|
+
export interface GetModelStreamingMetricsModelStreamingMetricsRequest {
|
|
1931
|
+
_selected_model_group?: string;
|
|
1932
|
+
startTime?: string;
|
|
1933
|
+
endTime?: string;
|
|
1934
|
+
}
|
|
1935
|
+
export const GetModelStreamingMetricsModelStreamingMetricsRequest =
|
|
1936
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1937
|
+
S.Struct({
|
|
1938
|
+
_selected_model_group: S.optional(S.String.pipe(T.Query())),
|
|
1939
|
+
startTime: S.optional(S.String.pipe(T.Query())),
|
|
1940
|
+
endTime: S.optional(S.String.pipe(T.Query())),
|
|
1941
|
+
}).pipe(
|
|
1942
|
+
T.Http({ method: "GET", uri: "/model/streaming_metrics", code: 200 }),
|
|
1943
|
+
),
|
|
1944
|
+
).annotate({
|
|
1945
|
+
identifier: "GetModelStreamingMetricsModelStreamingMetricsRequest",
|
|
1946
|
+
}) as any as S.Schema<GetModelStreamingMetricsModelStreamingMetricsRequest>;
|
|
1947
|
+
|
|
1948
|
+
export interface GetModelStreamingMetricsModelStreamingMetricsResponse {
|
|
1949
|
+
body: unknown;
|
|
1950
|
+
}
|
|
1951
|
+
export const GetModelStreamingMetricsModelStreamingMetricsResponse =
|
|
1952
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1953
|
+
S.Struct({
|
|
1954
|
+
body: S.Unknown,
|
|
1955
|
+
}),
|
|
1956
|
+
).annotate({
|
|
1957
|
+
identifier: "GetModelStreamingMetricsModelStreamingMetricsResponse",
|
|
1958
|
+
}) as any as S.Schema<GetModelStreamingMetricsModelStreamingMetricsResponse>;
|
|
1959
|
+
|
|
1960
|
+
export interface ListAccessGroupsAccessGroupListGetRequest {}
|
|
1961
|
+
export const ListAccessGroupsAccessGroupListGetRequest =
|
|
1962
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
1963
|
+
S.Struct({}).pipe(
|
|
1964
|
+
T.Http({ method: "GET", uri: "/access_group/list", code: 200 }),
|
|
1965
|
+
),
|
|
1966
|
+
).annotate({
|
|
1967
|
+
identifier: "ListAccessGroupsAccessGroupListGetRequest",
|
|
1968
|
+
}) as any as S.Schema<ListAccessGroupsAccessGroupListGetRequest>;
|
|
1969
|
+
|
|
1970
|
+
export type ListAccessGroupsResponseAccessGroupsList = Array<AccessGroupInfo>;
|
|
1971
|
+
export const ListAccessGroupsResponseAccessGroupsList = /*@__PURE__*/ S.Array(
|
|
1972
|
+
AccessGroupInfo,
|
|
1973
|
+
) as any as S.Schema<ListAccessGroupsResponseAccessGroupsList>;
|
|
1974
|
+
|
|
1975
|
+
export interface ListAccessGroupsResponse {
|
|
1976
|
+
access_groups: ListAccessGroupsResponseAccessGroupsList;
|
|
1977
|
+
}
|
|
1978
|
+
export const ListAccessGroupsResponse = /*@__PURE__*/ S.suspend(() =>
|
|
1979
|
+
S.Struct({
|
|
1980
|
+
access_groups: ListAccessGroupsResponseAccessGroupsList,
|
|
1981
|
+
}),
|
|
1982
|
+
).annotate({
|
|
1983
|
+
identifier: "ListAccessGroupsResponse",
|
|
1984
|
+
}) as any as S.Schema<ListAccessGroupsResponse>;
|
|
1985
|
+
|
|
1986
|
+
export interface ModelListModelsGetRequest {
|
|
1987
|
+
return_wildcard_routes?: boolean;
|
|
1988
|
+
team_id?: string;
|
|
1989
|
+
include_model_access_groups?: boolean;
|
|
1990
|
+
only_model_access_groups?: boolean;
|
|
1991
|
+
include_metadata?: boolean;
|
|
1992
|
+
fallback_type?: string;
|
|
1993
|
+
scope?: string;
|
|
1994
|
+
healthy_only?: boolean;
|
|
1995
|
+
}
|
|
1996
|
+
export const ModelListModelsGetRequest = /*@__PURE__*/ S.suspend(() =>
|
|
1997
|
+
S.Struct({
|
|
1998
|
+
return_wildcard_routes: S.optional(S.Boolean.pipe(T.Query())),
|
|
1999
|
+
team_id: S.optional(S.String.pipe(T.Query())),
|
|
2000
|
+
include_model_access_groups: S.optional(S.Boolean.pipe(T.Query())),
|
|
2001
|
+
only_model_access_groups: S.optional(S.Boolean.pipe(T.Query())),
|
|
2002
|
+
include_metadata: S.optional(S.Boolean.pipe(T.Query())),
|
|
2003
|
+
fallback_type: S.optional(S.String.pipe(T.Query())),
|
|
2004
|
+
scope: S.optional(S.String.pipe(T.Query())),
|
|
2005
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
2006
|
+
}).pipe(T.Http({ method: "GET", uri: "/models", code: 200 })),
|
|
2007
|
+
).annotate({
|
|
2008
|
+
identifier: "ModelListModelsGetRequest",
|
|
2009
|
+
}) as any as S.Schema<ModelListModelsGetRequest>;
|
|
2010
|
+
|
|
2011
|
+
export interface ModelListModelsGetResponse {
|
|
2012
|
+
body: unknown;
|
|
2013
|
+
}
|
|
2014
|
+
export const ModelListModelsGetResponse = /*@__PURE__*/ S.suspend(() =>
|
|
2015
|
+
S.Struct({
|
|
2016
|
+
body: S.Unknown,
|
|
2017
|
+
}),
|
|
2018
|
+
).annotate({
|
|
2019
|
+
identifier: "ModelListModelsGetResponse",
|
|
2020
|
+
}) as any as S.Schema<ModelListModelsGetResponse>;
|
|
2021
|
+
|
|
2022
|
+
export interface ModelListV1ModelsGetRequest {
|
|
2023
|
+
return_wildcard_routes?: boolean;
|
|
2024
|
+
team_id?: string;
|
|
2025
|
+
include_model_access_groups?: boolean;
|
|
2026
|
+
only_model_access_groups?: boolean;
|
|
2027
|
+
include_metadata?: boolean;
|
|
2028
|
+
fallback_type?: string;
|
|
2029
|
+
scope?: string;
|
|
2030
|
+
healthy_only?: boolean;
|
|
2031
|
+
}
|
|
2032
|
+
export const ModelListV1ModelsGetRequest = /*@__PURE__*/ S.suspend(() =>
|
|
2033
|
+
S.Struct({
|
|
2034
|
+
return_wildcard_routes: S.optional(S.Boolean.pipe(T.Query())),
|
|
2035
|
+
team_id: S.optional(S.String.pipe(T.Query())),
|
|
2036
|
+
include_model_access_groups: S.optional(S.Boolean.pipe(T.Query())),
|
|
2037
|
+
only_model_access_groups: S.optional(S.Boolean.pipe(T.Query())),
|
|
2038
|
+
include_metadata: S.optional(S.Boolean.pipe(T.Query())),
|
|
2039
|
+
fallback_type: S.optional(S.String.pipe(T.Query())),
|
|
2040
|
+
scope: S.optional(S.String.pipe(T.Query())),
|
|
2041
|
+
healthy_only: S.optional(S.Boolean.pipe(T.Query())),
|
|
2042
|
+
}).pipe(T.Http({ method: "GET", uri: "/v1/models", code: 200 })),
|
|
2043
|
+
).annotate({
|
|
2044
|
+
identifier: "ModelListV1ModelsGetRequest",
|
|
2045
|
+
}) as any as S.Schema<ModelListV1ModelsGetRequest>;
|
|
2046
|
+
|
|
2047
|
+
export interface ModelListV1ModelsGetResponse {
|
|
2048
|
+
body: unknown;
|
|
2049
|
+
}
|
|
2050
|
+
export const ModelListV1ModelsGetResponse = /*@__PURE__*/ S.suspend(() =>
|
|
2051
|
+
S.Struct({
|
|
2052
|
+
body: S.Unknown,
|
|
2053
|
+
}),
|
|
2054
|
+
).annotate({
|
|
2055
|
+
identifier: "ModelListV1ModelsGetResponse",
|
|
2056
|
+
}) as any as S.Schema<ModelListV1ModelsGetResponse>;
|
|
2057
|
+
|
|
2058
|
+
export type UpdateLiteLLMParamsAdaptiveRouterConfigMap = {
|
|
2059
|
+
[key: string]: unknown | undefined;
|
|
2060
|
+
};
|
|
2061
|
+
export const UpdateLiteLLMParamsAdaptiveRouterConfigMap =
|
|
2062
|
+
/*@__PURE__*/ S.Record(
|
|
2063
|
+
S.String,
|
|
2064
|
+
S.Unknown,
|
|
2065
|
+
) as any as S.Schema<UpdateLiteLLMParamsAdaptiveRouterConfigMap>;
|
|
2066
|
+
|
|
2067
|
+
export type UpdateLiteLLMParamsBedrockTagsList = Array<unknown>;
|
|
2068
|
+
export const UpdateLiteLLMParamsBedrockTagsList = /*@__PURE__*/ S.Array(
|
|
2069
|
+
S.Unknown,
|
|
2070
|
+
) as any as S.Schema<UpdateLiteLLMParamsBedrockTagsList>;
|
|
2071
|
+
|
|
2072
|
+
export type UpdateLiteLLMParamsComplexityRouterConfigMap = {
|
|
2073
|
+
[key: string]: unknown | undefined;
|
|
2074
|
+
};
|
|
2075
|
+
export const UpdateLiteLLMParamsComplexityRouterConfigMap =
|
|
2076
|
+
/*@__PURE__*/ S.Record(
|
|
2077
|
+
S.String,
|
|
2078
|
+
S.Unknown,
|
|
2079
|
+
) as any as S.Schema<UpdateLiteLLMParamsComplexityRouterConfigMap>;
|
|
2080
|
+
|
|
2081
|
+
export type UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem =
|
|
2082
|
+
| string
|
|
2083
|
+
| ConfigurableClientsideParamsCustomAuthInput;
|
|
2084
|
+
export const UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem =
|
|
2085
|
+
S.Unknown as any as S.Schema<UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem>;
|
|
2086
|
+
|
|
2087
|
+
export type UpdateLiteLLMParamsConfigurableClientsideAuthParamsList =
|
|
2088
|
+
Array<UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem>;
|
|
2089
|
+
export const UpdateLiteLLMParamsConfigurableClientsideAuthParamsList =
|
|
2090
|
+
/*@__PURE__*/ S.Array(
|
|
2091
|
+
UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem,
|
|
2092
|
+
) as any as S.Schema<UpdateLiteLLMParamsConfigurableClientsideAuthParamsList>;
|
|
2093
|
+
|
|
2094
|
+
export type UpdateLiteLLMParamsMilvusPartitionNamesList = Array<string>;
|
|
2095
|
+
export const UpdateLiteLLMParamsMilvusPartitionNamesList =
|
|
2096
|
+
/*@__PURE__*/ S.Array(
|
|
2097
|
+
S.String,
|
|
2098
|
+
) as any as S.Schema<UpdateLiteLLMParamsMilvusPartitionNamesList>;
|
|
2099
|
+
|
|
2100
|
+
export type UpdateLiteLLMParamsMockResponse = string | ModelResponse | unknown;
|
|
2101
|
+
export const UpdateLiteLLMParamsMockResponse =
|
|
2102
|
+
S.Unknown as any as S.Schema<UpdateLiteLLMParamsMockResponse>;
|
|
2103
|
+
|
|
2104
|
+
export type UpdateLiteLLMParamsModelInfoMap = {
|
|
2105
|
+
[key: string]: unknown | undefined;
|
|
2106
|
+
};
|
|
2107
|
+
export const UpdateLiteLLMParamsModelInfoMap = /*@__PURE__*/ S.Record(
|
|
2108
|
+
S.String,
|
|
2109
|
+
S.Unknown,
|
|
2110
|
+
) as any as S.Schema<UpdateLiteLLMParamsModelInfoMap>;
|
|
2111
|
+
|
|
2112
|
+
export type UpdateLiteLLMParamsQualityRouterConfigMap = {
|
|
2113
|
+
[key: string]: unknown | undefined;
|
|
2114
|
+
};
|
|
2115
|
+
export const UpdateLiteLLMParamsQualityRouterConfigMap = /*@__PURE__*/ S.Record(
|
|
2116
|
+
S.String,
|
|
2117
|
+
S.Unknown,
|
|
2118
|
+
) as any as S.Schema<UpdateLiteLLMParamsQualityRouterConfigMap>;
|
|
2119
|
+
|
|
2120
|
+
export type UpdateLiteLLMParamsSearchContextCostPerQueryMap = {
|
|
2121
|
+
[key: string]: unknown | undefined;
|
|
2122
|
+
};
|
|
2123
|
+
export const UpdateLiteLLMParamsSearchContextCostPerQueryMap =
|
|
2124
|
+
/*@__PURE__*/ S.Record(
|
|
2125
|
+
S.String,
|
|
2126
|
+
S.Unknown,
|
|
2127
|
+
) as any as S.Schema<UpdateLiteLLMParamsSearchContextCostPerQueryMap>;
|
|
2128
|
+
|
|
2129
|
+
export type UpdateLiteLLMParamsStreamTimeout = number | string;
|
|
2130
|
+
export const UpdateLiteLLMParamsStreamTimeout =
|
|
2131
|
+
S.Unknown as any as S.Schema<UpdateLiteLLMParamsStreamTimeout>;
|
|
2132
|
+
|
|
2133
|
+
export type UpdateLiteLLMParamsTagRegexList = Array<string>;
|
|
2134
|
+
export const UpdateLiteLLMParamsTagRegexList = /*@__PURE__*/ S.Array(
|
|
2135
|
+
S.String,
|
|
2136
|
+
) as any as S.Schema<UpdateLiteLLMParamsTagRegexList>;
|
|
2137
|
+
|
|
2138
|
+
export type UpdateLiteLLMParamsTagsList = Array<string>;
|
|
2139
|
+
export const UpdateLiteLLMParamsTagsList = /*@__PURE__*/ S.Array(
|
|
2140
|
+
S.String,
|
|
2141
|
+
) as any as S.Schema<UpdateLiteLLMParamsTagsList>;
|
|
2142
|
+
|
|
2143
|
+
export type UpdateLiteLLMParamsTieredPricingItemMap = {
|
|
2144
|
+
[key: string]: unknown | undefined;
|
|
2145
|
+
};
|
|
2146
|
+
export const UpdateLiteLLMParamsTieredPricingItemMap = /*@__PURE__*/ S.Record(
|
|
2147
|
+
S.String,
|
|
2148
|
+
S.Unknown,
|
|
2149
|
+
) as any as S.Schema<UpdateLiteLLMParamsTieredPricingItemMap>;
|
|
2150
|
+
|
|
2151
|
+
export type UpdateLiteLLMParamsTieredPricingList =
|
|
2152
|
+
Array<UpdateLiteLLMParamsTieredPricingItemMap>;
|
|
2153
|
+
export const UpdateLiteLLMParamsTieredPricingList = /*@__PURE__*/ S.Array(
|
|
2154
|
+
UpdateLiteLLMParamsTieredPricingItemMap,
|
|
2155
|
+
) as any as S.Schema<UpdateLiteLLMParamsTieredPricingList>;
|
|
2156
|
+
|
|
2157
|
+
export type UpdateLiteLLMParamsTimeout = number | string;
|
|
2158
|
+
export const UpdateLiteLLMParamsTimeout =
|
|
2159
|
+
S.Unknown as any as S.Schema<UpdateLiteLLMParamsTimeout>;
|
|
2160
|
+
|
|
2161
|
+
export type UpdateLiteLLMParamsVertexCredentialsCase1Map = {
|
|
2162
|
+
[key: string]: unknown | undefined;
|
|
2163
|
+
};
|
|
2164
|
+
export const UpdateLiteLLMParamsVertexCredentialsCase1Map =
|
|
2165
|
+
/*@__PURE__*/ S.Record(
|
|
2166
|
+
S.String,
|
|
2167
|
+
S.Unknown,
|
|
2168
|
+
) as any as S.Schema<UpdateLiteLLMParamsVertexCredentialsCase1Map>;
|
|
2169
|
+
|
|
2170
|
+
export type UpdateLiteLLMParamsVertexCredentials =
|
|
2171
|
+
| string
|
|
2172
|
+
| UpdateLiteLLMParamsVertexCredentialsCase1Map;
|
|
2173
|
+
export const UpdateLiteLLMParamsVertexCredentials =
|
|
2174
|
+
S.Unknown as any as S.Schema<UpdateLiteLLMParamsVertexCredentials>;
|
|
2175
|
+
|
|
2176
|
+
export interface UpdateLiteLLMParams {
|
|
2177
|
+
adaptive_router_config?: UpdateLiteLLMParamsAdaptiveRouterConfigMap | null;
|
|
2178
|
+
adaptive_router_default_model?: string | null;
|
|
2179
|
+
allow_client_keepalive_override?: boolean | null;
|
|
2180
|
+
annotation_cost_per_page?: number | null;
|
|
2181
|
+
api_base?: string | null;
|
|
2182
|
+
api_key?: string | null;
|
|
2183
|
+
api_version?: string | null;
|
|
2184
|
+
auto_router_config?: string | null;
|
|
2185
|
+
auto_router_config_path?: string | null;
|
|
2186
|
+
auto_router_default_model?: string | null;
|
|
2187
|
+
auto_router_embedding_model?: string | null;
|
|
2188
|
+
auto_router_max_input_chars?: number | null;
|
|
2189
|
+
aws_access_key_id?: string | null;
|
|
2190
|
+
aws_batch_role_arn?: string | null;
|
|
2191
|
+
aws_bedrock_project_id?: string | null;
|
|
2192
|
+
aws_bedrock_runtime_endpoint?: string | null;
|
|
2193
|
+
aws_external_id?: string | null;
|
|
2194
|
+
aws_profile_name?: string | null;
|
|
2195
|
+
aws_region_name?: string | null;
|
|
2196
|
+
aws_role_name?: string | null;
|
|
2197
|
+
aws_secret_access_key?: string | null;
|
|
2198
|
+
aws_session_name?: string | null;
|
|
2199
|
+
aws_session_token?: string | null;
|
|
2200
|
+
aws_sts_endpoint?: string | null;
|
|
2201
|
+
aws_web_identity_token?: string | null;
|
|
2202
|
+
azure_ad_token?: string | null;
|
|
2203
|
+
bedrock_tags?: UpdateLiteLLMParamsBedrockTagsList | null;
|
|
2204
|
+
budget_duration?: string | null;
|
|
2205
|
+
cache_creation_input_audio_token_cost?: number | null;
|
|
2206
|
+
cache_creation_input_token_cost?: number | null;
|
|
2207
|
+
cache_creation_input_token_cost_above_1hr?: number | null;
|
|
2208
|
+
cache_creation_input_token_cost_above_200k_tokens?: number | null;
|
|
2209
|
+
cache_creation_input_token_cost_above_272k_tokens?: number | null;
|
|
2210
|
+
cache_creation_input_token_cost_above_272k_tokens_flex?: number | null;
|
|
2211
|
+
cache_creation_input_token_cost_above_272k_tokens_priority?: number | null;
|
|
2212
|
+
cache_creation_input_token_cost_flex?: number | null;
|
|
2213
|
+
cache_creation_input_token_cost_priority?: number | null;
|
|
2214
|
+
cache_creation_input_token_cost_ultrafast?: number | null;
|
|
2215
|
+
cache_read_input_audio_token_cost?: number | null;
|
|
2216
|
+
cache_read_input_token_cost?: number | null;
|
|
2217
|
+
cache_read_input_token_cost_above_200k_tokens?: number | null;
|
|
2218
|
+
cache_read_input_token_cost_above_200k_tokens_priority?: number | null;
|
|
2219
|
+
cache_read_input_token_cost_above_272k_tokens?: number | null;
|
|
2220
|
+
cache_read_input_token_cost_above_272k_tokens_flex?: number | null;
|
|
2221
|
+
cache_read_input_token_cost_above_272k_tokens_priority?: number | null;
|
|
2222
|
+
cache_read_input_token_cost_above_512k_tokens?: number | null;
|
|
2223
|
+
cache_read_input_token_cost_flex?: number | null;
|
|
2224
|
+
cache_read_input_token_cost_priority?: number | null;
|
|
2225
|
+
cache_read_input_token_cost_ultrafast?: number | null;
|
|
2226
|
+
citation_cost_per_token?: number | null;
|
|
2227
|
+
complexity_router_config?: UpdateLiteLLMParamsComplexityRouterConfigMap | null;
|
|
2228
|
+
complexity_router_default_model?: string | null;
|
|
2229
|
+
configurable_clientside_auth_params?: UpdateLiteLLMParamsConfigurableClientsideAuthParamsList | null;
|
|
2230
|
+
custom_llm_provider?: string | null;
|
|
2231
|
+
default_api_key_rpm_limit?: number | null;
|
|
2232
|
+
default_api_key_tpm_limit?: number | null;
|
|
2233
|
+
gcs_bucket_name?: string | null;
|
|
2234
|
+
google_maps_grounding_cost_per_query?: number | null;
|
|
2235
|
+
input_cost_per_audio_per_second?: number | null;
|
|
2236
|
+
input_cost_per_audio_per_second_above_128k_tokens?: number | null;
|
|
2237
|
+
input_cost_per_audio_token?: number | null;
|
|
2238
|
+
input_cost_per_character?: number | null;
|
|
2239
|
+
input_cost_per_character_above_128k_tokens?: number | null;
|
|
2240
|
+
input_cost_per_image?: number | null;
|
|
2241
|
+
input_cost_per_image_above_128k_tokens?: number | null;
|
|
2242
|
+
input_cost_per_image_token?: number | null;
|
|
2243
|
+
input_cost_per_pixel?: number | null;
|
|
2244
|
+
input_cost_per_query?: number | null;
|
|
2245
|
+
input_cost_per_second?: number | null;
|
|
2246
|
+
input_cost_per_token?: number | null;
|
|
2247
|
+
input_cost_per_token_above_128k_tokens?: number | null;
|
|
2248
|
+
input_cost_per_token_above_200k_tokens?: number | null;
|
|
2249
|
+
input_cost_per_token_above_200k_tokens_priority?: number | null;
|
|
2250
|
+
input_cost_per_token_above_272k_tokens?: number | null;
|
|
2251
|
+
input_cost_per_token_above_272k_tokens_flex?: number | null;
|
|
2252
|
+
input_cost_per_token_above_272k_tokens_priority?: number | null;
|
|
2253
|
+
input_cost_per_token_above_512k_tokens?: number | null;
|
|
2254
|
+
input_cost_per_token_batches?: number | null;
|
|
2255
|
+
input_cost_per_token_cache_hit?: number | null;
|
|
2256
|
+
input_cost_per_token_flex?: number | null;
|
|
2257
|
+
input_cost_per_token_priority?: number | null;
|
|
2258
|
+
input_cost_per_token_ultrafast?: number | null;
|
|
2259
|
+
input_cost_per_video_per_second?: number | null;
|
|
2260
|
+
input_cost_per_video_per_second_above_128k_tokens?: number | null;
|
|
2261
|
+
input_cost_per_video_per_second_above_15s_interval?: number | null;
|
|
2262
|
+
input_cost_per_video_per_second_above_8s_interval?: number | null;
|
|
2263
|
+
input_cost_per_video_token?: number | null;
|
|
2264
|
+
itpm?: number | null;
|
|
2265
|
+
keepalive_seconds?: number | null;
|
|
2266
|
+
litellm_credential_name?: string | null;
|
|
2267
|
+
litellm_trace_id?: string | null;
|
|
2268
|
+
max_budget?: number | null;
|
|
2269
|
+
max_file_size_mb?: number | null;
|
|
2270
|
+
max_retries?: number | null;
|
|
2271
|
+
merge_reasoning_content_in_choices?: boolean | null;
|
|
2272
|
+
milvus_db_name?: string | null;
|
|
2273
|
+
milvus_partition_names?: UpdateLiteLLMParamsMilvusPartitionNamesList | null;
|
|
2274
|
+
milvus_text_field?: string | null;
|
|
2275
|
+
mock_response?: UpdateLiteLLMParamsMockResponse | null;
|
|
2276
|
+
model?: string | null;
|
|
2277
|
+
model_info?: UpdateLiteLLMParamsModelInfoMap | null;
|
|
2278
|
+
ocr_cost_per_credit?: number | null;
|
|
2279
|
+
ocr_cost_per_page?: number | null;
|
|
2280
|
+
organization?: string | null;
|
|
2281
|
+
otpm?: number | null;
|
|
2282
|
+
output_cost_per_audio_per_second?: number | null;
|
|
2283
|
+
output_cost_per_audio_token?: number | null;
|
|
2284
|
+
output_cost_per_character?: number | null;
|
|
2285
|
+
output_cost_per_character_above_128k_tokens?: number | null;
|
|
2286
|
+
output_cost_per_image?: number | null;
|
|
2287
|
+
output_cost_per_image_token?: number | null;
|
|
2288
|
+
output_cost_per_pixel?: number | null;
|
|
2289
|
+
output_cost_per_reasoning_token?: number | null;
|
|
2290
|
+
output_cost_per_reasoning_token_flex?: number | null;
|
|
2291
|
+
output_cost_per_reasoning_token_priority?: number | null;
|
|
2292
|
+
output_cost_per_second?: number | null;
|
|
2293
|
+
output_cost_per_second_1080p?: number | null;
|
|
2294
|
+
output_cost_per_second_480p?: number | null;
|
|
2295
|
+
output_cost_per_second_4k?: number | null;
|
|
2296
|
+
output_cost_per_token?: number | null;
|
|
2297
|
+
output_cost_per_token_above_128k_tokens?: number | null;
|
|
2298
|
+
output_cost_per_token_above_200k_tokens?: number | null;
|
|
2299
|
+
output_cost_per_token_above_200k_tokens_priority?: number | null;
|
|
2300
|
+
output_cost_per_token_above_272k_tokens?: number | null;
|
|
2301
|
+
output_cost_per_token_above_272k_tokens_flex?: number | null;
|
|
2302
|
+
output_cost_per_token_above_272k_tokens_priority?: number | null;
|
|
2303
|
+
output_cost_per_token_above_512k_tokens?: number | null;
|
|
2304
|
+
output_cost_per_token_batches?: number | null;
|
|
2305
|
+
output_cost_per_token_flex?: number | null;
|
|
2306
|
+
output_cost_per_token_priority?: number | null;
|
|
2307
|
+
output_cost_per_token_ultrafast?: number | null;
|
|
2308
|
+
output_cost_per_video_per_second?: number | null;
|
|
2309
|
+
output_cost_per_video_token?: number | null;
|
|
2310
|
+
output_vector_size?: number | null;
|
|
2311
|
+
quality_router_config?: UpdateLiteLLMParamsQualityRouterConfigMap | null;
|
|
2312
|
+
quality_router_default_model?: string | null;
|
|
2313
|
+
region_name?: string | null;
|
|
2314
|
+
regional_endpoint_uplift_multiplier?: number | null;
|
|
2315
|
+
regional_processing_uplift_multiplier_eu?: number | null;
|
|
2316
|
+
regional_processing_uplift_multiplier_us?: number | null;
|
|
2317
|
+
rpm?: number | null;
|
|
2318
|
+
s3_bucket_name?: string | null;
|
|
2319
|
+
s3_encryption_key_id?: string | null;
|
|
2320
|
+
s3_output_bucket_name?: string | null;
|
|
2321
|
+
s3_region_name?: string | null;
|
|
2322
|
+
search_context_cost_per_query?: UpdateLiteLLMParamsSearchContextCostPerQueryMap | null;
|
|
2323
|
+
stream_timeout?: UpdateLiteLLMParamsStreamTimeout | null;
|
|
2324
|
+
tag_regex?: UpdateLiteLLMParamsTagRegexList | null;
|
|
2325
|
+
tags?: UpdateLiteLLMParamsTagsList | null;
|
|
2326
|
+
tiered_pricing?: UpdateLiteLLMParamsTieredPricingList | null;
|
|
2327
|
+
timeout?: UpdateLiteLLMParamsTimeout | null;
|
|
2328
|
+
tpm?: number | null;
|
|
2329
|
+
use_chat_completions_api?: boolean | null;
|
|
2330
|
+
use_in_pass_through?: boolean | null;
|
|
2331
|
+
use_litellm_proxy?: boolean | null;
|
|
2332
|
+
/** Use stored xAI OAuth credentials when no xAI API key is configured. */
|
|
2333
|
+
use_xai_oauth?: boolean | null;
|
|
2334
|
+
valkey_embedding_field?: string | null;
|
|
2335
|
+
valkey_host?: string | null;
|
|
2336
|
+
valkey_password?: string | null;
|
|
2337
|
+
valkey_port?: number | null;
|
|
2338
|
+
valkey_ssl?: boolean | null;
|
|
2339
|
+
valkey_text_field?: string | null;
|
|
2340
|
+
vector_store_id?: string | null;
|
|
2341
|
+
vertex_credentials?: UpdateLiteLLMParamsVertexCredentials | null;
|
|
2342
|
+
vertex_location?: string | null;
|
|
2343
|
+
vertex_project?: string | null;
|
|
2344
|
+
watsonx_region_name?: string | null;
|
|
2345
|
+
}
|
|
2346
|
+
export const UpdateLiteLLMParams = /*@__PURE__*/ S.suspend(() =>
|
|
2347
|
+
S.Struct({
|
|
2348
|
+
adaptive_router_config: S.optional(
|
|
2349
|
+
S.NullOr(UpdateLiteLLMParamsAdaptiveRouterConfigMap),
|
|
2350
|
+
),
|
|
2351
|
+
adaptive_router_default_model: S.optional(S.NullOr(S.String)),
|
|
2352
|
+
allow_client_keepalive_override: S.optional(S.NullOr(S.Boolean)),
|
|
2353
|
+
annotation_cost_per_page: S.optional(S.NullOr(S.Number)),
|
|
2354
|
+
api_base: S.optional(S.NullOr(S.String)),
|
|
2355
|
+
api_key: S.optional(S.NullOr(S.String)),
|
|
2356
|
+
api_version: S.optional(S.NullOr(S.String)),
|
|
2357
|
+
auto_router_config: S.optional(S.NullOr(S.String)),
|
|
2358
|
+
auto_router_config_path: S.optional(S.NullOr(S.String)),
|
|
2359
|
+
auto_router_default_model: S.optional(S.NullOr(S.String)),
|
|
2360
|
+
auto_router_embedding_model: S.optional(S.NullOr(S.String)),
|
|
2361
|
+
auto_router_max_input_chars: S.optional(S.NullOr(S.Number)),
|
|
2362
|
+
aws_access_key_id: S.optional(S.NullOr(S.String)),
|
|
2363
|
+
aws_batch_role_arn: S.optional(S.NullOr(S.String)),
|
|
2364
|
+
aws_bedrock_project_id: S.optional(S.NullOr(S.String)),
|
|
2365
|
+
aws_bedrock_runtime_endpoint: S.optional(S.NullOr(S.String)),
|
|
2366
|
+
aws_external_id: S.optional(S.NullOr(S.String)),
|
|
2367
|
+
aws_profile_name: S.optional(S.NullOr(S.String)),
|
|
2368
|
+
aws_region_name: S.optional(S.NullOr(S.String)),
|
|
2369
|
+
aws_role_name: S.optional(S.NullOr(S.String)),
|
|
2370
|
+
aws_secret_access_key: S.optional(S.NullOr(S.String)),
|
|
2371
|
+
aws_session_name: S.optional(S.NullOr(S.String)),
|
|
2372
|
+
aws_session_token: S.optional(S.NullOr(S.String)),
|
|
2373
|
+
aws_sts_endpoint: S.optional(S.NullOr(S.String)),
|
|
2374
|
+
aws_web_identity_token: S.optional(S.NullOr(S.String)),
|
|
2375
|
+
azure_ad_token: S.optional(S.NullOr(S.String)),
|
|
2376
|
+
bedrock_tags: S.optional(S.NullOr(UpdateLiteLLMParamsBedrockTagsList)),
|
|
2377
|
+
budget_duration: S.optional(S.NullOr(S.String)),
|
|
2378
|
+
cache_creation_input_audio_token_cost: S.optional(S.NullOr(S.Number)),
|
|
2379
|
+
cache_creation_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
2380
|
+
cache_creation_input_token_cost_above_1hr: S.optional(S.NullOr(S.Number)),
|
|
2381
|
+
cache_creation_input_token_cost_above_200k_tokens: S.optional(
|
|
2382
|
+
S.NullOr(S.Number),
|
|
2383
|
+
),
|
|
2384
|
+
cache_creation_input_token_cost_above_272k_tokens: S.optional(
|
|
2385
|
+
S.NullOr(S.Number),
|
|
2386
|
+
),
|
|
2387
|
+
cache_creation_input_token_cost_above_272k_tokens_flex: S.optional(
|
|
2388
|
+
S.NullOr(S.Number),
|
|
2389
|
+
),
|
|
2390
|
+
cache_creation_input_token_cost_above_272k_tokens_priority: S.optional(
|
|
2391
|
+
S.NullOr(S.Number),
|
|
2392
|
+
),
|
|
2393
|
+
cache_creation_input_token_cost_flex: S.optional(S.NullOr(S.Number)),
|
|
2394
|
+
cache_creation_input_token_cost_priority: S.optional(S.NullOr(S.Number)),
|
|
2395
|
+
cache_creation_input_token_cost_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
2396
|
+
cache_read_input_audio_token_cost: S.optional(S.NullOr(S.Number)),
|
|
2397
|
+
cache_read_input_token_cost: S.optional(S.NullOr(S.Number)),
|
|
2398
|
+
cache_read_input_token_cost_above_200k_tokens: S.optional(
|
|
2399
|
+
S.NullOr(S.Number),
|
|
2400
|
+
),
|
|
2401
|
+
cache_read_input_token_cost_above_200k_tokens_priority: S.optional(
|
|
2402
|
+
S.NullOr(S.Number),
|
|
2403
|
+
),
|
|
2404
|
+
cache_read_input_token_cost_above_272k_tokens: S.optional(
|
|
2405
|
+
S.NullOr(S.Number),
|
|
2406
|
+
),
|
|
2407
|
+
cache_read_input_token_cost_above_272k_tokens_flex: S.optional(
|
|
2408
|
+
S.NullOr(S.Number),
|
|
2409
|
+
),
|
|
2410
|
+
cache_read_input_token_cost_above_272k_tokens_priority: S.optional(
|
|
2411
|
+
S.NullOr(S.Number),
|
|
2412
|
+
),
|
|
2413
|
+
cache_read_input_token_cost_above_512k_tokens: S.optional(
|
|
2414
|
+
S.NullOr(S.Number),
|
|
2415
|
+
),
|
|
2416
|
+
cache_read_input_token_cost_flex: S.optional(S.NullOr(S.Number)),
|
|
2417
|
+
cache_read_input_token_cost_priority: S.optional(S.NullOr(S.Number)),
|
|
2418
|
+
cache_read_input_token_cost_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
2419
|
+
citation_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
2420
|
+
complexity_router_config: S.optional(
|
|
2421
|
+
S.NullOr(UpdateLiteLLMParamsComplexityRouterConfigMap),
|
|
2422
|
+
),
|
|
2423
|
+
complexity_router_default_model: S.optional(S.NullOr(S.String)),
|
|
2424
|
+
configurable_clientside_auth_params: S.optional(
|
|
2425
|
+
S.NullOr(UpdateLiteLLMParamsConfigurableClientsideAuthParamsList),
|
|
2426
|
+
),
|
|
2427
|
+
custom_llm_provider: S.optional(S.NullOr(S.String)),
|
|
2428
|
+
default_api_key_rpm_limit: S.optional(S.NullOr(S.Number)),
|
|
2429
|
+
default_api_key_tpm_limit: S.optional(S.NullOr(S.Number)),
|
|
2430
|
+
gcs_bucket_name: S.optional(S.NullOr(S.String)),
|
|
2431
|
+
google_maps_grounding_cost_per_query: S.optional(S.NullOr(S.Number)),
|
|
2432
|
+
input_cost_per_audio_per_second: S.optional(S.NullOr(S.Number)),
|
|
2433
|
+
input_cost_per_audio_per_second_above_128k_tokens: S.optional(
|
|
2434
|
+
S.NullOr(S.Number),
|
|
2435
|
+
),
|
|
2436
|
+
input_cost_per_audio_token: S.optional(S.NullOr(S.Number)),
|
|
2437
|
+
input_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
2438
|
+
input_cost_per_character_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2439
|
+
input_cost_per_image: S.optional(S.NullOr(S.Number)),
|
|
2440
|
+
input_cost_per_image_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2441
|
+
input_cost_per_image_token: S.optional(S.NullOr(S.Number)),
|
|
2442
|
+
input_cost_per_pixel: S.optional(S.NullOr(S.Number)),
|
|
2443
|
+
input_cost_per_query: S.optional(S.NullOr(S.Number)),
|
|
2444
|
+
input_cost_per_second: S.optional(S.NullOr(S.Number)),
|
|
2445
|
+
input_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
2446
|
+
input_cost_per_token_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2447
|
+
input_cost_per_token_above_200k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2448
|
+
input_cost_per_token_above_200k_tokens_priority: S.optional(
|
|
2449
|
+
S.NullOr(S.Number),
|
|
2450
|
+
),
|
|
2451
|
+
input_cost_per_token_above_272k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2452
|
+
input_cost_per_token_above_272k_tokens_flex: S.optional(S.NullOr(S.Number)),
|
|
2453
|
+
input_cost_per_token_above_272k_tokens_priority: S.optional(
|
|
2454
|
+
S.NullOr(S.Number),
|
|
2455
|
+
),
|
|
2456
|
+
input_cost_per_token_above_512k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2457
|
+
input_cost_per_token_batches: S.optional(S.NullOr(S.Number)),
|
|
2458
|
+
input_cost_per_token_cache_hit: S.optional(S.NullOr(S.Number)),
|
|
2459
|
+
input_cost_per_token_flex: S.optional(S.NullOr(S.Number)),
|
|
2460
|
+
input_cost_per_token_priority: S.optional(S.NullOr(S.Number)),
|
|
2461
|
+
input_cost_per_token_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
2462
|
+
input_cost_per_video_per_second: S.optional(S.NullOr(S.Number)),
|
|
2463
|
+
input_cost_per_video_per_second_above_128k_tokens: S.optional(
|
|
2464
|
+
S.NullOr(S.Number),
|
|
2465
|
+
),
|
|
2466
|
+
input_cost_per_video_per_second_above_15s_interval: S.optional(
|
|
2467
|
+
S.NullOr(S.Number),
|
|
2468
|
+
),
|
|
2469
|
+
input_cost_per_video_per_second_above_8s_interval: S.optional(
|
|
2470
|
+
S.NullOr(S.Number),
|
|
2471
|
+
),
|
|
2472
|
+
input_cost_per_video_token: S.optional(S.NullOr(S.Number)),
|
|
2473
|
+
itpm: S.optional(S.NullOr(S.Number)),
|
|
2474
|
+
keepalive_seconds: S.optional(S.NullOr(S.Number)),
|
|
2475
|
+
litellm_credential_name: S.optional(S.NullOr(S.String)),
|
|
2476
|
+
litellm_trace_id: S.optional(S.NullOr(S.String)),
|
|
2477
|
+
max_budget: S.optional(S.NullOr(S.Number)),
|
|
2478
|
+
max_file_size_mb: S.optional(S.NullOr(S.Number)),
|
|
2479
|
+
max_retries: S.optional(S.NullOr(S.Number)),
|
|
2480
|
+
merge_reasoning_content_in_choices: S.optional(S.NullOr(S.Boolean)),
|
|
2481
|
+
milvus_db_name: S.optional(S.NullOr(S.String)),
|
|
2482
|
+
milvus_partition_names: S.optional(
|
|
2483
|
+
S.NullOr(UpdateLiteLLMParamsMilvusPartitionNamesList),
|
|
2484
|
+
),
|
|
2485
|
+
milvus_text_field: S.optional(S.NullOr(S.String)),
|
|
2486
|
+
mock_response: S.optional(S.NullOr(UpdateLiteLLMParamsMockResponse)),
|
|
2487
|
+
model: S.optional(S.NullOr(S.String)),
|
|
2488
|
+
model_info: S.optional(S.NullOr(UpdateLiteLLMParamsModelInfoMap)),
|
|
2489
|
+
ocr_cost_per_credit: S.optional(S.NullOr(S.Number)),
|
|
2490
|
+
ocr_cost_per_page: S.optional(S.NullOr(S.Number)),
|
|
2491
|
+
organization: S.optional(S.NullOr(S.String)),
|
|
2492
|
+
otpm: S.optional(S.NullOr(S.Number)),
|
|
2493
|
+
output_cost_per_audio_per_second: S.optional(S.NullOr(S.Number)),
|
|
2494
|
+
output_cost_per_audio_token: S.optional(S.NullOr(S.Number)),
|
|
2495
|
+
output_cost_per_character: S.optional(S.NullOr(S.Number)),
|
|
2496
|
+
output_cost_per_character_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2497
|
+
output_cost_per_image: S.optional(S.NullOr(S.Number)),
|
|
2498
|
+
output_cost_per_image_token: S.optional(S.NullOr(S.Number)),
|
|
2499
|
+
output_cost_per_pixel: S.optional(S.NullOr(S.Number)),
|
|
2500
|
+
output_cost_per_reasoning_token: S.optional(S.NullOr(S.Number)),
|
|
2501
|
+
output_cost_per_reasoning_token_flex: S.optional(S.NullOr(S.Number)),
|
|
2502
|
+
output_cost_per_reasoning_token_priority: S.optional(S.NullOr(S.Number)),
|
|
2503
|
+
output_cost_per_second: S.optional(S.NullOr(S.Number)),
|
|
2504
|
+
output_cost_per_second_1080p: S.optional(S.NullOr(S.Number)),
|
|
2505
|
+
output_cost_per_second_480p: S.optional(S.NullOr(S.Number)),
|
|
2506
|
+
output_cost_per_second_4k: S.optional(S.NullOr(S.Number)),
|
|
2507
|
+
output_cost_per_token: S.optional(S.NullOr(S.Number)),
|
|
2508
|
+
output_cost_per_token_above_128k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2509
|
+
output_cost_per_token_above_200k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2510
|
+
output_cost_per_token_above_200k_tokens_priority: S.optional(
|
|
2511
|
+
S.NullOr(S.Number),
|
|
2512
|
+
),
|
|
2513
|
+
output_cost_per_token_above_272k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2514
|
+
output_cost_per_token_above_272k_tokens_flex: S.optional(
|
|
2515
|
+
S.NullOr(S.Number),
|
|
2516
|
+
),
|
|
2517
|
+
output_cost_per_token_above_272k_tokens_priority: S.optional(
|
|
2518
|
+
S.NullOr(S.Number),
|
|
2519
|
+
),
|
|
2520
|
+
output_cost_per_token_above_512k_tokens: S.optional(S.NullOr(S.Number)),
|
|
2521
|
+
output_cost_per_token_batches: S.optional(S.NullOr(S.Number)),
|
|
2522
|
+
output_cost_per_token_flex: S.optional(S.NullOr(S.Number)),
|
|
2523
|
+
output_cost_per_token_priority: S.optional(S.NullOr(S.Number)),
|
|
2524
|
+
output_cost_per_token_ultrafast: S.optional(S.NullOr(S.Number)),
|
|
2525
|
+
output_cost_per_video_per_second: S.optional(S.NullOr(S.Number)),
|
|
2526
|
+
output_cost_per_video_token: S.optional(S.NullOr(S.Number)),
|
|
2527
|
+
output_vector_size: S.optional(S.NullOr(S.Number)),
|
|
2528
|
+
quality_router_config: S.optional(
|
|
2529
|
+
S.NullOr(UpdateLiteLLMParamsQualityRouterConfigMap),
|
|
2530
|
+
),
|
|
2531
|
+
quality_router_default_model: S.optional(S.NullOr(S.String)),
|
|
2532
|
+
region_name: S.optional(S.NullOr(S.String)),
|
|
2533
|
+
regional_endpoint_uplift_multiplier: S.optional(S.NullOr(S.Number)),
|
|
2534
|
+
regional_processing_uplift_multiplier_eu: S.optional(S.NullOr(S.Number)),
|
|
2535
|
+
regional_processing_uplift_multiplier_us: S.optional(S.NullOr(S.Number)),
|
|
2536
|
+
rpm: S.optional(S.NullOr(S.Number)),
|
|
2537
|
+
s3_bucket_name: S.optional(S.NullOr(S.String)),
|
|
2538
|
+
s3_encryption_key_id: S.optional(S.NullOr(S.String)),
|
|
2539
|
+
s3_output_bucket_name: S.optional(S.NullOr(S.String)),
|
|
2540
|
+
s3_region_name: S.optional(S.NullOr(S.String)),
|
|
2541
|
+
search_context_cost_per_query: S.optional(
|
|
2542
|
+
S.NullOr(UpdateLiteLLMParamsSearchContextCostPerQueryMap),
|
|
2543
|
+
),
|
|
2544
|
+
stream_timeout: S.optional(S.NullOr(UpdateLiteLLMParamsStreamTimeout)),
|
|
2545
|
+
tag_regex: S.optional(S.NullOr(UpdateLiteLLMParamsTagRegexList)),
|
|
2546
|
+
tags: S.optional(S.NullOr(UpdateLiteLLMParamsTagsList)),
|
|
2547
|
+
tiered_pricing: S.optional(S.NullOr(UpdateLiteLLMParamsTieredPricingList)),
|
|
2548
|
+
timeout: S.optional(S.NullOr(UpdateLiteLLMParamsTimeout)),
|
|
2549
|
+
tpm: S.optional(S.NullOr(S.Number)),
|
|
2550
|
+
use_chat_completions_api: S.optional(S.NullOr(S.Boolean)),
|
|
2551
|
+
use_in_pass_through: S.optional(S.NullOr(S.Boolean)),
|
|
2552
|
+
use_litellm_proxy: S.optional(S.NullOr(S.Boolean)),
|
|
2553
|
+
use_xai_oauth: S.optional(S.NullOr(S.Boolean)),
|
|
2554
|
+
valkey_embedding_field: S.optional(S.NullOr(S.String)),
|
|
2555
|
+
valkey_host: S.optional(S.NullOr(S.String)),
|
|
2556
|
+
valkey_password: S.optional(S.NullOr(S.String)),
|
|
2557
|
+
valkey_port: S.optional(S.NullOr(S.Number)),
|
|
2558
|
+
valkey_ssl: S.optional(S.NullOr(S.Boolean)),
|
|
2559
|
+
valkey_text_field: S.optional(S.NullOr(S.String)),
|
|
2560
|
+
vector_store_id: S.optional(S.NullOr(S.String)),
|
|
2561
|
+
vertex_credentials: S.optional(
|
|
2562
|
+
S.NullOr(UpdateLiteLLMParamsVertexCredentials),
|
|
2563
|
+
),
|
|
2564
|
+
vertex_location: S.optional(S.NullOr(S.String)),
|
|
2565
|
+
vertex_project: S.optional(S.NullOr(S.String)),
|
|
2566
|
+
watsonx_region_name: S.optional(S.NullOr(S.String)),
|
|
2567
|
+
}),
|
|
2568
|
+
).annotate({
|
|
2569
|
+
identifier: "UpdateLiteLLMParams",
|
|
2570
|
+
}) as any as S.Schema<UpdateLiteLLMParams>;
|
|
2571
|
+
|
|
2572
|
+
export interface PatchModelModelModelIdUpdatePatchRequest {
|
|
2573
|
+
model_id: string;
|
|
2574
|
+
blocked?: boolean | null;
|
|
2575
|
+
litellm_params?: UpdateLiteLLMParams | null;
|
|
2576
|
+
model_info?: LitellmTypesRouterModelInfo | null;
|
|
2577
|
+
model_name?: string | null;
|
|
2578
|
+
}
|
|
2579
|
+
export const PatchModelModelModelIdUpdatePatchRequest = /*@__PURE__*/ S.suspend(
|
|
2580
|
+
() =>
|
|
2581
|
+
S.Struct({
|
|
2582
|
+
model_id: S.String.pipe(T.Label()),
|
|
2583
|
+
blocked: S.optional(S.NullOr(S.Boolean)),
|
|
2584
|
+
litellm_params: S.optional(S.NullOr(UpdateLiteLLMParams)),
|
|
2585
|
+
model_info: S.optional(S.NullOr(LitellmTypesRouterModelInfo)),
|
|
2586
|
+
model_name: S.optional(S.NullOr(S.String)),
|
|
2587
|
+
}).pipe(
|
|
2588
|
+
T.Http({ method: "PATCH", uri: "/model/{model_id}/update", code: 200 }),
|
|
2589
|
+
),
|
|
2590
|
+
).annotate({
|
|
2591
|
+
identifier: "PatchModelModelModelIdUpdatePatchRequest",
|
|
2592
|
+
}) as any as S.Schema<PatchModelModelModelIdUpdatePatchRequest>;
|
|
2593
|
+
|
|
2594
|
+
export interface PatchModelModelModelIdUpdatePatchResponse {
|
|
2595
|
+
body: unknown;
|
|
2596
|
+
}
|
|
2597
|
+
export const PatchModelModelModelIdUpdatePatchResponse =
|
|
2598
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
2599
|
+
S.Struct({
|
|
2600
|
+
body: S.Unknown,
|
|
2601
|
+
}),
|
|
2602
|
+
).annotate({
|
|
2603
|
+
identifier: "PatchModelModelModelIdUpdatePatchResponse",
|
|
2604
|
+
}) as any as S.Schema<PatchModelModelModelIdUpdatePatchResponse>;
|
|
2605
|
+
|
|
2606
|
+
export interface PostBlockModelModelBlockRequest {
|
|
2607
|
+
model_id: string;
|
|
2608
|
+
}
|
|
2609
|
+
export const PostBlockModelModelBlockRequest = /*@__PURE__*/ S.suspend(() =>
|
|
2610
|
+
S.Struct({
|
|
2611
|
+
model_id: S.String,
|
|
2612
|
+
}).pipe(T.Http({ method: "POST", uri: "/model/block", code: 200 })),
|
|
2613
|
+
).annotate({
|
|
2614
|
+
identifier: "PostBlockModelModelBlockRequest",
|
|
2615
|
+
}) as any as S.Schema<PostBlockModelModelBlockRequest>;
|
|
2616
|
+
|
|
2617
|
+
export type LiteLLMProxyModelTableLitellmParamsMap = {
|
|
2618
|
+
[key: string]: unknown | undefined;
|
|
2619
|
+
};
|
|
2620
|
+
export const LiteLLMProxyModelTableLitellmParamsMap = /*@__PURE__*/ S.Record(
|
|
2621
|
+
S.String,
|
|
2622
|
+
S.Unknown,
|
|
2623
|
+
) as any as S.Schema<LiteLLMProxyModelTableLitellmParamsMap>;
|
|
2624
|
+
|
|
2625
|
+
export type LiteLLMProxyModelTableModelInfoMap = {
|
|
2626
|
+
[key: string]: unknown | undefined;
|
|
2627
|
+
};
|
|
2628
|
+
export const LiteLLMProxyModelTableModelInfoMap = /*@__PURE__*/ S.Record(
|
|
2629
|
+
S.String,
|
|
2630
|
+
S.Unknown,
|
|
2631
|
+
) as any as S.Schema<LiteLLMProxyModelTableModelInfoMap>;
|
|
2632
|
+
|
|
2633
|
+
export interface LiteLLMProxyModelTable {
|
|
2634
|
+
blocked?: boolean;
|
|
2635
|
+
created_at?: string | null;
|
|
2636
|
+
created_by?: string | null;
|
|
2637
|
+
litellm_params: LiteLLMProxyModelTableLitellmParamsMap;
|
|
2638
|
+
model_id: string;
|
|
2639
|
+
model_info?: LiteLLMProxyModelTableModelInfoMap | null;
|
|
2640
|
+
model_name: string;
|
|
2641
|
+
updated_at?: string | null;
|
|
2642
|
+
updated_by?: string | null;
|
|
2643
|
+
}
|
|
2644
|
+
export const LiteLLMProxyModelTable = /*@__PURE__*/ S.suspend(() =>
|
|
2645
|
+
S.Struct({
|
|
2646
|
+
blocked: S.optional(S.Boolean),
|
|
2647
|
+
created_at: S.optional(S.NullOr(S.String)),
|
|
2648
|
+
created_by: S.optional(S.NullOr(S.String)),
|
|
2649
|
+
litellm_params: LiteLLMProxyModelTableLitellmParamsMap,
|
|
2650
|
+
model_id: S.String,
|
|
2651
|
+
model_info: S.optional(S.NullOr(LiteLLMProxyModelTableModelInfoMap)),
|
|
2652
|
+
model_name: S.String,
|
|
2653
|
+
updated_at: S.optional(S.NullOr(S.String)),
|
|
2654
|
+
updated_by: S.optional(S.NullOr(S.String)),
|
|
2655
|
+
}),
|
|
2656
|
+
).annotate({
|
|
2657
|
+
identifier: "LiteLLMProxyModelTable",
|
|
2658
|
+
}) as any as S.Schema<LiteLLMProxyModelTable>;
|
|
2659
|
+
|
|
2660
|
+
export interface PostBlockModelModelBlockResponse {
|
|
2661
|
+
body: LiteLLMProxyModelTable | null;
|
|
2662
|
+
}
|
|
2663
|
+
export const PostBlockModelModelBlockResponse = /*@__PURE__*/ S.suspend(() =>
|
|
2664
|
+
S.Struct({
|
|
2665
|
+
body: S.NullOr(LiteLLMProxyModelTable),
|
|
2666
|
+
}),
|
|
2667
|
+
).annotate({
|
|
2668
|
+
identifier: "PostBlockModelModelBlockResponse",
|
|
2669
|
+
}) as any as S.Schema<PostBlockModelModelBlockResponse>;
|
|
2670
|
+
|
|
2671
|
+
/** An operator-defined tier: the name the LLM classifier must return and its rubric description. */
|
|
2672
|
+
export interface TierDefinition {
|
|
2673
|
+
/** What belongs in this tier; rendered as this tier's bullet in the classifier rubric. Required unless the name is a built-in tier (SIMPLE/MEDIUM/COMPLEX/REASONING), which inherits the built-in criteria when omitted */
|
|
2674
|
+
description?: string | null;
|
|
2675
|
+
/** Tier name; becomes a value the LLM classifier can return and a key of `tiers` */
|
|
2676
|
+
name: string;
|
|
2677
|
+
}
|
|
2678
|
+
export const TierDefinition = /*@__PURE__*/ S.suspend(() =>
|
|
2679
|
+
S.Struct({
|
|
2680
|
+
description: S.optional(S.NullOr(S.String)),
|
|
2681
|
+
name: S.String,
|
|
2682
|
+
}),
|
|
2683
|
+
).annotate({ identifier: "TierDefinition" }) as any as S.Schema<TierDefinition>;
|
|
2684
|
+
|
|
2685
|
+
export type PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList =
|
|
2686
|
+
Array<TierDefinition>;
|
|
2687
|
+
export const PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList =
|
|
2688
|
+
/*@__PURE__*/ S.Array(
|
|
2689
|
+
TierDefinition,
|
|
2690
|
+
) as any as S.Schema<PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList>;
|
|
2691
|
+
|
|
2692
|
+
export interface PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest {
|
|
2693
|
+
classification_prompt?: string | null;
|
|
2694
|
+
context_window_size?: number;
|
|
2695
|
+
tier_definitions: PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList;
|
|
2696
|
+
}
|
|
2697
|
+
export const PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest =
|
|
2698
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
2699
|
+
S.Struct({
|
|
2700
|
+
classification_prompt: S.optional(S.NullOr(S.String)),
|
|
2701
|
+
context_window_size: S.optional(S.Number),
|
|
2702
|
+
tier_definitions:
|
|
2703
|
+
PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList,
|
|
2704
|
+
}).pipe(
|
|
2705
|
+
T.Http({
|
|
2706
|
+
method: "POST",
|
|
2707
|
+
uri: "/auto_router/classifier/default_prompt",
|
|
2708
|
+
code: 200,
|
|
2709
|
+
}),
|
|
2710
|
+
),
|
|
2711
|
+
).annotate({
|
|
2712
|
+
identifier:
|
|
2713
|
+
"PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest",
|
|
2714
|
+
}) as any as S.Schema<PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest>;
|
|
2715
|
+
|
|
2716
|
+
/** When adaptive=True: 'all' scores every pool model with a tier-distance penalty (soft floors); 'classified_tier' Thompson-samples only inside the classified tier's pool */
|
|
2717
|
+
export type RequestComplexityRouterConfigAdaptiveEligible =
|
|
2718
|
+
| "all"
|
|
2719
|
+
| "classified_tier";
|
|
2720
|
+
export const RequestComplexityRouterConfigAdaptiveEligible = S.String;
|
|
2721
|
+
|
|
2722
|
+
export interface AdaptiveRouterWeights {
|
|
2723
|
+
cost?: number;
|
|
2724
|
+
quality?: number;
|
|
2725
|
+
}
|
|
2726
|
+
export const AdaptiveRouterWeights = /*@__PURE__*/ S.suspend(() =>
|
|
2727
|
+
S.Struct({
|
|
2728
|
+
cost: S.optional(S.Number),
|
|
2729
|
+
quality: S.optional(S.Number),
|
|
2730
|
+
}),
|
|
2731
|
+
).annotate({
|
|
2732
|
+
identifier: "AdaptiveRouterWeights",
|
|
2733
|
+
}) as any as S.Schema<AdaptiveRouterWeights>;
|
|
2734
|
+
|
|
2735
|
+
/** What classifies the request when the LLM classifier errors, times out, or returns an unparseable response. 'heuristic' runs the local complexity scorer, which is right when the classifier grades complexity too. 'default_model' skips scoring and routes to default_model, which is what a classifier on some other taxonomy wants: a prompt that grades data sensitivity has no use for a complexity score, and scoring one produces a tier unrelated to what the operator configured. Requires default_model when set to 'default_model'. Only applies when classifier_type is 'llm', 'custom', or 'heuristic_first'. */
|
|
2736
|
+
export type RequestComplexityRouterConfigClassifierFallback =
|
|
2737
|
+
| "heuristic"
|
|
2738
|
+
| "default_model";
|
|
2739
|
+
export const RequestComplexityRouterConfigClassifierFallback = S.String;
|
|
2740
|
+
|
|
2741
|
+
/** Configuration for the LLM-based complexity classifier. */
|
|
2742
|
+
export interface ClassifierLLMConfig {
|
|
2743
|
+
/** Which calibration examples the built-in rubric carries. 'agentic' anchors routine installs, builds, multi-file edits, and standard debugging at MEDIUM, so ordinary engineering does not route to the most expensive tier; it suits agent, terminal, and coding-assistant traffic as well as mixed traffic. 'chat' omits those engineering anchors, for a deployment serving only conversational traffic. 'business' carries business/sales anchors and business-flavored tier criteria that keep routine drafting and summarizing off the expensive tiers and reserve the top tier for committing to decisions under tradeoffs; it suits sales, support, and go-to-market traffic. Every preset keeps the same four tiers, so this moves where the boundary sits without changing the taxonomy. Leave unset for 'legacy', the rubric as it shipped before calibration examples existed, so an existing router's tier decisions and spend do not move on upgrade. Mutually exclusive with system_prompt, which replaces the rubric this would select. Only applies when classifier_type is 'llm'. */
|
|
2744
|
+
classification_rubric?: ClassificationRubric | (string & {}) | null;
|
|
2745
|
+
/** Model name (from the router's model_list) to call for classification */
|
|
2746
|
+
model: string;
|
|
2747
|
+
/** Replaces the built-in complexity rubric as the classifier's entire system role. When set, neither the default rubric nor the context-window closing line is appended, so the prompt owns the whole taxonomy and the tier names SIMPLE/MEDIUM/COMPLEX/REASONING become whatever buckets it defines: a prompt that classifies data sensitivity routes on that instead of on difficulty. Two consequences of full replacement. The default rubric's closing paragraph is the classifier's prompt-injection defense, telling it that the caller's quoted system prompt and prior turns are material to judge and never instructions; a replacement that omits it lets a caller ask for a tier and get it. And the heuristic fallback still scores complexity, so a router on some other taxonomy wants classifier_fallback='default_model'. Leave unset for the built-in rubric. Only applies when classifier_type is 'llm'. */
|
|
2748
|
+
system_prompt?: string | null;
|
|
2749
|
+
/** Timeout budget for the classification call, in milliseconds */
|
|
2750
|
+
timeout_ms?: number;
|
|
2751
|
+
}
|
|
2752
|
+
export const ClassifierLLMConfig = /*@__PURE__*/ S.suspend(() =>
|
|
2753
|
+
S.Struct({
|
|
2754
|
+
classification_rubric: S.optional(S.NullOr(ClassificationRubric)),
|
|
2755
|
+
model: S.String,
|
|
2756
|
+
system_prompt: S.optional(S.NullOr(S.String)),
|
|
2757
|
+
timeout_ms: S.optional(S.Number),
|
|
2758
|
+
}),
|
|
2759
|
+
).annotate({
|
|
2760
|
+
identifier: "ClassifierLLMConfig",
|
|
2761
|
+
}) as any as S.Schema<ClassifierLLMConfig>;
|
|
2762
|
+
|
|
2763
|
+
/** Classification strategy: local regex/keyword scoring, an LLM call, a custom classifier plugin, or 'heuristic_first', which scores locally and only pays for the LLM classifier when the local scorer does not confidently land a cheap tier */
|
|
2764
|
+
export type RequestComplexityRouterConfigClassifierType =
|
|
2765
|
+
| "heuristic"
|
|
2766
|
+
| "llm"
|
|
2767
|
+
| "custom"
|
|
2768
|
+
| "heuristic_first";
|
|
2769
|
+
export const RequestComplexityRouterConfigClassifierType = S.String;
|
|
2770
|
+
|
|
2771
|
+
export type RequestComplexityRouterConfigCodeKeywordsList = Array<string>;
|
|
2772
|
+
export const RequestComplexityRouterConfigCodeKeywordsList =
|
|
2773
|
+
/*@__PURE__*/ S.Array(
|
|
2774
|
+
S.String,
|
|
2775
|
+
) as any as S.Schema<RequestComplexityRouterConfigCodeKeywordsList>;
|
|
2776
|
+
|
|
2777
|
+
export type RequestComplexityRouterConfigCustomTechnicalKeywordsList =
|
|
2778
|
+
Array<string>;
|
|
2779
|
+
export const RequestComplexityRouterConfigCustomTechnicalKeywordsList =
|
|
2780
|
+
/*@__PURE__*/ S.Array(
|
|
2781
|
+
S.String,
|
|
2782
|
+
) as any as S.Schema<RequestComplexityRouterConfigCustomTechnicalKeywordsList>;
|
|
2783
|
+
|
|
2784
|
+
/** Weights for each scoring dimension */
|
|
2785
|
+
export type RequestComplexityRouterConfigDimensionWeightsMap = {
|
|
2786
|
+
[key: string]: number | undefined;
|
|
2787
|
+
};
|
|
2788
|
+
export const RequestComplexityRouterConfigDimensionWeightsMap =
|
|
2789
|
+
/*@__PURE__*/ S.Record(
|
|
2790
|
+
S.String,
|
|
2791
|
+
S.Number,
|
|
2792
|
+
) as any as S.Schema<RequestComplexityRouterConfigDimensionWeightsMap>;
|
|
2793
|
+
|
|
2794
|
+
export type RequestComplexityRouterConfigEscalationKeywordsList = Array<string>;
|
|
2795
|
+
export const RequestComplexityRouterConfigEscalationKeywordsList =
|
|
2796
|
+
/*@__PURE__*/ S.Array(
|
|
2797
|
+
S.String,
|
|
2798
|
+
) as any as S.Schema<RequestComplexityRouterConfigEscalationKeywordsList>;
|
|
2799
|
+
|
|
2800
|
+
export type RequestComplexityRouterConfigHousekeepingPatternsList =
|
|
2801
|
+
Array<string>;
|
|
2802
|
+
export const RequestComplexityRouterConfigHousekeepingPatternsList =
|
|
2803
|
+
/*@__PURE__*/ S.Array(
|
|
2804
|
+
S.String,
|
|
2805
|
+
) as any as S.Schema<RequestComplexityRouterConfigHousekeepingPatternsList>;
|
|
2806
|
+
|
|
2807
|
+
/** Keywords/phrases that trigger this rule (lexical or semantic match) */
|
|
2808
|
+
export type KeywordTierRuleKeywordsList = Array<string>;
|
|
2809
|
+
export const KeywordTierRuleKeywordsList = /*@__PURE__*/ S.Array(
|
|
2810
|
+
S.String,
|
|
2811
|
+
) as any as S.Schema<KeywordTierRuleKeywordsList>;
|
|
2812
|
+
|
|
2813
|
+
/** A deterministic override: if any keyword matches, route to this tier. */
|
|
2814
|
+
export interface KeywordTierRule {
|
|
2815
|
+
/** Keywords/phrases that trigger this rule (lexical or semantic match) */
|
|
2816
|
+
keywords: KeywordTierRuleKeywordsList;
|
|
2817
|
+
/** Tier to route to when this rule matches: a built-in tier name, or with tier_definitions set, one of the defined tier names */
|
|
2818
|
+
tier: string;
|
|
2819
|
+
}
|
|
2820
|
+
export const KeywordTierRule = /*@__PURE__*/ S.suspend(() =>
|
|
2821
|
+
S.Struct({
|
|
2822
|
+
keywords: KeywordTierRuleKeywordsList,
|
|
2823
|
+
tier: S.String,
|
|
2824
|
+
}),
|
|
2825
|
+
).annotate({
|
|
2826
|
+
identifier: "KeywordTierRule",
|
|
2827
|
+
}) as any as S.Schema<KeywordTierRule>;
|
|
2828
|
+
|
|
2829
|
+
export type RequestComplexityRouterConfigKeywordTierRulesList =
|
|
2830
|
+
Array<KeywordTierRule>;
|
|
2831
|
+
export const RequestComplexityRouterConfigKeywordTierRulesList =
|
|
2832
|
+
/*@__PURE__*/ S.Array(
|
|
2833
|
+
KeywordTierRule,
|
|
2834
|
+
) as any as S.Schema<RequestComplexityRouterConfigKeywordTierRulesList>;
|
|
2835
|
+
|
|
2836
|
+
export type RequestComplexityRouterConfigPlanModePatternsList = Array<string>;
|
|
2837
|
+
export const RequestComplexityRouterConfigPlanModePatternsList =
|
|
2838
|
+
/*@__PURE__*/ S.Array(
|
|
2839
|
+
S.String,
|
|
2840
|
+
) as any as S.Schema<RequestComplexityRouterConfigPlanModePatternsList>;
|
|
2841
|
+
|
|
2842
|
+
export type RequestComplexityRouterConfigReasoningKeywordsList = Array<string>;
|
|
2843
|
+
export const RequestComplexityRouterConfigReasoningKeywordsList =
|
|
2844
|
+
/*@__PURE__*/ S.Array(
|
|
2845
|
+
S.String,
|
|
2846
|
+
) as any as S.Schema<RequestComplexityRouterConfigReasoningKeywordsList>;
|
|
2847
|
+
|
|
2848
|
+
/** One open/close delimiter pair a harness wraps injected context in. Normalizing here rather than at the scan is what makes matching case-insensitive: markers reach the scan already lowered, so it lowercases only the haystack and never the needles. Stripping keeps YAML indentation whitespace from becoming part of the delimiter. */
|
|
2849
|
+
export interface ReminderMarkerPair {
|
|
2850
|
+
/** Closing delimiter, e.g. '</system-reminder>' */
|
|
2851
|
+
close: string;
|
|
2852
|
+
/** Opening delimiter, e.g. '<system-reminder>' */
|
|
2853
|
+
open: string;
|
|
2854
|
+
}
|
|
2855
|
+
export const ReminderMarkerPair = /*@__PURE__*/ S.suspend(() =>
|
|
2856
|
+
S.Struct({
|
|
2857
|
+
close: S.String,
|
|
2858
|
+
open: S.String,
|
|
2859
|
+
}),
|
|
2860
|
+
).annotate({
|
|
2861
|
+
identifier: "ReminderMarkerPair",
|
|
2862
|
+
}) as any as S.Schema<ReminderMarkerPair>;
|
|
2863
|
+
|
|
2864
|
+
export type RequestComplexityRouterConfigReminderMarkersList =
|
|
2865
|
+
Array<ReminderMarkerPair>;
|
|
2866
|
+
export const RequestComplexityRouterConfigReminderMarkersList =
|
|
2867
|
+
/*@__PURE__*/ S.Array(
|
|
2868
|
+
ReminderMarkerPair,
|
|
2869
|
+
) as any as S.Schema<RequestComplexityRouterConfigReminderMarkersList>;
|
|
2870
|
+
|
|
2871
|
+
export type RequestComplexityRouterConfigSimpleKeywordsList = Array<string>;
|
|
2872
|
+
export const RequestComplexityRouterConfigSimpleKeywordsList =
|
|
2873
|
+
/*@__PURE__*/ S.Array(
|
|
2874
|
+
S.String,
|
|
2875
|
+
) as any as S.Schema<RequestComplexityRouterConfigSimpleKeywordsList>;
|
|
2876
|
+
|
|
2877
|
+
export type RequestComplexityRouterConfigTechnicalKeywordsList = Array<string>;
|
|
2878
|
+
export const RequestComplexityRouterConfigTechnicalKeywordsList =
|
|
2879
|
+
/*@__PURE__*/ S.Array(
|
|
2880
|
+
S.String,
|
|
2881
|
+
) as any as S.Schema<RequestComplexityRouterConfigTechnicalKeywordsList>;
|
|
2882
|
+
|
|
2883
|
+
/** Score boundaries between tiers. These keys (simple_medium, medium_complex, complex_reasoning) name the gaps between the default tier names and are not renameable by tier_labels; they are scorer knobs persisted by name on every routing decision */
|
|
2884
|
+
export type RequestComplexityRouterConfigTierBoundariesMap = {
|
|
2885
|
+
[key: string]: number | undefined;
|
|
2886
|
+
};
|
|
2887
|
+
export const RequestComplexityRouterConfigTierBoundariesMap =
|
|
2888
|
+
/*@__PURE__*/ S.Record(
|
|
2889
|
+
S.String,
|
|
2890
|
+
S.Number,
|
|
2891
|
+
) as any as S.Schema<RequestComplexityRouterConfigTierBoundariesMap>;
|
|
2892
|
+
|
|
2893
|
+
export type RequestComplexityRouterConfigTierDefinitionsList =
|
|
2894
|
+
Array<TierDefinition>;
|
|
2895
|
+
export const RequestComplexityRouterConfigTierDefinitionsList =
|
|
2896
|
+
/*@__PURE__*/ S.Array(
|
|
2897
|
+
TierDefinition,
|
|
2898
|
+
) as any as S.Schema<RequestComplexityRouterConfigTierDefinitionsList>;
|
|
2899
|
+
|
|
2900
|
+
/** Display names for the complexity tiers, so a deployment can use its own vocabulary (e.g. Cheap/Standard/Premium/Deep) in the dashboard, spend logs, and the LLM classifier rubric. Purely operator-facing: config keys stay canonical (tiers, keyword_tier_rules[].tier, tier_boundaries), API callers never see these names, and the heuristic scorer never reads them. Unlisted tiers keep their canonical name. Partial maps are allowed. */
|
|
2901
|
+
export type RequestComplexityRouterConfigTierLabelsMap = {
|
|
2902
|
+
[key: string]: string | undefined;
|
|
2903
|
+
};
|
|
2904
|
+
export const RequestComplexityRouterConfigTierLabelsMap =
|
|
2905
|
+
/*@__PURE__*/ S.Record(
|
|
2906
|
+
S.String,
|
|
2907
|
+
S.String,
|
|
2908
|
+
) as any as S.Schema<RequestComplexityRouterConfigTierLabelsMap>;
|
|
2909
|
+
|
|
2910
|
+
export type ComplexityTierModelLitellmParamsMap = {
|
|
2911
|
+
[key: string]: unknown | undefined;
|
|
2912
|
+
};
|
|
2913
|
+
export const ComplexityTierModelLitellmParamsMap = /*@__PURE__*/ S.Record(
|
|
2914
|
+
S.String,
|
|
2915
|
+
S.Unknown,
|
|
2916
|
+
) as any as S.Schema<ComplexityTierModelLitellmParamsMap>;
|
|
2917
|
+
|
|
2918
|
+
export interface ComplexityTierModel {
|
|
2919
|
+
litellm_params?: ComplexityTierModelLitellmParamsMap;
|
|
2920
|
+
model_name: string;
|
|
2921
|
+
}
|
|
2922
|
+
export const ComplexityTierModel = /*@__PURE__*/ S.suspend(() =>
|
|
2923
|
+
S.Struct({
|
|
2924
|
+
litellm_params: S.optional(ComplexityTierModelLitellmParamsMap),
|
|
2925
|
+
model_name: S.String,
|
|
2926
|
+
}),
|
|
2927
|
+
).annotate({
|
|
2928
|
+
identifier: "ComplexityTierModel",
|
|
2929
|
+
}) as any as S.Schema<ComplexityTierModel>;
|
|
2930
|
+
|
|
2931
|
+
export type RequestComplexityRouterConfigTierModelConfigsValueList =
|
|
2932
|
+
Array<ComplexityTierModel>;
|
|
2933
|
+
export const RequestComplexityRouterConfigTierModelConfigsValueList =
|
|
2934
|
+
/*@__PURE__*/ S.Array(
|
|
2935
|
+
ComplexityTierModel,
|
|
2936
|
+
) as any as S.Schema<RequestComplexityRouterConfigTierModelConfigsValueList>;
|
|
2937
|
+
|
|
2938
|
+
export type RequestComplexityRouterConfigTierModelConfigsMap = {
|
|
2939
|
+
[key: string]:
|
|
2940
|
+
| RequestComplexityRouterConfigTierModelConfigsValueList
|
|
2941
|
+
| undefined;
|
|
2942
|
+
};
|
|
2943
|
+
export const RequestComplexityRouterConfigTierModelConfigsMap =
|
|
2944
|
+
/*@__PURE__*/ S.Record(
|
|
2945
|
+
S.String,
|
|
2946
|
+
RequestComplexityRouterConfigTierModelConfigsValueList,
|
|
2947
|
+
) as any as S.Schema<RequestComplexityRouterConfigTierModelConfigsMap>;
|
|
2948
|
+
|
|
2949
|
+
export type RequestComplexityRouterConfigTiersValueCase1List = Array<string>;
|
|
2950
|
+
export const RequestComplexityRouterConfigTiersValueCase1List =
|
|
2951
|
+
/*@__PURE__*/ S.Array(
|
|
2952
|
+
S.String,
|
|
2953
|
+
) as any as S.Schema<RequestComplexityRouterConfigTiersValueCase1List>;
|
|
2954
|
+
|
|
2955
|
+
export type RequestComplexityRouterConfigTiersValue =
|
|
2956
|
+
| string
|
|
2957
|
+
| RequestComplexityRouterConfigTiersValueCase1List;
|
|
2958
|
+
export const RequestComplexityRouterConfigTiersValue =
|
|
2959
|
+
S.Unknown as any as S.Schema<RequestComplexityRouterConfigTiersValue>;
|
|
2960
|
+
|
|
2961
|
+
/** Mapping of complexity tiers to a model or model pool. A list is randomly picked from when adaptive=False, and used as a soft-floor home pool when adaptive=True */
|
|
2962
|
+
export type RequestComplexityRouterConfigTiersMap = {
|
|
2963
|
+
[key: string]: RequestComplexityRouterConfigTiersValue | undefined;
|
|
2964
|
+
};
|
|
2965
|
+
export const RequestComplexityRouterConfigTiersMap = /*@__PURE__*/ S.Record(
|
|
2966
|
+
S.String,
|
|
2967
|
+
RequestComplexityRouterConfigTiersValue,
|
|
2968
|
+
) as any as S.Schema<RequestComplexityRouterConfigTiersMap>;
|
|
2969
|
+
|
|
2970
|
+
/** Token count thresholds for simple/complex classification */
|
|
2971
|
+
export type RequestComplexityRouterConfigTokenThresholdsMap = {
|
|
2972
|
+
[key: string]: number | undefined;
|
|
2973
|
+
};
|
|
2974
|
+
export const RequestComplexityRouterConfigTokenThresholdsMap =
|
|
2975
|
+
/*@__PURE__*/ S.Record(
|
|
2976
|
+
S.String,
|
|
2977
|
+
S.Number,
|
|
2978
|
+
) as any as S.Schema<RequestComplexityRouterConfigTokenThresholdsMap>;
|
|
2979
|
+
|
|
2980
|
+
/** The part of a complexity-router config a request can carry. `plugins` holds live RoutingPlugin objects, which no JSON body can express and which have no OpenAPI schema, so it is closed off here rather than left as an arbitrary-type field. */
|
|
2981
|
+
export interface RequestComplexityRouterConfig {
|
|
2982
|
+
/** Enable adaptive bandit selection with soft complexity floors */
|
|
2983
|
+
adaptive?: boolean;
|
|
2984
|
+
/** When adaptive=True: 'all' scores every pool model with a tier-distance penalty (soft floors); 'classified_tier' Thompson-samples only inside the classified tier's pool */
|
|
2985
|
+
adaptive_eligible?:
|
|
2986
|
+
| RequestComplexityRouterConfigAdaptiveEligible
|
|
2987
|
+
| (string & {});
|
|
2988
|
+
/** Quality vs cost weights for adaptive selection (used when adaptive=True) */
|
|
2989
|
+
adaptive_weights?: AdaptiveRouterWeights;
|
|
2990
|
+
/** Replaces the opening instructions of the LLM classifier rubric (the judging-criteria prose) for a custom tier set. The per-tier bullets and the trust-boundary paragraph telling the classifier to ignore tier requests embedded in quoted caller text are always appended after it and cannot be overridden. Requires tier_definitions; a built-in-tier router customizes its prompt via classifier_llm_config.system_prompt or classification_rubric instead. */
|
|
2991
|
+
classification_prompt?: string | null;
|
|
2992
|
+
/** Maximum characters of prior-turn text quoted to the LLM classifier, across the whole context window, per classification call. Turns are taken newest first and quoted whole while they fit, so a conversation small enough to quote entirely is never cut; once the budget runs out the older turns are dropped whole and only the turn straddling the boundary is truncated, into whatever space is left. The current ask and the caller's system prompt sit outside this budget and are always sent in full, as does the numbering each quoted turn carries. A budget under 120 leaves no room to quote a turn and suppresses the block; set classifier_context_window_size to 0 to turn context off deliberately. Only applies when classifier_type is 'llm'. */
|
|
2993
|
+
classifier_context_budget_chars?: number;
|
|
2994
|
+
/** Include assistant turns in the classifier context window, so difficulty stated by the model rather than by the user stays visible: a plan the assistant calls complex, which the user approves with 'yes', is classified on the work being approved instead of on the word 'yes'. When enabled, classifier_context_window_size counts the last N turns of the conversation across both roles rather than the last N user turns, and assistant text is sent to the classifier model, which may be a different deployment or provider than the routed completion model. Assistant replies spend classifier_context_budget_chars alongside user turns, so raise it if the oldest turns stop being quoted once replies join the window. Off by default because enabling it shifts tier decisions, and therefore spend, for an already-deployed router. Only applies when classifier_type is 'llm'. */
|
|
2995
|
+
classifier_context_include_assistant_turns?: boolean;
|
|
2996
|
+
/** Optional cap on each individual prior turn's text, applied before classifier_context_budget_chars bounds the block. Unset by default, so one long turn may spend the whole budget, which is usually what a follow-up needs; set it when no single turn should dominate the context the classifier sees. A capped turn keeps its opening and its ending with the middle elided. Only applies when classifier_type is 'llm'. */
|
|
2997
|
+
classifier_context_per_turn_chars?: number | null;
|
|
2998
|
+
/** Number of prior user turns (tool output and harness reminders excluded) to include as context in the LLM classifier prompt, so a follow-up like 'now do the same for the streaming path' is classified against what it refers to. Counts turns of both roles when classifier_context_include_assistant_turns is enabled. These turns are sent to the classifier model, which may be a different deployment or provider than the routed completion model; that call already carries the current user ask and the caller's system prompt in full. Set to 0 to send neither prior turns nor any conversation context beyond the current ask. Only applies when classifier_type is 'llm'. */
|
|
2999
|
+
classifier_context_window_size?: number;
|
|
3000
|
+
/** What classifies the request when the LLM classifier errors, times out, or returns an unparseable response. 'heuristic' runs the local complexity scorer, which is right when the classifier grades complexity too. 'default_model' skips scoring and routes to default_model, which is what a classifier on some other taxonomy wants: a prompt that grades data sensitivity has no use for a complexity score, and scoring one produces a tier unrelated to what the operator configured. Requires default_model when set to 'default_model'. Only applies when classifier_type is 'llm', 'custom', or 'heuristic_first'. */
|
|
3001
|
+
classifier_fallback?:
|
|
3002
|
+
| RequestComplexityRouterConfigClassifierFallback
|
|
3003
|
+
| (string & {});
|
|
3004
|
+
/** Configuration for the LLM classifier; required when classifier_type is 'llm' or 'heuristic_first' */
|
|
3005
|
+
classifier_llm_config?: ClassifierLLMConfig | null;
|
|
3006
|
+
/** Not settable over HTTP; the classifier plugin is a runtime object */
|
|
3007
|
+
classifier_plugin?: unknown | null;
|
|
3008
|
+
/** Timeout budget for the classifier plugin call, in milliseconds. On expiry the fallback path decides the tier. Only applies when classifier_type is 'custom'. */
|
|
3009
|
+
classifier_plugin_timeout_ms?: number;
|
|
3010
|
+
/** Classification strategy: local regex/keyword scoring, an LLM call, a custom classifier plugin, or 'heuristic_first', which scores locally and only pays for the LLM classifier when the local scorer does not confidently land a cheap tier */
|
|
3011
|
+
classifier_type?: RequestComplexityRouterConfigClassifierType | (string & {});
|
|
3012
|
+
/** Keywords indicating code-related content */
|
|
3013
|
+
code_keywords?: RequestComplexityRouterConfigCodeKeywordsList | null;
|
|
3014
|
+
/** Domain-specific technical keywords appended to the effective base list (technical_keywords if set, otherwise DEFAULT_TECHNICAL_KEYWORDS). Order is preserved; duplicates are removed case-insensitively against the base list and within this list. */
|
|
3015
|
+
custom_technical_keywords?: RequestComplexityRouterConfigCustomTechnicalKeywordsList | null;
|
|
3016
|
+
/** Default model to use if tier cannot be determined */
|
|
3017
|
+
default_model?: string | null;
|
|
3018
|
+
/** When True and a session_id is resolvable on the request, pin the deployment chosen inside each routed model group and reuse it whenever the session returns to that group, without pinning which group the session routes to. Independent of session_affinity, which pins the model group instead (and always carries this deployment pin with it): with session_affinity off, every turn is still classified on its own merits while a session that escalates to a stronger tier and comes back still lands on the deployment it used before, which is what keeps a provider prompt cache warm. Pins are held per model group, so switching tiers does not disturb the pin left behind in the previous group. On by default because re-shuffling a conversation across deployments of the same model discards that cache for no benefit; set False to keep every turn load-balanced across the group, which is what a deployment set with tight per-deployment rate limits wants. Inert when no session_id is resolvable, since there is nothing to key a pin on, and suppressed when plugins are configured, for the same reason session_affinity is. */
|
|
3019
|
+
deployment_affinity?: boolean;
|
|
3020
|
+
/** Weights for each scoring dimension */
|
|
3021
|
+
dimension_weights?: RequestComplexityRouterConfigDimensionWeightsMap;
|
|
3022
|
+
/** Embedding model (LiteLLM model name) used when semantic_keyword_matching is enabled */
|
|
3023
|
+
embedding_model?: string | null;
|
|
3024
|
+
/** Case-sensitive phrases a user can include to force a bump to the next-higher complexity tier when they aren't satisfied with results (they can force a stronger model, but not choose which one). Defaults to ['LITELLM ESCALATE'] when unset; set to an empty list to disable. */
|
|
3025
|
+
escalation_keywords?: RequestComplexityRouterConfigEscalationKeywordsList | null;
|
|
3026
|
+
/** Tier routed to when the LLM classifier fails (timeout, provider error, or an unparseable reply). Required with tier_definitions and must name a defined tier; the heuristic scorer cannot produce custom tiers, so this replaces the heuristic fallback for custom tier sets. */
|
|
3027
|
+
fallback_tier?: string | null;
|
|
3028
|
+
/** The highest tier the local scorer may decide on its own; required when classifier_type is 'heuristic_first' and rejected otherwise. A request whose heuristic tier is at or below this one skips the LLM classifier and routes straight to that heuristic tier, so the classifier call is only paid for on traffic the scorer could not place cheaply. The scorer must also have produced at least one signal: a prompt where no dimension fired scores 0.0 and would otherwise land SIMPLE by default rather than by evidence, which is how a chained router would silently send unclassified traffic to the cheapest model. Names a built-in tier, and may not name the highest one, since that would make the LLM classifier unreachable. */
|
|
3029
|
+
heuristic_first_max_tier?: string | null;
|
|
3030
|
+
/** Additional case-sensitive literal sentinels that mark a request as client housekeeping, on top of the built-in conversation-title ones. For clients whose wording the built-ins don't cover, or after a client release changes its strings. */
|
|
3031
|
+
housekeeping_patterns?: RequestComplexityRouterConfigHousekeepingPatternsList | null;
|
|
3032
|
+
/** Rules that force a specific tier when their keywords match the prompt */
|
|
3033
|
+
keyword_tier_rules?: RequestComplexityRouterConfigKeywordTierRulesList | null;
|
|
3034
|
+
/** Minimum cosine similarity for a semantic keyword match */
|
|
3035
|
+
match_threshold?: number;
|
|
3036
|
+
/** When set, requests carrying a coding-agent plan-mode sentinel (Claude Code plan mode, VS Code Copilot Plan mode, Copilot CLI's exit_plan_mode tool) are routed to at least this tier: the classified tier still wins when it is higher, and the floor also overrides a session-affinity pin to a lower tier for exactly the turns carrying the sentinel, without rewriting the pin -- the first turn after plan mode exits routes as if plan mode had never happened. Names a built-in tier, or with tier_definitions set, one of the defined tier names (list order is ascending severity, same as keyword_tier_rules). Unset disables detection entirely. The sentinels ride in client-injected prompt text, so a caller who pastes one can spend up to this tier's models -- never down, and never outside the configured pools. */
|
|
3037
|
+
plan_mode_min_tier?: string | null;
|
|
3038
|
+
/** Additional case-sensitive literal sentinels that mark a request as plan mode, on top of the built-in Claude Code and Copilot ones. For clients whose plan-mode wording the built-ins don't cover, or after a client release changes its strings. */
|
|
3039
|
+
plan_mode_patterns?: RequestComplexityRouterConfigPlanModePatternsList | null;
|
|
3040
|
+
/** Not settable over HTTP; routing plugins are runtime objects */
|
|
3041
|
+
plugins?: unknown | null;
|
|
3042
|
+
/** Keywords indicating reasoning-required content */
|
|
3043
|
+
reasoning_keywords?: RequestComplexityRouterConfigReasoningKeywordsList | null;
|
|
3044
|
+
/** Minimum weighted score a request must reach before 2+ reasoning markers may promote it to the reasoning tier. Unset tracks tier_boundaries.simple_medium, so the override never rescues a request the scorer placed in the cheapest tier; 0 restores the unconditional override */
|
|
3045
|
+
reasoning_override_min_score?: number | null;
|
|
3046
|
+
/** Override the delimiter pairs used to recognize and strip harness-injected reminder blocks before classification. A harness that wraps injected context differently per agent type (main, subagent, cron) lists every pair it emits. Replaces, rather than adds to, the built-in default of ('<system-reminder>', '</system-reminder>'), so a harness that also emits that pair lists it too. Matching is case-insensitive. */
|
|
3047
|
+
reminder_markers?: RequestComplexityRouterConfigReminderMarkersList | null;
|
|
3048
|
+
/** Return the resolved raw model name in the response model field instead of the client-requested complexity-router alias */
|
|
3049
|
+
return_raw_model_name?: boolean;
|
|
3050
|
+
/** Route a coding agent's own housekeeping calls to the cheapest configured tier without classifying them. A client names the conversation by quoting the whole session and asking for a title, so the ask reads as the session's engineering work and lands on the most expensive tier, which is the reverse of what the call is worth. Detection is a literal match against client-owned sentinels on the newest ask only, so it cannot fire on an earlier turn, and it never lowers what anyone else asked for: a keyword_tier_rule or a session pin still decides instead, and an escalation keyword or the plan-mode floor still raises the tier from here. Only the classifier is displaced, and its call is skipped, so a matched request costs nothing to route. Set false to classify these calls like any other. */
|
|
3051
|
+
route_housekeeping_to_cheapest_tier?: boolean;
|
|
3052
|
+
/** Match keyword_tier_rules by embedding similarity instead of literal text */
|
|
3053
|
+
semantic_keyword_matching?: boolean;
|
|
3054
|
+
/** When True and a session_id is resolvable on the request, pin the model chosen on the session's first turn and reuse it for every later turn, skipping re-classification. Off by default so every turn is classified on its own merits and routed to the cheapest adequate tier. Set True to keep a multi-turn session on one model, which preserves provider prompt caches and avoids cross-model conversation-history errors. Always implies the deployment pin regardless of deployment_affinity: the session sticks to one deployment of the pinned model, since freezing the model while re-shuffling its deployments would still go cache-cold. */
|
|
3055
|
+
session_affinity?: boolean;
|
|
3056
|
+
/** TTL for the session affinity pin; refreshed on every cache hit. Bounds both the session_affinity model pin and the deployment_affinity deployment pin, so it measures idle time for the session's routing decisions rather than total session length */
|
|
3057
|
+
session_affinity_ttl_seconds?: number;
|
|
3058
|
+
/** Keywords indicating simple/basic queries */
|
|
3059
|
+
simple_keywords?: RequestComplexityRouterConfigSimpleKeywordsList | null;
|
|
3060
|
+
/** Keywords indicating technical content */
|
|
3061
|
+
technical_keywords?: RequestComplexityRouterConfigTechnicalKeywordsList | null;
|
|
3062
|
+
/** Score boundaries between tiers. These keys (simple_medium, medium_complex, complex_reasoning) name the gaps between the default tier names and are not renameable by tier_labels; they are scorer knobs persisted by name on every routing decision */
|
|
3063
|
+
tier_boundaries?: RequestComplexityRouterConfigTierBoundariesMap;
|
|
3064
|
+
/** Operator-defined tier set replacing the built-in SIMPLE/MEDIUM/COMPLEX/REASONING. Each entry's name becomes a value the LLM classifier can return and its description becomes that tier's rubric bullet; entries named after a built-in tier may omit the description and inherit the built-in criteria. List order is ascending severity and decides which tier wins when several keyword_tier_rules match. Requires classifier_type 'llm' or 'custom', a fallback_tier, and `tiers` keys matching the defined names exactly. Escalation, adaptive selection, session affinity, plugins, tier_labels, and the calibration-example rubric presets are unavailable with a custom tier set: the first four are built on the built-in tier ladder, and the last two rename or exemplify tiers the set replaces. */
|
|
3065
|
+
tier_definitions?: RequestComplexityRouterConfigTierDefinitionsList | null;
|
|
3066
|
+
/** Score penalty per tier-step away from the classified tier when adaptive=True */
|
|
3067
|
+
tier_distance_penalty?: number;
|
|
3068
|
+
/** Display names for the complexity tiers, so a deployment can use its own vocabulary (e.g. Cheap/Standard/Premium/Deep) in the dashboard, spend logs, and the LLM classifier rubric. Purely operator-facing: config keys stay canonical (tiers, keyword_tier_rules[].tier, tier_boundaries), API callers never see these names, and the heuristic scorer never reads them. Unlisted tiers keep their canonical name. Partial maps are allowed. */
|
|
3069
|
+
tier_labels?: RequestComplexityRouterConfigTierLabelsMap;
|
|
3070
|
+
tier_model_configs?: RequestComplexityRouterConfigTierModelConfigsMap;
|
|
3071
|
+
/** Mapping of complexity tiers to a model or model pool. A list is randomly picked from when adaptive=False, and used as a soft-floor home pool when adaptive=True */
|
|
3072
|
+
tiers?: RequestComplexityRouterConfigTiersMap;
|
|
3073
|
+
/** Token count thresholds for simple/complex classification */
|
|
3074
|
+
token_thresholds?: RequestComplexityRouterConfigTokenThresholdsMap;
|
|
3075
|
+
}
|
|
3076
|
+
export const RequestComplexityRouterConfig = /*@__PURE__*/ S.suspend(() =>
|
|
3077
|
+
S.Struct({
|
|
3078
|
+
adaptive: S.optional(S.Boolean),
|
|
3079
|
+
adaptive_eligible: S.optional(
|
|
3080
|
+
RequestComplexityRouterConfigAdaptiveEligible,
|
|
3081
|
+
),
|
|
3082
|
+
adaptive_weights: S.optional(AdaptiveRouterWeights),
|
|
3083
|
+
classification_prompt: S.optional(S.NullOr(S.String)),
|
|
3084
|
+
classifier_context_budget_chars: S.optional(S.Number),
|
|
3085
|
+
classifier_context_include_assistant_turns: S.optional(S.Boolean),
|
|
3086
|
+
classifier_context_per_turn_chars: S.optional(S.NullOr(S.Number)),
|
|
3087
|
+
classifier_context_window_size: S.optional(S.Number),
|
|
3088
|
+
classifier_fallback: S.optional(
|
|
3089
|
+
RequestComplexityRouterConfigClassifierFallback,
|
|
3090
|
+
),
|
|
3091
|
+
classifier_llm_config: S.optional(S.NullOr(ClassifierLLMConfig)),
|
|
3092
|
+
classifier_plugin: S.optional(S.NullOr(S.Unknown)),
|
|
3093
|
+
classifier_plugin_timeout_ms: S.optional(S.Number),
|
|
3094
|
+
classifier_type: S.optional(RequestComplexityRouterConfigClassifierType),
|
|
3095
|
+
code_keywords: S.optional(
|
|
3096
|
+
S.NullOr(RequestComplexityRouterConfigCodeKeywordsList),
|
|
3097
|
+
),
|
|
3098
|
+
custom_technical_keywords: S.optional(
|
|
3099
|
+
S.NullOr(RequestComplexityRouterConfigCustomTechnicalKeywordsList),
|
|
3100
|
+
),
|
|
3101
|
+
default_model: S.optional(S.NullOr(S.String)),
|
|
3102
|
+
deployment_affinity: S.optional(S.Boolean),
|
|
3103
|
+
dimension_weights: S.optional(
|
|
3104
|
+
RequestComplexityRouterConfigDimensionWeightsMap,
|
|
3105
|
+
),
|
|
3106
|
+
embedding_model: S.optional(S.NullOr(S.String)),
|
|
3107
|
+
escalation_keywords: S.optional(
|
|
3108
|
+
S.NullOr(RequestComplexityRouterConfigEscalationKeywordsList),
|
|
3109
|
+
),
|
|
3110
|
+
fallback_tier: S.optional(S.NullOr(S.String)),
|
|
3111
|
+
heuristic_first_max_tier: S.optional(S.NullOr(S.String)),
|
|
3112
|
+
housekeeping_patterns: S.optional(
|
|
3113
|
+
S.NullOr(RequestComplexityRouterConfigHousekeepingPatternsList),
|
|
3114
|
+
),
|
|
3115
|
+
keyword_tier_rules: S.optional(
|
|
3116
|
+
S.NullOr(RequestComplexityRouterConfigKeywordTierRulesList),
|
|
3117
|
+
),
|
|
3118
|
+
match_threshold: S.optional(S.Number),
|
|
3119
|
+
plan_mode_min_tier: S.optional(S.NullOr(S.String)),
|
|
3120
|
+
plan_mode_patterns: S.optional(
|
|
3121
|
+
S.NullOr(RequestComplexityRouterConfigPlanModePatternsList),
|
|
3122
|
+
),
|
|
3123
|
+
plugins: S.optional(S.NullOr(S.Unknown)),
|
|
3124
|
+
reasoning_keywords: S.optional(
|
|
3125
|
+
S.NullOr(RequestComplexityRouterConfigReasoningKeywordsList),
|
|
3126
|
+
),
|
|
3127
|
+
reasoning_override_min_score: S.optional(S.NullOr(S.Number)),
|
|
3128
|
+
reminder_markers: S.optional(
|
|
3129
|
+
S.NullOr(RequestComplexityRouterConfigReminderMarkersList),
|
|
3130
|
+
),
|
|
3131
|
+
return_raw_model_name: S.optional(S.Boolean),
|
|
3132
|
+
route_housekeeping_to_cheapest_tier: S.optional(S.Boolean),
|
|
3133
|
+
semantic_keyword_matching: S.optional(S.Boolean),
|
|
3134
|
+
session_affinity: S.optional(S.Boolean),
|
|
3135
|
+
session_affinity_ttl_seconds: S.optional(S.Number),
|
|
3136
|
+
simple_keywords: S.optional(
|
|
3137
|
+
S.NullOr(RequestComplexityRouterConfigSimpleKeywordsList),
|
|
3138
|
+
),
|
|
3139
|
+
technical_keywords: S.optional(
|
|
3140
|
+
S.NullOr(RequestComplexityRouterConfigTechnicalKeywordsList),
|
|
3141
|
+
),
|
|
3142
|
+
tier_boundaries: S.optional(RequestComplexityRouterConfigTierBoundariesMap),
|
|
3143
|
+
tier_definitions: S.optional(
|
|
3144
|
+
S.NullOr(RequestComplexityRouterConfigTierDefinitionsList),
|
|
3145
|
+
),
|
|
3146
|
+
tier_distance_penalty: S.optional(S.Number),
|
|
3147
|
+
tier_labels: S.optional(RequestComplexityRouterConfigTierLabelsMap),
|
|
3148
|
+
tier_model_configs: S.optional(
|
|
3149
|
+
RequestComplexityRouterConfigTierModelConfigsMap,
|
|
3150
|
+
),
|
|
3151
|
+
tiers: S.optional(RequestComplexityRouterConfigTiersMap),
|
|
3152
|
+
token_thresholds: S.optional(
|
|
3153
|
+
RequestComplexityRouterConfigTokenThresholdsMap,
|
|
3154
|
+
),
|
|
3155
|
+
}),
|
|
3156
|
+
).annotate({
|
|
3157
|
+
identifier: "RequestComplexityRouterConfig",
|
|
3158
|
+
}) as any as S.Schema<RequestComplexityRouterConfig>;
|
|
3159
|
+
|
|
3160
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap =
|
|
3161
|
+
{ [key: string]: unknown | undefined };
|
|
3162
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap =
|
|
3163
|
+
/*@__PURE__*/ S.Record(
|
|
3164
|
+
S.String,
|
|
3165
|
+
S.Unknown,
|
|
3166
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap>;
|
|
3167
|
+
|
|
3168
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList =
|
|
3169
|
+
Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap>;
|
|
3170
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList =
|
|
3171
|
+
/*@__PURE__*/ S.Array(
|
|
3172
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap,
|
|
3173
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList>;
|
|
3174
|
+
|
|
3175
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap =
|
|
3176
|
+
{ [key: string]: unknown | undefined };
|
|
3177
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap =
|
|
3178
|
+
/*@__PURE__*/ S.Record(
|
|
3179
|
+
S.String,
|
|
3180
|
+
S.Unknown,
|
|
3181
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap>;
|
|
3182
|
+
|
|
3183
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List =
|
|
3184
|
+
Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap>;
|
|
3185
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List =
|
|
3186
|
+
/*@__PURE__*/ S.Array(
|
|
3187
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap,
|
|
3188
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List>;
|
|
3189
|
+
|
|
3190
|
+
/** The top-level system prompt an Anthropic /v1/messages body carries beside its messages */
|
|
3191
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem =
|
|
3192
|
+
| string
|
|
3193
|
+
| PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List;
|
|
3194
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem =
|
|
3195
|
+
S.Unknown as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem>;
|
|
3196
|
+
|
|
3197
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap =
|
|
3198
|
+
{ [key: string]: unknown | undefined };
|
|
3199
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap =
|
|
3200
|
+
/*@__PURE__*/ S.Record(
|
|
3201
|
+
S.String,
|
|
3202
|
+
S.Unknown,
|
|
3203
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap>;
|
|
3204
|
+
|
|
3205
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList =
|
|
3206
|
+
Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap>;
|
|
3207
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList =
|
|
3208
|
+
/*@__PURE__*/ S.Array(
|
|
3209
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap,
|
|
3210
|
+
) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList>;
|
|
3211
|
+
|
|
3212
|
+
export interface PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest {
|
|
3213
|
+
/** The complexity router config to route against, in the shape /model/new accepts */
|
|
3214
|
+
complexity_router_config: RequestComplexityRouterConfig;
|
|
3215
|
+
/** Model to route to when no tier resolves, i.e. complexity_router_default_model */
|
|
3216
|
+
default_model?: string | null;
|
|
3217
|
+
/** The full message list to route, exactly as the serving path would receive it. Mutually exclusive with prompt */
|
|
3218
|
+
messages?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList | null;
|
|
3219
|
+
/** A single ask to route, as an end user would send it. Mutually exclusive with messages */
|
|
3220
|
+
prompt?: string | null;
|
|
3221
|
+
/** Name reported as the router in the routing decision. Display only */
|
|
3222
|
+
router_name?: string;
|
|
3223
|
+
/** The top-level system prompt an Anthropic /v1/messages body carries beside its messages */
|
|
3224
|
+
system?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem | null;
|
|
3225
|
+
/** Team the router is being created for. Required for a team admin, who may only test their own team's routers */
|
|
3226
|
+
team_id?: string | null;
|
|
3227
|
+
/** The tool definitions the request advertises, which decide whether the plan-mode floor applies */
|
|
3228
|
+
tools?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList | null;
|
|
3229
|
+
}
|
|
3230
|
+
export const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest =
|
|
3231
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3232
|
+
S.Struct({
|
|
3233
|
+
complexity_router_config: RequestComplexityRouterConfig,
|
|
3234
|
+
default_model: S.optional(S.NullOr(S.String)),
|
|
3235
|
+
messages: S.optional(
|
|
3236
|
+
S.NullOr(
|
|
3237
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList,
|
|
3238
|
+
),
|
|
3239
|
+
),
|
|
3240
|
+
prompt: S.optional(S.NullOr(S.String)),
|
|
3241
|
+
router_name: S.optional(S.String),
|
|
3242
|
+
system: S.optional(
|
|
3243
|
+
S.NullOr(
|
|
3244
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem,
|
|
3245
|
+
),
|
|
3246
|
+
),
|
|
3247
|
+
team_id: S.optional(S.NullOr(S.String)),
|
|
3248
|
+
tools: S.optional(
|
|
3249
|
+
S.NullOr(
|
|
3250
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList,
|
|
3251
|
+
),
|
|
3252
|
+
),
|
|
3253
|
+
}).pipe(
|
|
3254
|
+
T.Http({ method: "POST", uri: "/auto_router/test_routing", code: 200 }),
|
|
3255
|
+
),
|
|
3256
|
+
).annotate({
|
|
3257
|
+
identifier: "PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest",
|
|
3258
|
+
}) as any as S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest>;
|
|
3259
|
+
|
|
3260
|
+
export type StandardLoggingRoutingDecisionCause =
|
|
3261
|
+
| "heuristic_scorer"
|
|
3262
|
+
| "reasoning_override"
|
|
3263
|
+
| "llm_classifier"
|
|
3264
|
+
| "heuristic_first_short_circuit"
|
|
3265
|
+
| "classifier_plugin"
|
|
3266
|
+
| "classifier_fallback"
|
|
3267
|
+
| "default_model_fallback"
|
|
3268
|
+
| "literal_keyword_match"
|
|
3269
|
+
| "semantic_keyword_match"
|
|
3270
|
+
| "plan_mode"
|
|
3271
|
+
| "housekeeping"
|
|
3272
|
+
| "session_affinity_pin"
|
|
3273
|
+
| "session_affinity_escalation"
|
|
3274
|
+
| "default_fallback"
|
|
3275
|
+
| "keyword"
|
|
3276
|
+
| "quality_tier"
|
|
3277
|
+
| "bandit";
|
|
3278
|
+
export const StandardLoggingRoutingDecisionCause = S.String;
|
|
3279
|
+
|
|
3280
|
+
export type StandardLoggingRoutingDecisionRouterType =
|
|
3281
|
+
| "complexity"
|
|
3282
|
+
| "adaptive"
|
|
3283
|
+
| "quality";
|
|
3284
|
+
export const StandardLoggingRoutingDecisionRouterType = S.String;
|
|
3285
|
+
|
|
3286
|
+
export type StandardLoggingRoutingDecisionSignalsList = Array<string>;
|
|
3287
|
+
export const StandardLoggingRoutingDecisionSignalsList = /*@__PURE__*/ S.Array(
|
|
3288
|
+
S.String,
|
|
3289
|
+
) as any as S.Schema<StandardLoggingRoutingDecisionSignalsList>;
|
|
3290
|
+
|
|
3291
|
+
/** Snapshot of the complexity scorer's tier boundaries at decision time, so a historical spend log row stays explainable after the router config changes. */
|
|
3292
|
+
export interface StandardLoggingRoutingDecisionTierBoundaries {
|
|
3293
|
+
complex_reasoning: number;
|
|
3294
|
+
medium_complex: number;
|
|
3295
|
+
simple_medium: number;
|
|
3296
|
+
}
|
|
3297
|
+
export const StandardLoggingRoutingDecisionTierBoundaries =
|
|
3298
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3299
|
+
S.Struct({
|
|
3300
|
+
complex_reasoning: S.Number,
|
|
3301
|
+
medium_complex: S.Number,
|
|
3302
|
+
simple_medium: S.Number,
|
|
3303
|
+
}),
|
|
3304
|
+
).annotate({
|
|
3305
|
+
identifier: "StandardLoggingRoutingDecisionTierBoundaries",
|
|
3306
|
+
}) as any as S.Schema<StandardLoggingRoutingDecisionTierBoundaries>;
|
|
3307
|
+
|
|
3308
|
+
export type StandardLoggingRoutingDecisionTierLitellmParamsMap = {
|
|
3309
|
+
[key: string]: unknown | undefined;
|
|
3310
|
+
};
|
|
3311
|
+
export const StandardLoggingRoutingDecisionTierLitellmParamsMap =
|
|
3312
|
+
/*@__PURE__*/ S.Record(
|
|
3313
|
+
S.String,
|
|
3314
|
+
S.Unknown,
|
|
3315
|
+
) as any as S.Schema<StandardLoggingRoutingDecisionTierLitellmParamsMap>;
|
|
3316
|
+
|
|
3317
|
+
/** Per-request provenance for a pre-routing strategy (auto-router) decision. */
|
|
3318
|
+
export interface StandardLoggingRoutingDecision {
|
|
3319
|
+
cause?: StandardLoggingRoutingDecisionCause;
|
|
3320
|
+
classifier_cost?: number;
|
|
3321
|
+
classifier_model?: string;
|
|
3322
|
+
conversation_continuing?: boolean;
|
|
3323
|
+
escalated?: boolean;
|
|
3324
|
+
escalation_keyword?: string;
|
|
3325
|
+
matched_keyword?: string;
|
|
3326
|
+
reasoning_override_min_score?: number;
|
|
3327
|
+
request_type?: string;
|
|
3328
|
+
routed_model?: string;
|
|
3329
|
+
router_model_name?: string;
|
|
3330
|
+
router_type?: StandardLoggingRoutingDecisionRouterType;
|
|
3331
|
+
savings_baseline_deployment_id?: string;
|
|
3332
|
+
savings_baseline_model?: string;
|
|
3333
|
+
score?: number;
|
|
3334
|
+
signals?: StandardLoggingRoutingDecisionSignalsList;
|
|
3335
|
+
tier?: string;
|
|
3336
|
+
tier_boundaries?: StandardLoggingRoutingDecisionTierBoundaries;
|
|
3337
|
+
tier_label?: string;
|
|
3338
|
+
tier_litellm_params?: StandardLoggingRoutingDecisionTierLitellmParamsMap;
|
|
3339
|
+
}
|
|
3340
|
+
export const StandardLoggingRoutingDecision = /*@__PURE__*/ S.suspend(() =>
|
|
3341
|
+
S.Struct({
|
|
3342
|
+
cause: S.optional(StandardLoggingRoutingDecisionCause),
|
|
3343
|
+
classifier_cost: S.optional(S.Number),
|
|
3344
|
+
classifier_model: S.optional(S.String),
|
|
3345
|
+
conversation_continuing: S.optional(S.Boolean),
|
|
3346
|
+
escalated: S.optional(S.Boolean),
|
|
3347
|
+
escalation_keyword: S.optional(S.String),
|
|
3348
|
+
matched_keyword: S.optional(S.String),
|
|
3349
|
+
reasoning_override_min_score: S.optional(S.Number),
|
|
3350
|
+
request_type: S.optional(S.String),
|
|
3351
|
+
routed_model: S.optional(S.String),
|
|
3352
|
+
router_model_name: S.optional(S.String),
|
|
3353
|
+
router_type: S.optional(StandardLoggingRoutingDecisionRouterType),
|
|
3354
|
+
savings_baseline_deployment_id: S.optional(S.String),
|
|
3355
|
+
savings_baseline_model: S.optional(S.String),
|
|
3356
|
+
score: S.optional(S.Number),
|
|
3357
|
+
signals: S.optional(StandardLoggingRoutingDecisionSignalsList),
|
|
3358
|
+
tier: S.optional(S.String),
|
|
3359
|
+
tier_boundaries: S.optional(StandardLoggingRoutingDecisionTierBoundaries),
|
|
3360
|
+
tier_label: S.optional(S.String),
|
|
3361
|
+
tier_litellm_params: S.optional(
|
|
3362
|
+
StandardLoggingRoutingDecisionTierLitellmParamsMap,
|
|
3363
|
+
),
|
|
3364
|
+
}),
|
|
3365
|
+
).annotate({
|
|
3366
|
+
identifier: "StandardLoggingRoutingDecision",
|
|
3367
|
+
}) as any as S.Schema<StandardLoggingRoutingDecision>;
|
|
3368
|
+
|
|
3369
|
+
/** Where one prompt would have been routed, and why. */
|
|
3370
|
+
export interface AutoRouterRoutingTestResponse {
|
|
3371
|
+
/** The model group the router picked */
|
|
3372
|
+
routed_model: string;
|
|
3373
|
+
/** Whether routed_model is a model group available to the caller, scoped to team_id when given. Never confirms models the caller could not use */
|
|
3374
|
+
routed_model_configured: boolean;
|
|
3375
|
+
/** The decision record this request would have written to its log row */
|
|
3376
|
+
routing_decision: StandardLoggingRoutingDecision;
|
|
3377
|
+
}
|
|
3378
|
+
export const AutoRouterRoutingTestResponse = /*@__PURE__*/ S.suspend(() =>
|
|
3379
|
+
S.Struct({
|
|
3380
|
+
routed_model: S.String,
|
|
3381
|
+
routed_model_configured: S.Boolean,
|
|
3382
|
+
routing_decision: StandardLoggingRoutingDecision,
|
|
3383
|
+
}),
|
|
3384
|
+
).annotate({
|
|
3385
|
+
identifier: "AutoRouterRoutingTestResponse",
|
|
3386
|
+
}) as any as S.Schema<AutoRouterRoutingTestResponse>;
|
|
3387
|
+
|
|
3388
|
+
export interface PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest {}
|
|
3389
|
+
export const PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest =
|
|
3390
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3391
|
+
S.Struct({}).pipe(
|
|
3392
|
+
T.Http({
|
|
3393
|
+
method: "POST",
|
|
3394
|
+
uri: "/reload/anthropic_beta_headers",
|
|
3395
|
+
code: 200,
|
|
3396
|
+
}),
|
|
3397
|
+
),
|
|
3398
|
+
).annotate({
|
|
3399
|
+
identifier:
|
|
3400
|
+
"PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest",
|
|
3401
|
+
}) as any as S.Schema<PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest>;
|
|
3402
|
+
|
|
3403
|
+
export interface PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse {
|
|
3404
|
+
body: unknown;
|
|
3405
|
+
}
|
|
3406
|
+
export const PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse =
|
|
3407
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3408
|
+
S.Struct({
|
|
3409
|
+
body: S.Unknown,
|
|
3410
|
+
}),
|
|
3411
|
+
).annotate({
|
|
3412
|
+
identifier:
|
|
3413
|
+
"PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse",
|
|
3414
|
+
}) as any as S.Schema<PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse>;
|
|
3415
|
+
|
|
3416
|
+
export interface PostReloadModelCostMapReloadModelCostMapRequest {}
|
|
3417
|
+
export const PostReloadModelCostMapReloadModelCostMapRequest =
|
|
3418
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3419
|
+
S.Struct({}).pipe(
|
|
3420
|
+
T.Http({ method: "POST", uri: "/reload/model_cost_map", code: 200 }),
|
|
3421
|
+
),
|
|
3422
|
+
).annotate({
|
|
3423
|
+
identifier: "PostReloadModelCostMapReloadModelCostMapRequest",
|
|
3424
|
+
}) as any as S.Schema<PostReloadModelCostMapReloadModelCostMapRequest>;
|
|
3425
|
+
|
|
3426
|
+
export interface PostReloadModelCostMapReloadModelCostMapResponse {
|
|
3427
|
+
body: unknown;
|
|
3428
|
+
}
|
|
3429
|
+
export const PostReloadModelCostMapReloadModelCostMapResponse =
|
|
3430
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3431
|
+
S.Struct({
|
|
3432
|
+
body: S.Unknown,
|
|
3433
|
+
}),
|
|
3434
|
+
).annotate({
|
|
3435
|
+
identifier: "PostReloadModelCostMapReloadModelCostMapResponse",
|
|
3436
|
+
}) as any as S.Schema<PostReloadModelCostMapReloadModelCostMapResponse>;
|
|
3437
|
+
|
|
3438
|
+
export interface PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest {
|
|
3439
|
+
hours: number;
|
|
3440
|
+
}
|
|
3441
|
+
export const PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest =
|
|
3442
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3443
|
+
S.Struct({
|
|
3444
|
+
hours: S.Number.pipe(T.Query()),
|
|
3445
|
+
}).pipe(
|
|
3446
|
+
T.Http({
|
|
3447
|
+
method: "POST",
|
|
3448
|
+
uri: "/schedule/anthropic_beta_headers_reload",
|
|
3449
|
+
code: 200,
|
|
3450
|
+
}),
|
|
3451
|
+
),
|
|
3452
|
+
).annotate({
|
|
3453
|
+
identifier:
|
|
3454
|
+
"PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest",
|
|
3455
|
+
}) as any as S.Schema<PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest>;
|
|
3456
|
+
|
|
3457
|
+
export interface PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse {
|
|
3458
|
+
body: unknown;
|
|
3459
|
+
}
|
|
3460
|
+
export const PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse =
|
|
3461
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3462
|
+
S.Struct({
|
|
3463
|
+
body: S.Unknown,
|
|
3464
|
+
}),
|
|
3465
|
+
).annotate({
|
|
3466
|
+
identifier:
|
|
3467
|
+
"PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse",
|
|
3468
|
+
}) as any as S.Schema<PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse>;
|
|
3469
|
+
|
|
3470
|
+
export interface PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest {
|
|
3471
|
+
hours: number;
|
|
3472
|
+
}
|
|
3473
|
+
export const PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest =
|
|
3474
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3475
|
+
S.Struct({
|
|
3476
|
+
hours: S.Number.pipe(T.Query()),
|
|
3477
|
+
}).pipe(
|
|
3478
|
+
T.Http({
|
|
3479
|
+
method: "POST",
|
|
3480
|
+
uri: "/schedule/model_cost_map_reload",
|
|
3481
|
+
code: 200,
|
|
3482
|
+
}),
|
|
3483
|
+
),
|
|
3484
|
+
).annotate({
|
|
3485
|
+
identifier:
|
|
3486
|
+
"PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest",
|
|
3487
|
+
}) as any as S.Schema<PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest>;
|
|
3488
|
+
|
|
3489
|
+
export interface PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse {
|
|
3490
|
+
body: unknown;
|
|
3491
|
+
}
|
|
3492
|
+
export const PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse =
|
|
3493
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3494
|
+
S.Struct({
|
|
3495
|
+
body: S.Unknown,
|
|
3496
|
+
}),
|
|
3497
|
+
).annotate({
|
|
3498
|
+
identifier:
|
|
3499
|
+
"PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse",
|
|
3500
|
+
}) as any as S.Schema<PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse>;
|
|
3501
|
+
|
|
3502
|
+
export interface PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest {
|
|
3503
|
+
access_group: string;
|
|
3504
|
+
budget_duration?: string | null;
|
|
3505
|
+
budget_id?: string | null;
|
|
3506
|
+
max_budget?: number | null;
|
|
3507
|
+
soft_budget?: number | null;
|
|
3508
|
+
}
|
|
3509
|
+
export const PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest =
|
|
3510
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3511
|
+
S.Struct({
|
|
3512
|
+
access_group: S.String.pipe(T.Label()),
|
|
3513
|
+
budget_duration: S.optional(S.NullOr(S.String)),
|
|
3514
|
+
budget_id: S.optional(S.NullOr(S.String)),
|
|
3515
|
+
max_budget: S.optional(S.NullOr(S.Number)),
|
|
3516
|
+
soft_budget: S.optional(S.NullOr(S.Number)),
|
|
3517
|
+
}).pipe(
|
|
3518
|
+
T.Http({
|
|
3519
|
+
method: "PUT",
|
|
3520
|
+
uri: "/access_group/{access_group}/budget",
|
|
3521
|
+
code: 200,
|
|
3522
|
+
}),
|
|
3523
|
+
),
|
|
3524
|
+
).annotate({
|
|
3525
|
+
identifier: "PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest",
|
|
3526
|
+
}) as any as S.Schema<PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest>;
|
|
3527
|
+
|
|
3528
|
+
export interface UnblockModelModelUnblockPostRequest {
|
|
3529
|
+
model_id: string;
|
|
3530
|
+
}
|
|
3531
|
+
export const UnblockModelModelUnblockPostRequest = /*@__PURE__*/ S.suspend(() =>
|
|
3532
|
+
S.Struct({
|
|
3533
|
+
model_id: S.String,
|
|
3534
|
+
}).pipe(T.Http({ method: "POST", uri: "/model/unblock", code: 200 })),
|
|
3535
|
+
).annotate({
|
|
3536
|
+
identifier: "UnblockModelModelUnblockPostRequest",
|
|
3537
|
+
}) as any as S.Schema<UnblockModelModelUnblockPostRequest>;
|
|
3538
|
+
|
|
3539
|
+
export interface UnblockModelModelUnblockPostResponse {
|
|
3540
|
+
body: LiteLLMProxyModelTable | null;
|
|
3541
|
+
}
|
|
3542
|
+
export const UnblockModelModelUnblockPostResponse = /*@__PURE__*/ S.suspend(
|
|
3543
|
+
() =>
|
|
3544
|
+
S.Struct({
|
|
3545
|
+
body: S.NullOr(LiteLLMProxyModelTable),
|
|
3546
|
+
}),
|
|
3547
|
+
).annotate({
|
|
3548
|
+
identifier: "UnblockModelModelUnblockPostResponse",
|
|
3549
|
+
}) as any as S.Schema<UnblockModelModelUnblockPostResponse>;
|
|
3550
|
+
|
|
3551
|
+
export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList =
|
|
3552
|
+
Array<string>;
|
|
3553
|
+
export const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList =
|
|
3554
|
+
/*@__PURE__*/ S.Array(
|
|
3555
|
+
S.String,
|
|
3556
|
+
) as any as S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList>;
|
|
3557
|
+
|
|
3558
|
+
export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList =
|
|
3559
|
+
Array<string>;
|
|
3560
|
+
export const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList =
|
|
3561
|
+
/*@__PURE__*/ S.Array(
|
|
3562
|
+
S.String,
|
|
3563
|
+
) as any as S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList>;
|
|
3564
|
+
|
|
3565
|
+
export interface UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest {
|
|
3566
|
+
access_group: string;
|
|
3567
|
+
model_ids?: UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList | null;
|
|
3568
|
+
model_names?: UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList | null;
|
|
3569
|
+
}
|
|
3570
|
+
export const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest =
|
|
3571
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3572
|
+
S.Struct({
|
|
3573
|
+
access_group: S.String.pipe(T.Label()),
|
|
3574
|
+
model_ids: S.optional(
|
|
3575
|
+
S.NullOr(
|
|
3576
|
+
UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList,
|
|
3577
|
+
),
|
|
3578
|
+
),
|
|
3579
|
+
model_names: S.optional(
|
|
3580
|
+
S.NullOr(
|
|
3581
|
+
UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList,
|
|
3582
|
+
),
|
|
3583
|
+
),
|
|
3584
|
+
}).pipe(
|
|
3585
|
+
T.Http({
|
|
3586
|
+
method: "PUT",
|
|
3587
|
+
uri: "/access_group/{access_group}/update",
|
|
3588
|
+
code: 200,
|
|
3589
|
+
}),
|
|
3590
|
+
),
|
|
3591
|
+
).annotate({
|
|
3592
|
+
identifier: "UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest",
|
|
3593
|
+
}) as any as S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest>;
|
|
3594
|
+
|
|
3595
|
+
export interface UpdateModelModelUpdatePostRequest {
|
|
3596
|
+
blocked?: boolean | null;
|
|
3597
|
+
litellm_params?: UpdateLiteLLMParams | null;
|
|
3598
|
+
model_info?: LitellmTypesRouterModelInfo | null;
|
|
3599
|
+
model_name?: string | null;
|
|
3600
|
+
}
|
|
3601
|
+
export const UpdateModelModelUpdatePostRequest = /*@__PURE__*/ S.suspend(() =>
|
|
3602
|
+
S.Struct({
|
|
3603
|
+
blocked: S.optional(S.NullOr(S.Boolean)),
|
|
3604
|
+
litellm_params: S.optional(S.NullOr(UpdateLiteLLMParams)),
|
|
3605
|
+
model_info: S.optional(S.NullOr(LitellmTypesRouterModelInfo)),
|
|
3606
|
+
model_name: S.optional(S.NullOr(S.String)),
|
|
3607
|
+
}).pipe(T.Http({ method: "POST", uri: "/model/update", code: 200 })),
|
|
3608
|
+
).annotate({
|
|
3609
|
+
identifier: "UpdateModelModelUpdatePostRequest",
|
|
3610
|
+
}) as any as S.Schema<UpdateModelModelUpdatePostRequest>;
|
|
3611
|
+
|
|
3612
|
+
export interface UpdateModelModelUpdatePostResponse {
|
|
3613
|
+
body: unknown;
|
|
3614
|
+
}
|
|
3615
|
+
export const UpdateModelModelUpdatePostResponse = /*@__PURE__*/ S.suspend(() =>
|
|
3616
|
+
S.Struct({
|
|
3617
|
+
body: S.Unknown,
|
|
3618
|
+
}),
|
|
3619
|
+
).annotate({
|
|
3620
|
+
identifier: "UpdateModelModelUpdatePostResponse",
|
|
3621
|
+
}) as any as S.Schema<UpdateModelModelUpdatePostResponse>;
|
|
3622
|
+
|
|
3623
|
+
/** List of model group names to make public */
|
|
3624
|
+
export type UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList =
|
|
3625
|
+
Array<string>;
|
|
3626
|
+
export const UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList =
|
|
3627
|
+
/*@__PURE__*/ S.Array(
|
|
3628
|
+
S.String,
|
|
3629
|
+
) as any as S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList>;
|
|
3630
|
+
|
|
3631
|
+
export interface UpdatePublicModelGroupsModelGroupMakePublicPostRequest {
|
|
3632
|
+
/** List of model group names to make public */
|
|
3633
|
+
model_groups: UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList;
|
|
3634
|
+
}
|
|
3635
|
+
export const UpdatePublicModelGroupsModelGroupMakePublicPostRequest =
|
|
3636
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3637
|
+
S.Struct({
|
|
3638
|
+
model_groups:
|
|
3639
|
+
UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList,
|
|
3640
|
+
}).pipe(
|
|
3641
|
+
T.Http({ method: "POST", uri: "/model_group/make_public", code: 200 }),
|
|
3642
|
+
),
|
|
3643
|
+
).annotate({
|
|
3644
|
+
identifier: "UpdatePublicModelGroupsModelGroupMakePublicPostRequest",
|
|
3645
|
+
}) as any as S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostRequest>;
|
|
3646
|
+
|
|
3647
|
+
export interface UpdatePublicModelGroupsModelGroupMakePublicPostResponse {
|
|
3648
|
+
body: unknown;
|
|
3649
|
+
}
|
|
3650
|
+
export const UpdatePublicModelGroupsModelGroupMakePublicPostResponse =
|
|
3651
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3652
|
+
S.Struct({
|
|
3653
|
+
body: S.Unknown,
|
|
3654
|
+
}),
|
|
3655
|
+
).annotate({
|
|
3656
|
+
identifier: "UpdatePublicModelGroupsModelGroupMakePublicPostResponse",
|
|
3657
|
+
}) as any as S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostResponse>;
|
|
3658
|
+
|
|
3659
|
+
export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map =
|
|
3660
|
+
{ [key: string]: unknown | undefined };
|
|
3661
|
+
export const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map =
|
|
3662
|
+
/*@__PURE__*/ S.Record(
|
|
3663
|
+
S.String,
|
|
3664
|
+
S.Unknown,
|
|
3665
|
+
) as any as S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map>;
|
|
3666
|
+
|
|
3667
|
+
export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue =
|
|
3668
|
+
| string
|
|
3669
|
+
| UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map;
|
|
3670
|
+
export const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue =
|
|
3671
|
+
S.Unknown as any as S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue>;
|
|
3672
|
+
|
|
3673
|
+
export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap =
|
|
3674
|
+
{
|
|
3675
|
+
[key: string]:
|
|
3676
|
+
| UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue
|
|
3677
|
+
| undefined;
|
|
3678
|
+
};
|
|
3679
|
+
export const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap =
|
|
3680
|
+
/*@__PURE__*/ S.Record(
|
|
3681
|
+
S.String,
|
|
3682
|
+
UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue,
|
|
3683
|
+
) as any as S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap>;
|
|
3684
|
+
|
|
3685
|
+
export interface UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest {
|
|
3686
|
+
useful_links: UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap;
|
|
3687
|
+
}
|
|
3688
|
+
export const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest =
|
|
3689
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3690
|
+
S.Struct({
|
|
3691
|
+
useful_links:
|
|
3692
|
+
UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap,
|
|
3693
|
+
}).pipe(
|
|
3694
|
+
T.Http({
|
|
3695
|
+
method: "POST",
|
|
3696
|
+
uri: "/model_hub/update_useful_links",
|
|
3697
|
+
code: 200,
|
|
3698
|
+
}),
|
|
3699
|
+
),
|
|
3700
|
+
).annotate({
|
|
3701
|
+
identifier: "UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest",
|
|
3702
|
+
}) as any as S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest>;
|
|
3703
|
+
|
|
3704
|
+
export interface UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse {
|
|
3705
|
+
body: unknown;
|
|
3706
|
+
}
|
|
3707
|
+
export const UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse =
|
|
3708
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3709
|
+
S.Struct({
|
|
3710
|
+
body: S.Unknown,
|
|
3711
|
+
}),
|
|
3712
|
+
).annotate({
|
|
3713
|
+
identifier: "UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse",
|
|
3714
|
+
}) as any as S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse>;
|
|
3715
|
+
|
|
3716
|
+
export type ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap =
|
|
3717
|
+
{ [key: string]: unknown | undefined };
|
|
3718
|
+
export const ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap =
|
|
3719
|
+
/*@__PURE__*/ S.Record(
|
|
3720
|
+
S.String,
|
|
3721
|
+
S.Unknown,
|
|
3722
|
+
) as any as S.Schema<ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap>;
|
|
3723
|
+
|
|
3724
|
+
export interface ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest {
|
|
3725
|
+
complexity_router_config: ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap;
|
|
3726
|
+
/** Team the router is being created for. Required for a team admin, who may only validate their own team's routers */
|
|
3727
|
+
team_id?: string | null;
|
|
3728
|
+
}
|
|
3729
|
+
export const ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest =
|
|
3730
|
+
/*@__PURE__*/ S.suspend(() =>
|
|
3731
|
+
S.Struct({
|
|
3732
|
+
complexity_router_config:
|
|
3733
|
+
ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap,
|
|
3734
|
+
team_id: S.optional(S.NullOr(S.String)),
|
|
3735
|
+
}).pipe(
|
|
3736
|
+
T.Http({
|
|
3737
|
+
method: "POST",
|
|
3738
|
+
uri: "/auto_router/validate_complexity_router_config",
|
|
3739
|
+
code: 200,
|
|
3740
|
+
}),
|
|
3741
|
+
),
|
|
3742
|
+
).annotate({
|
|
3743
|
+
identifier:
|
|
3744
|
+
"ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest",
|
|
3745
|
+
}) as any as S.Schema<ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest>;
|
|
3746
|
+
|
|
3747
|
+
export interface ComplexityRouterConfigValidationResponse {
|
|
3748
|
+
error?: string | null;
|
|
3749
|
+
valid: boolean;
|
|
3750
|
+
}
|
|
3751
|
+
export const ComplexityRouterConfigValidationResponse = /*@__PURE__*/ S.suspend(
|
|
3752
|
+
() =>
|
|
3753
|
+
S.Struct({
|
|
3754
|
+
error: S.optional(S.NullOr(S.String)),
|
|
3755
|
+
valid: S.Boolean,
|
|
3756
|
+
}),
|
|
3757
|
+
).annotate({
|
|
3758
|
+
identifier: "ComplexityRouterConfigValidationResponse",
|
|
3759
|
+
}) as any as S.Schema<ComplexityRouterConfigValidationResponse>;
|
|
3760
|
+
|
|
3761
|
+
export type AddNewModelModelNewPostError = UnprocessableEntity | LitellmOpError;
|
|
3762
|
+
/** Add New Model Allows adding new models to the model list in the config.yaml */
|
|
3763
|
+
export const addNewModelModelNewPost: API.OperationMethod<
|
|
3764
|
+
AddNewModelModelNewPostRequest,
|
|
3765
|
+
AddNewModelModelNewPostResponse,
|
|
3766
|
+
AddNewModelModelNewPostError,
|
|
3767
|
+
LitellmOpContext
|
|
3768
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3769
|
+
input: AddNewModelModelNewPostRequest,
|
|
3770
|
+
output: AddNewModelModelNewPostResponse,
|
|
3771
|
+
errors: [UnprocessableEntity],
|
|
3772
|
+
protocol: LitellmProtocol,
|
|
3773
|
+
retry: Retry.Retry,
|
|
3774
|
+
}));
|
|
3775
|
+
|
|
3776
|
+
export type CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteError =
|
|
3777
|
+
LitellmOpError;
|
|
3778
|
+
/** Cancel Anthropic Beta Headers Reload ADMIN ONLY / MASTER KEY Only Endpoint Cancel the scheduled periodic reload of the Anthropic beta headers configuration. */
|
|
3779
|
+
export const cancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDelete: API.OperationMethod<
|
|
3780
|
+
CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest,
|
|
3781
|
+
CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse,
|
|
3782
|
+
CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteError,
|
|
3783
|
+
LitellmOpContext
|
|
3784
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3785
|
+
input:
|
|
3786
|
+
CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest,
|
|
3787
|
+
output:
|
|
3788
|
+
CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse,
|
|
3789
|
+
errors: [],
|
|
3790
|
+
protocol: LitellmProtocol,
|
|
3791
|
+
retry: Retry.Retry,
|
|
3792
|
+
}));
|
|
3793
|
+
|
|
3794
|
+
export type CancelModelCostMapReloadScheduleModelCostMapReloadDeleteError =
|
|
3795
|
+
LitellmOpError;
|
|
3796
|
+
/** Cancel Model Cost Map Reload ADMIN ONLY / MASTER KEY Only Endpoint Cancel the scheduled periodic reload of the model cost map. */
|
|
3797
|
+
export const cancelModelCostMapReloadScheduleModelCostMapReloadDelete: API.OperationMethod<
|
|
3798
|
+
CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest,
|
|
3799
|
+
CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse,
|
|
3800
|
+
CancelModelCostMapReloadScheduleModelCostMapReloadDeleteError,
|
|
3801
|
+
LitellmOpContext
|
|
3802
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3803
|
+
input: CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest,
|
|
3804
|
+
output: CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse,
|
|
3805
|
+
errors: [],
|
|
3806
|
+
protocol: LitellmProtocol,
|
|
3807
|
+
retry: Retry.Retry,
|
|
3808
|
+
}));
|
|
3809
|
+
|
|
3810
|
+
export type CreateModelGroupAccessGroupNewPostError =
|
|
3811
|
+
| UnprocessableEntity
|
|
3812
|
+
| LitellmOpError;
|
|
3813
|
+
/** Create Model Group Create a new access group containing multiple model names. An access group is a named collection of model groups that can be referenced by teams/keys for simplified access control. Example: ```bash curl -X POST 'http://localhost:4000/access_group/new' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "access_group": "production-models", "model_names": ["gpt-4", "claude-3-opus", "gemini-pro"] }' ``` Parameters: - access_group: str - The access group name (e.g., "production-models") - model_names: List[str] - List of existing model groups to include Returns: - NewModelGroupResponse with the created access group details Raises: - HTTPException 400: If any model names don't exist - HTTPException 500: If database operations fail */
|
|
3814
|
+
export const createModelGroupAccessGroupNewPost: API.OperationMethod<
|
|
3815
|
+
CreateModelGroupAccessGroupNewPostRequest,
|
|
3816
|
+
NewModelGroupResponse,
|
|
3817
|
+
CreateModelGroupAccessGroupNewPostError,
|
|
3818
|
+
LitellmOpContext
|
|
3819
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3820
|
+
input: CreateModelGroupAccessGroupNewPostRequest,
|
|
3821
|
+
output: NewModelGroupResponse,
|
|
3822
|
+
errors: [UnprocessableEntity],
|
|
3823
|
+
protocol: LitellmProtocol,
|
|
3824
|
+
retry: Retry.Retry,
|
|
3825
|
+
}));
|
|
3826
|
+
|
|
3827
|
+
export type DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteError =
|
|
3828
|
+
| UnprocessableEntity
|
|
3829
|
+
| LitellmOpError;
|
|
3830
|
+
/** Delete Access Group Delete an access group. Removes the access group from all deployments that have it. Example: ```bash curl -X DELETE 'http://localhost:4000/access_group/production-models/delete' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - DeleteModelGroupResponse with deletion details Raises: - HTTPException 404: If access group not found */
|
|
3831
|
+
export const deleteAccessGroupAccessGroupAccessGroupDeleteDelete: API.OperationMethod<
|
|
3832
|
+
DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest,
|
|
3833
|
+
DeleteModelGroupResponse,
|
|
3834
|
+
DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteError,
|
|
3835
|
+
LitellmOpContext
|
|
3836
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3837
|
+
input: DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest,
|
|
3838
|
+
output: DeleteModelGroupResponse,
|
|
3839
|
+
errors: [UnprocessableEntity],
|
|
3840
|
+
protocol: LitellmProtocol,
|
|
3841
|
+
retry: Retry.Retry,
|
|
3842
|
+
}));
|
|
3843
|
+
|
|
3844
|
+
export type DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteError =
|
|
3845
|
+
| UnprocessableEntity
|
|
3846
|
+
| LitellmOpError;
|
|
3847
|
+
/** Delete Access Group Budget Clear the shared budget of an access group, leaving the group itself in place. Example: ```bash curl -X DELETE 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - DeleteAccessGroupBudgetResponse; budget_deleted is false when there was nothing to clear Raises: - HTTPException 404: If access group not found */
|
|
3848
|
+
export const deleteAccessGroupBudgetAccessGroupAccessGroupBudgetDelete: API.OperationMethod<
|
|
3849
|
+
DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest,
|
|
3850
|
+
DeleteAccessGroupBudgetResponse,
|
|
3851
|
+
DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteError,
|
|
3852
|
+
LitellmOpContext
|
|
3853
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3854
|
+
input: DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest,
|
|
3855
|
+
output: DeleteAccessGroupBudgetResponse,
|
|
3856
|
+
errors: [UnprocessableEntity],
|
|
3857
|
+
protocol: LitellmProtocol,
|
|
3858
|
+
retry: Retry.Retry,
|
|
3859
|
+
}));
|
|
3860
|
+
|
|
3861
|
+
export type DeleteModelModelDeletePostError =
|
|
3862
|
+
| BadRequest
|
|
3863
|
+
| UnprocessableEntity
|
|
3864
|
+
| LitellmOpError;
|
|
3865
|
+
/** Delete Model Allows deleting models in the model list in the config.yaml */
|
|
3866
|
+
export const deleteModelModelDeletePost: API.OperationMethod<
|
|
3867
|
+
DeleteModelModelDeletePostRequest,
|
|
3868
|
+
DeleteModelModelDeletePostResponse,
|
|
3869
|
+
DeleteModelModelDeletePostError,
|
|
3870
|
+
LitellmOpContext
|
|
3871
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3872
|
+
input: DeleteModelModelDeletePostRequest,
|
|
3873
|
+
output: DeleteModelModelDeletePostResponse,
|
|
3874
|
+
errors: [BadRequest, UnprocessableEntity],
|
|
3875
|
+
protocol: LitellmProtocol,
|
|
3876
|
+
retry: Retry.Retry,
|
|
3877
|
+
}));
|
|
3878
|
+
|
|
3879
|
+
export type GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetError =
|
|
3880
|
+
| UnprocessableEntity
|
|
3881
|
+
| LitellmOpError;
|
|
3882
|
+
/** Get Access Group Budget Get the shared budget of an access group, and the spend drawn against it. Example: ```bash curl -X GET 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - AccessGroupBudgetResponse; budget is null when the group has no budget set Raises: - HTTPException 404: If access group not found */
|
|
3883
|
+
export const getAccessGroupBudgetAccessGroupAccessGroupBudgetGet: API.OperationMethod<
|
|
3884
|
+
GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest,
|
|
3885
|
+
AccessGroupBudgetResponse,
|
|
3886
|
+
GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetError,
|
|
3887
|
+
LitellmOpContext
|
|
3888
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3889
|
+
input: GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest,
|
|
3890
|
+
output: AccessGroupBudgetResponse,
|
|
3891
|
+
errors: [UnprocessableEntity],
|
|
3892
|
+
protocol: LitellmProtocol,
|
|
3893
|
+
retry: Retry.Retry,
|
|
3894
|
+
}));
|
|
3895
|
+
|
|
3896
|
+
export type GetAccessGroupInfoAccessGroupAccessGroupInfoGetError =
|
|
3897
|
+
| UnprocessableEntity
|
|
3898
|
+
| LitellmOpError;
|
|
3899
|
+
/** Get Access Group Info Get information about a specific access group. Example: ```bash curl -X GET 'http://localhost:4000/access_group/production-models/info' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - AccessGroupInfo with the access group details, its shared budget and its spend Raises: - HTTPException 404: If access group not found */
|
|
3900
|
+
export const getAccessGroupInfoAccessGroupAccessGroupInfoGet: API.OperationMethod<
|
|
3901
|
+
GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest,
|
|
3902
|
+
AccessGroupInfo,
|
|
3903
|
+
GetAccessGroupInfoAccessGroupAccessGroupInfoGetError,
|
|
3904
|
+
LitellmOpContext
|
|
3905
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3906
|
+
input: GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest,
|
|
3907
|
+
output: AccessGroupInfo,
|
|
3908
|
+
errors: [UnprocessableEntity],
|
|
3909
|
+
protocol: LitellmProtocol,
|
|
3910
|
+
retry: Retry.Retry,
|
|
3911
|
+
}));
|
|
3912
|
+
|
|
3913
|
+
export type GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetError =
|
|
3914
|
+
LitellmOpError;
|
|
3915
|
+
/** Get Anthropic Beta Headers Reload Status ADMIN ONLY / MASTER KEY Only Endpoint Get the status of the scheduled Anthropic beta headers reload job. */
|
|
3916
|
+
export const getAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGet: API.OperationMethod<
|
|
3917
|
+
GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest,
|
|
3918
|
+
GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse,
|
|
3919
|
+
GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetError,
|
|
3920
|
+
LitellmOpContext
|
|
3921
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3922
|
+
input:
|
|
3923
|
+
GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest,
|
|
3924
|
+
output:
|
|
3925
|
+
GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse,
|
|
3926
|
+
errors: [],
|
|
3927
|
+
protocol: LitellmProtocol,
|
|
3928
|
+
retry: Retry.Retry,
|
|
3929
|
+
}));
|
|
3930
|
+
|
|
3931
|
+
export type GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetError =
|
|
3932
|
+
| UnprocessableEntity
|
|
3933
|
+
| LitellmOpError;
|
|
3934
|
+
/** Get Auto Router Classifier Default Prompt Get the built-in system prompt used by an auto-router's LLM classifier */
|
|
3935
|
+
export const getAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGet: API.OperationMethod<
|
|
3936
|
+
GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest,
|
|
3937
|
+
AutoRouterClassifierDefaultPromptResponse,
|
|
3938
|
+
GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetError,
|
|
3939
|
+
LitellmOpContext
|
|
3940
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3941
|
+
input:
|
|
3942
|
+
GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest,
|
|
3943
|
+
output: AutoRouterClassifierDefaultPromptResponse,
|
|
3944
|
+
errors: [UnprocessableEntity],
|
|
3945
|
+
protocol: LitellmProtocol,
|
|
3946
|
+
retry: Retry.Retry,
|
|
3947
|
+
}));
|
|
3948
|
+
|
|
3949
|
+
export type GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetError =
|
|
3950
|
+
LitellmOpError;
|
|
3951
|
+
/** Get Model Cost Map Reload Status ADMIN ONLY / MASTER KEY Only Endpoint Get the status of the scheduled model cost map reload job. */
|
|
3952
|
+
export const getModelCostMapReloadStatusScheduleModelCostMapReloadStatusGet: API.OperationMethod<
|
|
3953
|
+
GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest,
|
|
3954
|
+
GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse,
|
|
3955
|
+
GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetError,
|
|
3956
|
+
LitellmOpContext
|
|
3957
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3958
|
+
input: GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest,
|
|
3959
|
+
output:
|
|
3960
|
+
GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse,
|
|
3961
|
+
errors: [],
|
|
3962
|
+
protocol: LitellmProtocol,
|
|
3963
|
+
retry: Retry.Retry,
|
|
3964
|
+
}));
|
|
3965
|
+
|
|
3966
|
+
export type GetModelCostMapSourceModelCostMapSourceGetError = LitellmOpError;
|
|
3967
|
+
/** Get Model Cost Map Source ADMIN ONLY / MASTER KEY Only Endpoint Returns information about where the current model cost/pricing data was loaded from. Response fields: - source: "local" (bundled backup) or "remote" (fetched from URL) - url: the remote URL that was attempted (null when env-forced local) - is_env_forced: true if LITELLM_LOCAL_MODEL_COST_MAP=True forced local usage - fallback_reason: human-readable reason why remote failed (null on success) - model_count: number of models in the currently loaded cost map */
|
|
3968
|
+
export const getModelCostMapSourceModelCostMapSourceGet: API.OperationMethod<
|
|
3969
|
+
GetModelCostMapSourceModelCostMapSourceGetRequest,
|
|
3970
|
+
GetModelCostMapSourceModelCostMapSourceGetResponse,
|
|
3971
|
+
GetModelCostMapSourceModelCostMapSourceGetError,
|
|
3972
|
+
LitellmOpContext
|
|
3973
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3974
|
+
input: GetModelCostMapSourceModelCostMapSourceGetRequest,
|
|
3975
|
+
output: GetModelCostMapSourceModelCostMapSourceGetResponse,
|
|
3976
|
+
errors: [],
|
|
3977
|
+
protocol: LitellmProtocol,
|
|
3978
|
+
retry: Retry.Retry,
|
|
3979
|
+
}));
|
|
3980
|
+
|
|
3981
|
+
export type GetModelDeprecationsModelDeprecationError =
|
|
3982
|
+
| UnprocessableEntity
|
|
3983
|
+
| LitellmOpError;
|
|
3984
|
+
/** Model Deprecations List models with known deprecation/sunset dates, bucketed by urgency. Reads `deprecation_date` metadata from `model_prices_and_context_window.json` (and any per-deployment `model_info.deprecation_date` overrides) for the models configured on this proxy. Parameters: warn_within_days: Window (in days) used to bucket "imminent" models, 30 by default. Returns: A payload with three lists of `ModelDeprecationInfo` entries: - `deprecated`: deprecation date is in the past, so these requests may fail at any time. - `imminent`: deprecation date is within `warn_within_days` from today. - `upcoming`: deprecation date is further out. Example: ```shell curl -X GET 'http://localhost:4000/model/deprecations' \ -H 'Authorization: Bearer sk-1234' ``` */
|
|
3985
|
+
export const getModelDeprecationsModelDeprecation: API.OperationMethod<
|
|
3986
|
+
GetModelDeprecationsModelDeprecationRequest,
|
|
3987
|
+
ModelDeprecationResponse,
|
|
3988
|
+
GetModelDeprecationsModelDeprecationError,
|
|
3989
|
+
LitellmOpContext
|
|
3990
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
3991
|
+
input: GetModelDeprecationsModelDeprecationRequest,
|
|
3992
|
+
output: ModelDeprecationResponse,
|
|
3993
|
+
errors: [UnprocessableEntity],
|
|
3994
|
+
protocol: LitellmProtocol,
|
|
3995
|
+
retry: Retry.Retry,
|
|
3996
|
+
}));
|
|
3997
|
+
|
|
3998
|
+
export type GetModelDeprecationsV1ModelDeprecationError =
|
|
3999
|
+
| UnprocessableEntity
|
|
4000
|
+
| LitellmOpError;
|
|
4001
|
+
/** Model Deprecations List models with known deprecation/sunset dates, bucketed by urgency. Reads `deprecation_date` metadata from `model_prices_and_context_window.json` (and any per-deployment `model_info.deprecation_date` overrides) for the models configured on this proxy. Parameters: warn_within_days: Window (in days) used to bucket "imminent" models, 30 by default. Returns: A payload with three lists of `ModelDeprecationInfo` entries: - `deprecated`: deprecation date is in the past, so these requests may fail at any time. - `imminent`: deprecation date is within `warn_within_days` from today. - `upcoming`: deprecation date is further out. Example: ```shell curl -X GET 'http://localhost:4000/model/deprecations' \ -H 'Authorization: Bearer sk-1234' ``` */
|
|
4002
|
+
export const getModelDeprecationsV1ModelDeprecation: API.OperationMethod<
|
|
4003
|
+
GetModelDeprecationsV1ModelDeprecationRequest,
|
|
4004
|
+
ModelDeprecationResponse,
|
|
4005
|
+
GetModelDeprecationsV1ModelDeprecationError,
|
|
4006
|
+
LitellmOpContext
|
|
4007
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4008
|
+
input: GetModelDeprecationsV1ModelDeprecationRequest,
|
|
4009
|
+
output: ModelDeprecationResponse,
|
|
4010
|
+
errors: [UnprocessableEntity],
|
|
4011
|
+
protocol: LitellmProtocol,
|
|
4012
|
+
retry: Retry.Retry,
|
|
4013
|
+
}));
|
|
4014
|
+
|
|
4015
|
+
export type GetModelGroupInfoModelGroupInfoError =
|
|
4016
|
+
| UnprocessableEntity
|
|
4017
|
+
| LitellmOpError;
|
|
4018
|
+
/** Model Group Info Get information about all the deployments on litellm proxy, including config.yaml descriptions (except api key and api base) - /model_group/info returns all model groups. End users of proxy should use /model_group/info since those models will be used for /chat/completions, /embeddings, etc. - /model_group/info?model_group=rerank-english-v3.0 returns all model groups for a specific model group (`model_name` in config.yaml) Example Request (All Models): ```shell curl -X 'GET' 'http://localhost:4000/model_group/info' -H 'accept: application/json' -H 'x-api-key: sk-1234' ``` Example Request (Specific Model Group): ```shell curl -X 'GET' 'http://localhost:4000/model_group/info?model_group=rerank-english-v3.0' -H 'accept: application/json' -H 'Authorization: Bearer sk-1234' ``` Example Request (Specific Wildcard Model Group): (e.g. `model_name: openai/*` on config.yaml) ```shell curl -X 'GET' 'http://localhost:4000/model_group/info?model_group=openai/tts-1' -H 'accept: application/json' -H 'Authorization: Bearersk-1234' ``` Learn how to use and set wildcard models [here](https://docs.litellm.ai/docs/wildcard_routing) Example Response: ```json { "data": [ { "model_group": "rerank-english-v3.0", "providers": [ "cohere" ], "max_input_tokens": null, "max_output_tokens": null, "input_cost_per_token": 0.0, "output_cost_per_token": 0.0, "mode": null, "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": false, "supports_function_calling": false, "supported_openai_params": [ "stream", "temperature", "max_tokens", "logit_bias", "top_p", "frequency_penalty", "presence_penalty", "stop", "n", "extra_headers" ] }, { "model_group": "gpt-3.5-turbo", "providers": [ "openai" ], "max_input_tokens": 16385.0, "max_output_tokens": 4096.0, "input_cost_per_token": 1.5e-06, "output_cost_per_token": 2e-06, "mode": "chat", "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": false, "supports_function_calling": true, "supported_openai_params": [ "frequency_penalty", "logit_bias", "logprobs", "top_logprobs", "max_tokens", "max_completion_tokens", "n", "presence_penalty", "seed", "stop", "stream", "stream_options", "temperature", "top_p", "tools", "tool_choice", "function_call", "functions", "max_retries", "extra_headers", "parallel_tool_calls", "response_format" ] }, { "model_group": "llava-hf", "providers": [ "openai" ], "max_input_tokens": null, "max_output_tokens": null, "input_cost_per_token": 0.0, "output_cost_per_token": 0.0, "mode": null, "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": true, "supports_function_calling": false, "supported_openai_params": [ "frequency_penalty", "logit_bias", "logprobs", "top_logprobs", "max_tokens", "max_completion_tokens", "n", "presence_penalty", "seed", "stop", "stream", "stream_options", "temperature", "top_p", "tools", "tool_choice", "function_call", "functions", "max_retries", "extra_headers", "parallel_tool_calls", "response_format" ] } ] } ``` */
|
|
4019
|
+
export const getModelGroupInfoModelGroupInfo: API.OperationMethod<
|
|
4020
|
+
GetModelGroupInfoModelGroupInfoRequest,
|
|
4021
|
+
GetModelGroupInfoModelGroupInfoResponse,
|
|
4022
|
+
GetModelGroupInfoModelGroupInfoError,
|
|
4023
|
+
LitellmOpContext
|
|
4024
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4025
|
+
input: GetModelGroupInfoModelGroupInfoRequest,
|
|
4026
|
+
output: GetModelGroupInfoModelGroupInfoResponse,
|
|
4027
|
+
errors: [UnprocessableEntity],
|
|
4028
|
+
protocol: LitellmProtocol,
|
|
4029
|
+
retry: Retry.Retry,
|
|
4030
|
+
}));
|
|
4031
|
+
|
|
4032
|
+
export type GetModelInfoModelsModelIdError =
|
|
4033
|
+
| UnprocessableEntity
|
|
4034
|
+
| LitellmOpError;
|
|
4035
|
+
/** Model Info Retrieve information about a specific model accessible to your API key. Returns model details only if the model is available to your API key/team. Returns 404 if the model doesn't exist or is not accessible. Follows OpenAI API specification for individual model retrieval. https://platform.openai.com/docs/api-reference/models/retrieve Query parameters mirror `/v1/models` so the same caller context (team scoping, health filtering, paused deployments) drives both endpoints; the listing's public id must resolve to the same internal deployment here. */
|
|
4036
|
+
export const getModelInfoModelsModelId: API.OperationMethod<
|
|
4037
|
+
GetModelInfoModelsModelIdRequest,
|
|
4038
|
+
GetModelInfoModelsModelIdResponse,
|
|
4039
|
+
GetModelInfoModelsModelIdError,
|
|
4040
|
+
LitellmOpContext
|
|
4041
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4042
|
+
input: GetModelInfoModelsModelIdRequest,
|
|
4043
|
+
output: GetModelInfoModelsModelIdResponse,
|
|
4044
|
+
errors: [UnprocessableEntity],
|
|
4045
|
+
protocol: LitellmProtocol,
|
|
4046
|
+
retry: Retry.Retry,
|
|
4047
|
+
}));
|
|
4048
|
+
|
|
4049
|
+
export type GetModelInfoV1ModelInfoError = UnprocessableEntity | LitellmOpError;
|
|
4050
|
+
/** Model Info V1 Provides more info about each model in /models, including config.yaml descriptions (except api key and api base) Parameters: litellm_model_id: Optional[str] = None (this is the value of `x-litellm-model-id` returned in response headers) - When litellm_model_id is passed, it will return the info for that specific model - When litellm_model_id is not passed, it will return the info for all models - include_team_models: When true, filter to deployments the caller can use (same as /v2/model/info). - teamId: Filter to models accessible by the given team. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks, matching `/v1/models?healthy_only=true`. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true`, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Ignored when `litellm_model_id` is passed, since that is a direct lookup of one deployment rather than a listing. Hiding is presentation-only: a hidden model can still be called directly. Each model in the list response includes `model_info.access_via_team_ids` and `model_info.direct_access` when the proxy database is connected. Returns: Returns a dictionary containing information about each model. Example Response: ```json { "data": [ { "model_name": "fake-openai-endpoint", "litellm_params": { "api_base": "https://exampleopenaiendpoint-production.up.railway.app/", "model": "openai/fake" }, "model_info": { "id": "112f74fab24a7a5245d2ced3536dd8f5f9192c57ee6e332af0f0512e08bed5af", "db_model": false } } ] } ``` */
|
|
4051
|
+
export const getModelInfoV1ModelInfo: API.OperationMethod<
|
|
4052
|
+
GetModelInfoV1ModelInfoRequest,
|
|
4053
|
+
GetModelInfoV1ModelInfoResponse,
|
|
4054
|
+
GetModelInfoV1ModelInfoError,
|
|
4055
|
+
LitellmOpContext
|
|
4056
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4057
|
+
input: GetModelInfoV1ModelInfoRequest,
|
|
4058
|
+
output: GetModelInfoV1ModelInfoResponse,
|
|
4059
|
+
errors: [UnprocessableEntity],
|
|
4060
|
+
protocol: LitellmProtocol,
|
|
4061
|
+
retry: Retry.Retry,
|
|
4062
|
+
}));
|
|
4063
|
+
|
|
4064
|
+
export type GetModelInfoV1ModelsModelIdError =
|
|
4065
|
+
| UnprocessableEntity
|
|
4066
|
+
| LitellmOpError;
|
|
4067
|
+
/** Model Info Retrieve information about a specific model accessible to your API key. Returns model details only if the model is available to your API key/team. Returns 404 if the model doesn't exist or is not accessible. Follows OpenAI API specification for individual model retrieval. https://platform.openai.com/docs/api-reference/models/retrieve Query parameters mirror `/v1/models` so the same caller context (team scoping, health filtering, paused deployments) drives both endpoints; the listing's public id must resolve to the same internal deployment here. */
|
|
4068
|
+
export const getModelInfoV1ModelsModelId: API.OperationMethod<
|
|
4069
|
+
GetModelInfoV1ModelsModelIdRequest,
|
|
4070
|
+
GetModelInfoV1ModelsModelIdResponse,
|
|
4071
|
+
GetModelInfoV1ModelsModelIdError,
|
|
4072
|
+
LitellmOpContext
|
|
4073
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4074
|
+
input: GetModelInfoV1ModelsModelIdRequest,
|
|
4075
|
+
output: GetModelInfoV1ModelsModelIdResponse,
|
|
4076
|
+
errors: [UnprocessableEntity],
|
|
4077
|
+
protocol: LitellmProtocol,
|
|
4078
|
+
retry: Retry.Retry,
|
|
4079
|
+
}));
|
|
4080
|
+
|
|
4081
|
+
export type GetModelInfoV1V1ModelInfoError =
|
|
4082
|
+
| UnprocessableEntity
|
|
4083
|
+
| LitellmOpError;
|
|
4084
|
+
/** Model Info V1 Provides more info about each model in /models, including config.yaml descriptions (except api key and api base) Parameters: litellm_model_id: Optional[str] = None (this is the value of `x-litellm-model-id` returned in response headers) - When litellm_model_id is passed, it will return the info for that specific model - When litellm_model_id is not passed, it will return the info for all models - include_team_models: When true, filter to deployments the caller can use (same as /v2/model/info). - teamId: Filter to models accessible by the given team. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks, matching `/v1/models?healthy_only=true`. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true`, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Ignored when `litellm_model_id` is passed, since that is a direct lookup of one deployment rather than a listing. Hiding is presentation-only: a hidden model can still be called directly. Each model in the list response includes `model_info.access_via_team_ids` and `model_info.direct_access` when the proxy database is connected. Returns: Returns a dictionary containing information about each model. Example Response: ```json { "data": [ { "model_name": "fake-openai-endpoint", "litellm_params": { "api_base": "https://exampleopenaiendpoint-production.up.railway.app/", "model": "openai/fake" }, "model_info": { "id": "112f74fab24a7a5245d2ced3536dd8f5f9192c57ee6e332af0f0512e08bed5af", "db_model": false } } ] } ``` */
|
|
4085
|
+
export const getModelInfoV1V1ModelInfo: API.OperationMethod<
|
|
4086
|
+
GetModelInfoV1V1ModelInfoRequest,
|
|
4087
|
+
GetModelInfoV1V1ModelInfoResponse,
|
|
4088
|
+
GetModelInfoV1V1ModelInfoError,
|
|
4089
|
+
LitellmOpContext
|
|
4090
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4091
|
+
input: GetModelInfoV1V1ModelInfoRequest,
|
|
4092
|
+
output: GetModelInfoV1V1ModelInfoResponse,
|
|
4093
|
+
errors: [UnprocessableEntity],
|
|
4094
|
+
protocol: LitellmProtocol,
|
|
4095
|
+
retry: Retry.Retry,
|
|
4096
|
+
}));
|
|
4097
|
+
|
|
4098
|
+
export type GetModelInfoV2V2ModelInfoError =
|
|
4099
|
+
| UnprocessableEntity
|
|
4100
|
+
| LitellmOpError;
|
|
4101
|
+
/** Model Info V2 Paginated model metadata for proxy deployments (pricing, provider, team access). Returns configured router deployments with enriched `model_info` (costs, provider, context window, etc.). Sensitive fields such as API keys and api_base are omitted. Query parameters: model: Filter to a single public `model_name`. user_models_only: When true, only return models created by the calling user. include_team_models: When true, populate `access_via_team_ids` and `direct_access` on each model and filter to deployments the caller can use. page / size: Pagination controls (defaults: page=1, size=50). search: Case-insensitive partial match on model name or team public name. modelId: Return a single deployment by LiteLLM model id. teamId: Filter to models with direct access or team membership for this team id. sortBy / sortOrder: Sort by model_name, created_at, updated_at, costs, or status. Example request: ``` curl -X GET 'http://localhost:4000/v2/model/info?include_team_models=true&page=1&size=50' \ --header 'Authorization: Bearer sk-1234' ``` Example response: ```json { "data": [ { "model_name": "gpt-4", "litellm_params": {"model": "openai/gpt-4.1"}, "model_info": { "id": "abc123", "litellm_provider": "openai", "access_via_team_ids": ["team-1"], "direct_access": true } } ], "total_count": 1, "current_page": 1, "total_pages": 1, "size": 50 } ``` */
|
|
4102
|
+
export const getModelInfoV2V2ModelInfo: API.OperationMethod<
|
|
4103
|
+
GetModelInfoV2V2ModelInfoRequest,
|
|
4104
|
+
GetModelInfoV2V2ModelInfoResponse,
|
|
4105
|
+
GetModelInfoV2V2ModelInfoError,
|
|
4106
|
+
LitellmOpContext
|
|
4107
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4108
|
+
input: GetModelInfoV2V2ModelInfoRequest,
|
|
4109
|
+
output: GetModelInfoV2V2ModelInfoResponse,
|
|
4110
|
+
errors: [UnprocessableEntity],
|
|
4111
|
+
protocol: LitellmProtocol,
|
|
4112
|
+
retry: Retry.Retry,
|
|
4113
|
+
}));
|
|
4114
|
+
|
|
4115
|
+
export type GetModelMetricsExceptionsModelMetricsExceptionError =
|
|
4116
|
+
| UnprocessableEntity
|
|
4117
|
+
| LitellmOpError;
|
|
4118
|
+
/** Model Metrics Exceptions View number of failed requests per model on config.yaml */
|
|
4119
|
+
export const getModelMetricsExceptionsModelMetricsException: API.OperationMethod<
|
|
4120
|
+
GetModelMetricsExceptionsModelMetricsExceptionRequest,
|
|
4121
|
+
GetModelMetricsExceptionsModelMetricsExceptionResponse,
|
|
4122
|
+
GetModelMetricsExceptionsModelMetricsExceptionError,
|
|
4123
|
+
LitellmOpContext
|
|
4124
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4125
|
+
input: GetModelMetricsExceptionsModelMetricsExceptionRequest,
|
|
4126
|
+
output: GetModelMetricsExceptionsModelMetricsExceptionResponse,
|
|
4127
|
+
errors: [UnprocessableEntity],
|
|
4128
|
+
protocol: LitellmProtocol,
|
|
4129
|
+
retry: Retry.Retry,
|
|
4130
|
+
}));
|
|
4131
|
+
|
|
4132
|
+
export type GetModelMetricsModelMetricsError =
|
|
4133
|
+
| UnprocessableEntity
|
|
4134
|
+
| LitellmOpError;
|
|
4135
|
+
/** Model Metrics View number of requests & avg latency per model on config.yaml */
|
|
4136
|
+
export const getModelMetricsModelMetrics: API.OperationMethod<
|
|
4137
|
+
GetModelMetricsModelMetricsRequest,
|
|
4138
|
+
GetModelMetricsModelMetricsResponse,
|
|
4139
|
+
GetModelMetricsModelMetricsError,
|
|
4140
|
+
LitellmOpContext
|
|
4141
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4142
|
+
input: GetModelMetricsModelMetricsRequest,
|
|
4143
|
+
output: GetModelMetricsModelMetricsResponse,
|
|
4144
|
+
errors: [UnprocessableEntity],
|
|
4145
|
+
protocol: LitellmProtocol,
|
|
4146
|
+
retry: Retry.Retry,
|
|
4147
|
+
}));
|
|
4148
|
+
|
|
4149
|
+
export type GetModelMetricsSlowResponsesModelMetricsSlowResponseError =
|
|
4150
|
+
| UnprocessableEntity
|
|
4151
|
+
| LitellmOpError;
|
|
4152
|
+
/** Model Metrics Slow Responses View number of hanging requests per model_group */
|
|
4153
|
+
export const getModelMetricsSlowResponsesModelMetricsSlowResponse: API.OperationMethod<
|
|
4154
|
+
GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest,
|
|
4155
|
+
GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse,
|
|
4156
|
+
GetModelMetricsSlowResponsesModelMetricsSlowResponseError,
|
|
4157
|
+
LitellmOpContext
|
|
4158
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4159
|
+
input: GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest,
|
|
4160
|
+
output: GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse,
|
|
4161
|
+
errors: [UnprocessableEntity],
|
|
4162
|
+
protocol: LitellmProtocol,
|
|
4163
|
+
retry: Retry.Retry,
|
|
4164
|
+
}));
|
|
4165
|
+
|
|
4166
|
+
export type GetModelSettingsModelSettingsError = LitellmOpError;
|
|
4167
|
+
/** Model Settings Returns provider name, description, and required parameters for each provider */
|
|
4168
|
+
export const getModelSettingsModelSettings: API.OperationMethod<
|
|
4169
|
+
GetModelSettingsModelSettingsRequest,
|
|
4170
|
+
GetModelSettingsModelSettingsResponse,
|
|
4171
|
+
GetModelSettingsModelSettingsError,
|
|
4172
|
+
LitellmOpContext
|
|
4173
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4174
|
+
input: GetModelSettingsModelSettingsRequest,
|
|
4175
|
+
output: GetModelSettingsModelSettingsResponse,
|
|
4176
|
+
errors: [],
|
|
4177
|
+
protocol: LitellmProtocol,
|
|
4178
|
+
retry: Retry.Retry,
|
|
4179
|
+
}));
|
|
4180
|
+
|
|
4181
|
+
export type GetModelStreamingMetricsModelStreamingMetricsError =
|
|
4182
|
+
| UnprocessableEntity
|
|
4183
|
+
| LitellmOpError;
|
|
4184
|
+
/** Model Streaming Metrics View time to first token for models in spend logs */
|
|
4185
|
+
export const getModelStreamingMetricsModelStreamingMetrics: API.OperationMethod<
|
|
4186
|
+
GetModelStreamingMetricsModelStreamingMetricsRequest,
|
|
4187
|
+
GetModelStreamingMetricsModelStreamingMetricsResponse,
|
|
4188
|
+
GetModelStreamingMetricsModelStreamingMetricsError,
|
|
4189
|
+
LitellmOpContext
|
|
4190
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4191
|
+
input: GetModelStreamingMetricsModelStreamingMetricsRequest,
|
|
4192
|
+
output: GetModelStreamingMetricsModelStreamingMetricsResponse,
|
|
4193
|
+
errors: [UnprocessableEntity],
|
|
4194
|
+
protocol: LitellmProtocol,
|
|
4195
|
+
retry: Retry.Retry,
|
|
4196
|
+
}));
|
|
4197
|
+
|
|
4198
|
+
export type ListAccessGroupsAccessGroupListGetError = LitellmOpError;
|
|
4199
|
+
/** List Access Groups List all access groups. Returns a list of all access groups with their model names, deployment counts, shared budget and the spend drawn against it. Example: ```bash curl -X GET 'http://localhost:4000/access_group/list' \ -H 'Authorization: Bearer sk-1234' ``` Returns: - ListAccessGroupsResponse with all access groups */
|
|
4200
|
+
export const listAccessGroupsAccessGroupListGet: API.OperationMethod<
|
|
4201
|
+
ListAccessGroupsAccessGroupListGetRequest,
|
|
4202
|
+
ListAccessGroupsResponse,
|
|
4203
|
+
ListAccessGroupsAccessGroupListGetError,
|
|
4204
|
+
LitellmOpContext
|
|
4205
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4206
|
+
input: ListAccessGroupsAccessGroupListGetRequest,
|
|
4207
|
+
output: ListAccessGroupsResponse,
|
|
4208
|
+
errors: [],
|
|
4209
|
+
protocol: LitellmProtocol,
|
|
4210
|
+
retry: Retry.Retry,
|
|
4211
|
+
}));
|
|
4212
|
+
|
|
4213
|
+
export type ModelListModelsGetError =
|
|
4214
|
+
| BadRequest
|
|
4215
|
+
| UnprocessableEntity
|
|
4216
|
+
| LitellmOpError;
|
|
4217
|
+
/** Model List Use `/model/info` - to get detailed model information, example - pricing, mode, etc. This is just for compatibility with openai projects like aider. Query Parameters: - include_metadata: Include additional metadata in the response with fallback information - fallback_type: Type of fallbacks to include ("general", "context_window", "content_policy") Defaults to "general" when include_metadata=true - scope: Optional scope parameter. Currently only accepts "expand". When scope=expand is passed, proxy admins, team admins, and org admins will receive all proxy models as if they are a proxy admin. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true` in general_settings, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Models expanded from wildcard routes (e.g. `openai/*`) are not filtered, and nothing is hidden when `allowed_fails_policy` is configured (cooldown remains the sole exclusion mechanism). Hiding is presentation-only: a hidden model can still be called directly. */
|
|
4218
|
+
export const modelListModelsGet: API.OperationMethod<
|
|
4219
|
+
ModelListModelsGetRequest,
|
|
4220
|
+
ModelListModelsGetResponse,
|
|
4221
|
+
ModelListModelsGetError,
|
|
4222
|
+
LitellmOpContext
|
|
4223
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4224
|
+
input: ModelListModelsGetRequest,
|
|
4225
|
+
output: ModelListModelsGetResponse,
|
|
4226
|
+
errors: [BadRequest, UnprocessableEntity],
|
|
4227
|
+
protocol: LitellmProtocol,
|
|
4228
|
+
retry: Retry.Retry,
|
|
4229
|
+
}));
|
|
4230
|
+
|
|
4231
|
+
export type ModelListV1ModelsGetError = UnprocessableEntity | LitellmOpError;
|
|
4232
|
+
/** Model List Use `/model/info` - to get detailed model information, example - pricing, mode, etc. This is just for compatibility with openai projects like aider. Query Parameters: - include_metadata: Include additional metadata in the response with fallback information - fallback_type: Type of fallbacks to include ("general", "context_window", "content_policy") Defaults to "general" when include_metadata=true - scope: Optional scope parameter. Currently only accepts "expand". When scope=expand is passed, proxy admins, team admins, and org admins will receive all proxy models as if they are a proxy admin. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true` in general_settings, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Models expanded from wildcard routes (e.g. `openai/*`) are not filtered, and nothing is hidden when `allowed_fails_policy` is configured (cooldown remains the sole exclusion mechanism). Hiding is presentation-only: a hidden model can still be called directly. */
|
|
4233
|
+
export const modelListV1ModelsGet: API.OperationMethod<
|
|
4234
|
+
ModelListV1ModelsGetRequest,
|
|
4235
|
+
ModelListV1ModelsGetResponse,
|
|
4236
|
+
ModelListV1ModelsGetError,
|
|
4237
|
+
LitellmOpContext
|
|
4238
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4239
|
+
input: ModelListV1ModelsGetRequest,
|
|
4240
|
+
output: ModelListV1ModelsGetResponse,
|
|
4241
|
+
errors: [UnprocessableEntity],
|
|
4242
|
+
protocol: LitellmProtocol,
|
|
4243
|
+
retry: Retry.Retry,
|
|
4244
|
+
}));
|
|
4245
|
+
|
|
4246
|
+
export type PatchModelModelModelIdUpdatePatchError =
|
|
4247
|
+
| UnprocessableEntity
|
|
4248
|
+
| LitellmOpError;
|
|
4249
|
+
/** Patch Model PATCH Endpoint for partial model updates. Only updates the fields specified in the request while preserving other existing values. Follows proper PATCH semantics by only modifying provided fields. Args: model_id: The ID of the model to update patch_data: The fields to update and their new values user_api_key_dict: User authentication information Returns: Updated model information Raises: ProxyException: For various error conditions including authentication and database errors */
|
|
4250
|
+
export const patchModelModelModelIdUpdatePatch: API.OperationMethod<
|
|
4251
|
+
PatchModelModelModelIdUpdatePatchRequest,
|
|
4252
|
+
PatchModelModelModelIdUpdatePatchResponse,
|
|
4253
|
+
PatchModelModelModelIdUpdatePatchError,
|
|
4254
|
+
LitellmOpContext
|
|
4255
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4256
|
+
input: PatchModelModelModelIdUpdatePatchRequest,
|
|
4257
|
+
output: PatchModelModelModelIdUpdatePatchResponse,
|
|
4258
|
+
errors: [UnprocessableEntity],
|
|
4259
|
+
protocol: LitellmProtocol,
|
|
4260
|
+
retry: Retry.Retry,
|
|
4261
|
+
}));
|
|
4262
|
+
|
|
4263
|
+
export type PostBlockModelModelBlockError =
|
|
4264
|
+
| UnprocessableEntity
|
|
4265
|
+
| LitellmOpError;
|
|
4266
|
+
/** Block Model Block a DB-stored model deployment from serving requests. Parameters: - model_id: str - The model deployment id to block. */
|
|
4267
|
+
export const postBlockModelModelBlock: API.OperationMethod<
|
|
4268
|
+
PostBlockModelModelBlockRequest,
|
|
4269
|
+
PostBlockModelModelBlockResponse,
|
|
4270
|
+
PostBlockModelModelBlockError,
|
|
4271
|
+
LitellmOpContext
|
|
4272
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4273
|
+
input: PostBlockModelModelBlockRequest,
|
|
4274
|
+
output: PostBlockModelModelBlockResponse,
|
|
4275
|
+
errors: [UnprocessableEntity],
|
|
4276
|
+
protocol: LitellmProtocol,
|
|
4277
|
+
retry: Retry.Retry,
|
|
4278
|
+
}));
|
|
4279
|
+
|
|
4280
|
+
export type PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptError =
|
|
4281
|
+
| UnprocessableEntity
|
|
4282
|
+
| LitellmOpError;
|
|
4283
|
+
/** Preview Auto Router Classifier Prompt Get the system prompt an auto-router's LLM classifier sends for an edited tier set */
|
|
4284
|
+
export const postPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPrompt: API.OperationMethod<
|
|
4285
|
+
PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest,
|
|
4286
|
+
AutoRouterClassifierDefaultPromptResponse,
|
|
4287
|
+
PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptError,
|
|
4288
|
+
LitellmOpContext
|
|
4289
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4290
|
+
input:
|
|
4291
|
+
PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest,
|
|
4292
|
+
output: AutoRouterClassifierDefaultPromptResponse,
|
|
4293
|
+
errors: [UnprocessableEntity],
|
|
4294
|
+
protocol: LitellmProtocol,
|
|
4295
|
+
retry: Retry.Retry,
|
|
4296
|
+
}));
|
|
4297
|
+
|
|
4298
|
+
export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingError =
|
|
4299
|
+
| UnprocessableEntity
|
|
4300
|
+
| LitellmOpError;
|
|
4301
|
+
/** Preview Auto Router Routing Route a single request through a complexity-router config and report where it landed. Answers "which model would this request get?" for a config that only exists in a form, so an auto router can be checked before it is created. The request is classified by the same pre-routing hook a live request runs, over the same messages, system prompt and tool definitions, then dropped: nothing is sent to the model it routed to, and no auto router is created. A heuristic config therefore spends nothing, while an `llm` classifier or semantic keyword matching bills its classifier/embedding call to the calling key, like Test Connection does. Send `messages` to classify a real turn, with `system` and `tools` beside it when the surface carries them top level, as Anthropic /v1/messages does. `prompt` is the single-ask shorthand and routes as one user turn with nothing around it. **Example Request:** ```json { "messages": [ {"role": "system", "content": "You are a database migration assistant"}, {"role": "user", "content": "the index is not unique"}, {"role": "assistant", "content": "Then two workers can both insert. Add a unique index"}, {"role": "user", "content": "ok do it"} ], "tools": [{"type": "function", "function": {"name": "Bash", "description": "Run a command"}}], "complexity_router_config": { "tiers": {"SIMPLE": ["gpt-4o-mini"], "REASONING": ["o3"]}, "classifier_type": "heuristic" } } ``` */
|
|
4302
|
+
export const postPreviewAutoRouterRoutingAutoRouterTestRouting: API.OperationMethod<
|
|
4303
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest,
|
|
4304
|
+
AutoRouterRoutingTestResponse,
|
|
4305
|
+
PostPreviewAutoRouterRoutingAutoRouterTestRoutingError,
|
|
4306
|
+
LitellmOpContext
|
|
4307
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4308
|
+
input: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest,
|
|
4309
|
+
output: AutoRouterRoutingTestResponse,
|
|
4310
|
+
errors: [UnprocessableEntity],
|
|
4311
|
+
protocol: LitellmProtocol,
|
|
4312
|
+
retry: Retry.Retry,
|
|
4313
|
+
}));
|
|
4314
|
+
|
|
4315
|
+
export type PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderError =
|
|
4316
|
+
LitellmOpError;
|
|
4317
|
+
/** Reload Anthropic Beta Headers ADMIN ONLY / MASTER KEY Only Endpoint Manually reload the Anthropic beta headers configuration from the remote source. This will fetch fresh configuration from the anthropic_beta_headers_config.json file. */
|
|
4318
|
+
export const postReloadAnthropicBetaHeadersReloadAnthropicBetaHeader: API.OperationMethod<
|
|
4319
|
+
PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest,
|
|
4320
|
+
PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse,
|
|
4321
|
+
PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderError,
|
|
4322
|
+
LitellmOpContext
|
|
4323
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4324
|
+
input: PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest,
|
|
4325
|
+
output: PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse,
|
|
4326
|
+
errors: [],
|
|
4327
|
+
protocol: LitellmProtocol,
|
|
4328
|
+
retry: Retry.Retry,
|
|
4329
|
+
}));
|
|
4330
|
+
|
|
4331
|
+
export type PostReloadModelCostMapReloadModelCostMapError = LitellmOpError;
|
|
4332
|
+
/** Reload Model Cost Map ADMIN ONLY / MASTER KEY Only Endpoint Manually reload the model cost map from the remote source. This will fetch fresh pricing data from the model_prices_and_context_window.json file. */
|
|
4333
|
+
export const postReloadModelCostMapReloadModelCostMap: API.OperationMethod<
|
|
4334
|
+
PostReloadModelCostMapReloadModelCostMapRequest,
|
|
4335
|
+
PostReloadModelCostMapReloadModelCostMapResponse,
|
|
4336
|
+
PostReloadModelCostMapReloadModelCostMapError,
|
|
4337
|
+
LitellmOpContext
|
|
4338
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4339
|
+
input: PostReloadModelCostMapReloadModelCostMapRequest,
|
|
4340
|
+
output: PostReloadModelCostMapReloadModelCostMapResponse,
|
|
4341
|
+
errors: [],
|
|
4342
|
+
protocol: LitellmProtocol,
|
|
4343
|
+
retry: Retry.Retry,
|
|
4344
|
+
}));
|
|
4345
|
+
|
|
4346
|
+
export type PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadError =
|
|
4347
|
+
| UnprocessableEntity
|
|
4348
|
+
| LitellmOpError;
|
|
4349
|
+
/** Schedule Anthropic Beta Headers Reload ADMIN ONLY / MASTER KEY Only Endpoint Schedule periodic reload of the Anthropic beta headers configuration. This will create a background job that reloads the configuration every specified hours. */
|
|
4350
|
+
export const postScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReload: API.OperationMethod<
|
|
4351
|
+
PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest,
|
|
4352
|
+
PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse,
|
|
4353
|
+
PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadError,
|
|
4354
|
+
LitellmOpContext
|
|
4355
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4356
|
+
input:
|
|
4357
|
+
PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest,
|
|
4358
|
+
output:
|
|
4359
|
+
PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse,
|
|
4360
|
+
errors: [UnprocessableEntity],
|
|
4361
|
+
protocol: LitellmProtocol,
|
|
4362
|
+
retry: Retry.Retry,
|
|
4363
|
+
}));
|
|
4364
|
+
|
|
4365
|
+
export type PostScheduleModelCostMapReloadScheduleModelCostMapReloadError =
|
|
4366
|
+
| UnprocessableEntity
|
|
4367
|
+
| LitellmOpError;
|
|
4368
|
+
/** Schedule Model Cost Map Reload ADMIN ONLY / MASTER KEY Only Endpoint Schedule periodic reload of the model cost map. This will create a background job that reloads the model cost map every specified hours. */
|
|
4369
|
+
export const postScheduleModelCostMapReloadScheduleModelCostMapReload: API.OperationMethod<
|
|
4370
|
+
PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest,
|
|
4371
|
+
PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse,
|
|
4372
|
+
PostScheduleModelCostMapReloadScheduleModelCostMapReloadError,
|
|
4373
|
+
LitellmOpContext
|
|
4374
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4375
|
+
input: PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest,
|
|
4376
|
+
output: PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse,
|
|
4377
|
+
errors: [UnprocessableEntity],
|
|
4378
|
+
protocol: LitellmProtocol,
|
|
4379
|
+
retry: Retry.Retry,
|
|
4380
|
+
}));
|
|
4381
|
+
|
|
4382
|
+
export type PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetError =
|
|
4383
|
+
| UnprocessableEntity
|
|
4384
|
+
| LitellmOpError;
|
|
4385
|
+
/** Set Access Group Budget Set or replace the shared budget of an access group. Idempotent. Every key that can reach a model in the group draws from this one budget. Example: ```bash curl -X PUT 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "max_budget": 100.0, "budget_duration": "30d" }' ``` Parameters: - access_group: str - The access group name (URL path parameter) - max_budget: Optional[float] - Requests fail once the group's shared spend exceeds this - soft_budget: Optional[float] - Fires an alert when reached; requests still succeed - budget_duration: Optional[str] - Frequency of resetting the group's spend (e.g. '30d') - budget_id: Optional[str] - Link an existing budget instead of creating one Returns: - AccessGroupBudgetResponse with the stored budget and current spend Raises: - HTTPException 400: If no budget field is given, or budget_duration cannot be parsed - HTTPException 404: If access group not found */
|
|
4386
|
+
export const putSetAccessGroupBudgetAccessGroupAccessGroupBudget: API.OperationMethod<
|
|
4387
|
+
PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest,
|
|
4388
|
+
AccessGroupBudgetResponse,
|
|
4389
|
+
PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetError,
|
|
4390
|
+
LitellmOpContext
|
|
4391
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4392
|
+
input: PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest,
|
|
4393
|
+
output: AccessGroupBudgetResponse,
|
|
4394
|
+
errors: [UnprocessableEntity],
|
|
4395
|
+
protocol: LitellmProtocol,
|
|
4396
|
+
retry: Retry.Retry,
|
|
4397
|
+
}));
|
|
4398
|
+
|
|
4399
|
+
export type UnblockModelModelUnblockPostError =
|
|
4400
|
+
| UnprocessableEntity
|
|
4401
|
+
| LitellmOpError;
|
|
4402
|
+
/** Unblock Model Unblock a DB-stored model deployment so it can serve requests again. Parameters: - model_id: str - The model deployment id to unblock. */
|
|
4403
|
+
export const unblockModelModelUnblockPost: API.OperationMethod<
|
|
4404
|
+
UnblockModelModelUnblockPostRequest,
|
|
4405
|
+
UnblockModelModelUnblockPostResponse,
|
|
4406
|
+
UnblockModelModelUnblockPostError,
|
|
4407
|
+
LitellmOpContext
|
|
4408
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4409
|
+
input: UnblockModelModelUnblockPostRequest,
|
|
4410
|
+
output: UnblockModelModelUnblockPostResponse,
|
|
4411
|
+
errors: [UnprocessableEntity],
|
|
4412
|
+
protocol: LitellmProtocol,
|
|
4413
|
+
retry: Retry.Retry,
|
|
4414
|
+
}));
|
|
4415
|
+
|
|
4416
|
+
export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutError =
|
|
4417
|
+
| UnprocessableEntity
|
|
4418
|
+
| LitellmOpError;
|
|
4419
|
+
/** Update Access Group Update an access group's model names. This will: 1. Remove the access group from all current deployments 2. Add the access group to all deployments for the new model_names list Example: ```bash curl -X PUT 'http://localhost:4000/access_group/production-models/update' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "model_names": ["gpt-4", "claude-3-sonnet"] }' ``` Parameters: - access_group: str - The access group name (URL path parameter) - model_names: List[str] - New list of model groups to include Returns: - NewModelGroupResponse with the updated access group details Raises: - HTTPException 400: If any model names don't exist - HTTPException 404: If access group not found */
|
|
4420
|
+
export const updateAccessGroupAccessGroupAccessGroupUpdatePut: API.OperationMethod<
|
|
4421
|
+
UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest,
|
|
4422
|
+
NewModelGroupResponse,
|
|
4423
|
+
UpdateAccessGroupAccessGroupAccessGroupUpdatePutError,
|
|
4424
|
+
LitellmOpContext
|
|
4425
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4426
|
+
input: UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest,
|
|
4427
|
+
output: NewModelGroupResponse,
|
|
4428
|
+
errors: [UnprocessableEntity],
|
|
4429
|
+
protocol: LitellmProtocol,
|
|
4430
|
+
retry: Retry.Retry,
|
|
4431
|
+
}));
|
|
4432
|
+
|
|
4433
|
+
export type UpdateModelModelUpdatePostError =
|
|
4434
|
+
| BadRequest
|
|
4435
|
+
| UnprocessableEntity
|
|
4436
|
+
| LitellmOpError;
|
|
4437
|
+
/** Update Model Edit existing model params */
|
|
4438
|
+
export const updateModelModelUpdatePost: API.OperationMethod<
|
|
4439
|
+
UpdateModelModelUpdatePostRequest,
|
|
4440
|
+
UpdateModelModelUpdatePostResponse,
|
|
4441
|
+
UpdateModelModelUpdatePostError,
|
|
4442
|
+
LitellmOpContext
|
|
4443
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4444
|
+
input: UpdateModelModelUpdatePostRequest,
|
|
4445
|
+
output: UpdateModelModelUpdatePostResponse,
|
|
4446
|
+
errors: [BadRequest, UnprocessableEntity],
|
|
4447
|
+
protocol: LitellmProtocol,
|
|
4448
|
+
retry: Retry.Retry,
|
|
4449
|
+
}));
|
|
4450
|
+
|
|
4451
|
+
export type UpdatePublicModelGroupsModelGroupMakePublicPostError =
|
|
4452
|
+
| UnprocessableEntity
|
|
4453
|
+
| LitellmOpError;
|
|
4454
|
+
/** Update Public Model Groups Update which model groups are public */
|
|
4455
|
+
export const updatePublicModelGroupsModelGroupMakePublicPost: API.OperationMethod<
|
|
4456
|
+
UpdatePublicModelGroupsModelGroupMakePublicPostRequest,
|
|
4457
|
+
UpdatePublicModelGroupsModelGroupMakePublicPostResponse,
|
|
4458
|
+
UpdatePublicModelGroupsModelGroupMakePublicPostError,
|
|
4459
|
+
LitellmOpContext
|
|
4460
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4461
|
+
input: UpdatePublicModelGroupsModelGroupMakePublicPostRequest,
|
|
4462
|
+
output: UpdatePublicModelGroupsModelGroupMakePublicPostResponse,
|
|
4463
|
+
errors: [UnprocessableEntity],
|
|
4464
|
+
protocol: LitellmProtocol,
|
|
4465
|
+
retry: Retry.Retry,
|
|
4466
|
+
}));
|
|
4467
|
+
|
|
4468
|
+
export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostError =
|
|
4469
|
+
| UnprocessableEntity
|
|
4470
|
+
| LitellmOpError;
|
|
4471
|
+
/** Update Useful Links Update useful links */
|
|
4472
|
+
export const updateUsefulLinksModelHubUpdateUsefulLinksPost: API.OperationMethod<
|
|
4473
|
+
UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest,
|
|
4474
|
+
UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse,
|
|
4475
|
+
UpdateUsefulLinksModelHubUpdateUsefulLinksPostError,
|
|
4476
|
+
LitellmOpContext
|
|
4477
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4478
|
+
input: UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest,
|
|
4479
|
+
output: UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse,
|
|
4480
|
+
errors: [UnprocessableEntity],
|
|
4481
|
+
protocol: LitellmProtocol,
|
|
4482
|
+
retry: Retry.Retry,
|
|
4483
|
+
}));
|
|
4484
|
+
|
|
4485
|
+
export type ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostError =
|
|
4486
|
+
| UnprocessableEntity
|
|
4487
|
+
| LitellmOpError;
|
|
4488
|
+
/** Validate Complexity Router Config Validate a complexity-router config without saving it. Runs the same check every write path runs (the router's own pydantic model), so a form can show the backend's exact verdict while the operator is still editing rather than after a rejected save. Gated exactly like the save it rehearses: a proxy admin, or a team admin naming their own team. Nothing is created, routed, or billed. */
|
|
4489
|
+
export const validateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPost: API.OperationMethod<
|
|
4490
|
+
ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest,
|
|
4491
|
+
ComplexityRouterConfigValidationResponse,
|
|
4492
|
+
ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostError,
|
|
4493
|
+
LitellmOpContext
|
|
4494
|
+
> = /*@__PURE__*/ API.make(() => ({
|
|
4495
|
+
input:
|
|
4496
|
+
ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest,
|
|
4497
|
+
output: ComplexityRouterConfigValidationResponse,
|
|
4498
|
+
errors: [UnprocessableEntity],
|
|
4499
|
+
protocol: LitellmProtocol,
|
|
4500
|
+
retry: Retry.Retry,
|
|
4501
|
+
}));
|