@homeflare/distilled-litellm 0.2.0 → 0.3.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +41 -12
- package/dist/services/a2a.js +1 -1
- package/dist/services/a2a_registration.js +1 -1
- package/dist/services/access_groups.d.ts +18 -0
- package/dist/services/access_groups.d.ts.map +1 -1
- package/dist/services/access_groups.js +15 -1
- package/dist/services/access_groups.js.map +1 -1
- package/dist/services/adaptive_router.js +1 -1
- package/dist/services/agents.d.ts +10 -0
- package/dist/services/agents.d.ts.map +1 -1
- package/dist/services/agents.js +10 -1
- package/dist/services/agents.js.map +1 -1
- package/dist/services/alerting.js +1 -1
- package/dist/services/anthropic_passthrough.js +1 -1
- package/dist/services/anthropic_skills.d.ts +7 -1
- package/dist/services/anthropic_skills.d.ts.map +1 -1
- package/dist/services/anthropic_skills.js +6 -2
- package/dist/services/anthropic_skills.js.map +1 -1
- package/dist/services/assistants.js +1 -1
- package/dist/services/audio.js +1 -1
- package/dist/services/audit_logging.d.ts +2 -0
- package/dist/services/audit_logging.d.ts.map +1 -1
- package/dist/services/audit_logging.js +2 -1
- package/dist/services/audit_logging.js.map +1 -1
- package/dist/services/auto_router.d.ts +162 -62
- package/dist/services/auto_router.d.ts.map +1 -1
- package/dist/services/auto_router.js +92 -31
- package/dist/services/auto_router.js.map +1 -1
- package/dist/services/batch.js +1 -1
- package/dist/services/beta_agents.js +1 -1
- package/dist/services/beta_mcp.js +1 -1
- package/dist/services/budget_management.d.ts +8 -3
- package/dist/services/budget_management.d.ts.map +1 -1
- package/dist/services/budget_management.js +7 -4
- package/dist/services/budget_management.js.map +1 -1
- package/dist/services/budget_spend_tracking.d.ts +24 -5
- package/dist/services/budget_spend_tracking.d.ts.map +1 -1
- package/dist/services/budget_spend_tracking.js +17 -4
- package/dist/services/budget_spend_tracking.js.map +1 -1
- package/dist/services/cache_settings.js +1 -1
- package/dist/services/caching.js +1 -1
- package/dist/services/chat_completions.js +1 -1
- package/dist/services/claude_code_marketplace.d.ts +11 -10
- package/dist/services/claude_code_marketplace.d.ts.map +1 -1
- package/dist/services/claude_code_marketplace.js +8 -6
- package/dist/services/claude_code_marketplace.js.map +1 -1
- package/dist/services/cloudzero.js +1 -1
- package/dist/services/completions.js +1 -1
- package/dist/services/compliance.js +1 -1
- package/dist/services/config_overrides.d.ts +5 -1
- package/dist/services/config_overrides.d.ts.map +1 -1
- package/dist/services/config_overrides.js +3 -1
- package/dist/services/config_overrides.js.map +1 -1
- package/dist/services/config_yaml.d.ts +120 -6
- package/dist/services/config_yaml.d.ts.map +1 -1
- package/dist/services/config_yaml.js +67 -1
- package/dist/services/config_yaml.js.map +1 -1
- package/dist/services/containers.d.ts +8 -2
- package/dist/services/containers.d.ts.map +1 -1
- package/dist/services/containers.js +13 -5
- package/dist/services/containers.js.map +1 -1
- package/dist/services/coordination_redis_settings.js +1 -1
- package/dist/services/cost_tracking.d.ts +91 -1
- package/dist/services/cost_tracking.d.ts.map +1 -1
- package/dist/services/cost_tracking.js +82 -2
- package/dist/services/cost_tracking.js.map +1 -1
- package/dist/services/credential_management.d.ts +2 -1
- package/dist/services/credential_management.d.ts.map +1 -1
- package/dist/services/credential_management.js +3 -2
- package/dist/services/credential_management.js.map +1 -1
- package/dist/services/customer_management.d.ts +25 -2
- package/dist/services/customer_management.d.ts.map +1 -1
- package/dist/services/customer_management.js +22 -3
- package/dist/services/customer_management.js.map +1 -1
- package/dist/services/email_management.js +1 -1
- package/dist/services/embeddings.js +1 -1
- package/dist/services/evals.js +1 -1
- package/dist/services/experimental.js +1 -1
- package/dist/services/fallback_management.js +1 -1
- package/dist/services/files.js +1 -1
- package/dist/services/fine_tuning.js +1 -1
- package/dist/services/gemini_agents.js +1 -1
- package/dist/services/google_genai_endpoints.js +1 -1
- package/dist/services/guardrails.d.ts +80 -7
- package/dist/services/guardrails.d.ts.map +1 -1
- package/dist/services/guardrails.js +37 -1
- package/dist/services/guardrails.js.map +1 -1
- package/dist/services/health.d.ts +2 -2
- package/dist/services/health.d.ts.map +1 -1
- package/dist/services/health.js +2 -2
- package/dist/services/health.js.map +1 -1
- package/dist/services/images.js +1 -1
- package/dist/services/index.d.ts +1 -15
- package/dist/services/index.d.ts.map +1 -1
- package/dist/services/index.js +1 -15
- package/dist/services/index.js.map +1 -1
- package/dist/services/internal_user_management.d.ts +227 -39
- package/dist/services/internal_user_management.d.ts.map +1 -1
- package/dist/services/internal_user_management.js +194 -37
- package/dist/services/internal_user_management.js.map +1 -1
- package/dist/services/invite_links.js +1 -1
- package/dist/services/jwt_mappings.d.ts +3 -0
- package/dist/services/jwt_mappings.d.ts.map +1 -1
- package/dist/services/jwt_mappings.js +4 -1
- package/dist/services/jwt_mappings.js.map +1 -1
- package/dist/services/key_management.d.ts +93 -49
- package/dist/services/key_management.d.ts.map +1 -1
- package/dist/services/key_management.js +75 -39
- package/dist/services/key_management.js.map +1 -1
- package/dist/services/langfuse_passthrough.js +1 -1
- package/dist/services/llm_passthrough.d.ts +1199 -0
- package/dist/services/llm_passthrough.d.ts.map +1 -0
- package/dist/services/llm_passthrough.js +2150 -0
- package/dist/services/llm_passthrough.js.map +1 -0
- package/dist/services/llm_utils.js +1 -1
- package/dist/services/logging_callbacks.js +1 -1
- package/dist/services/mcp_byok_oauth.d.ts +18 -11
- package/dist/services/mcp_byok_oauth.d.ts.map +1 -1
- package/dist/services/mcp_byok_oauth.js +27 -24
- package/dist/services/mcp_byok_oauth.js.map +1 -1
- package/dist/services/mcp_discoverable.d.ts +34 -0
- package/dist/services/mcp_discoverable.d.ts.map +1 -1
- package/dist/services/mcp_discoverable.js +51 -1
- package/dist/services/mcp_discoverable.js.map +1 -1
- package/dist/services/mcp_management.d.ts +90 -2
- package/dist/services/mcp_management.d.ts.map +1 -1
- package/dist/services/mcp_management.js +115 -3
- package/dist/services/mcp_management.js.map +1 -1
- package/dist/services/mcp_rest.d.ts +13 -0
- package/dist/services/mcp_rest.d.ts.map +1 -1
- package/dist/services/mcp_rest.js +10 -2
- package/dist/services/mcp_rest.js.map +1 -1
- package/dist/services/memory_management.d.ts +2 -0
- package/dist/services/memory_management.d.ts.map +1 -1
- package/dist/services/memory_management.js +2 -1
- package/dist/services/memory_management.js.map +1 -1
- package/dist/services/misc.d.ts +86 -1
- package/dist/services/misc.d.ts.map +1 -1
- package/dist/services/misc.js +126 -2
- package/dist/services/misc.js.map +1 -1
- package/dist/services/model_management.d.ts +310 -19
- package/dist/services/model_management.d.ts.map +1 -1
- package/dist/services/model_management.js +240 -7
- package/dist/services/model_management.js.map +1 -1
- package/dist/services/moderations.js +1 -1
- package/dist/services/ocr.js +1 -1
- package/dist/services/open_ai_pass_through.d.ts +35 -90
- package/dist/services/open_ai_pass_through.d.ts.map +1 -1
- package/dist/services/open_ai_pass_through.js +41 -139
- package/dist/services/open_ai_pass_through.js.map +1 -1
- package/dist/services/organization_management.d.ts +21 -1
- package/dist/services/organization_management.d.ts.map +1 -1
- package/dist/services/organization_management.js +20 -2
- package/dist/services/organization_management.js.map +1 -1
- package/dist/services/plugins.js +1 -1
- package/dist/services/policies.d.ts +2 -0
- package/dist/services/policies.d.ts.map +1 -1
- package/dist/services/policies.js +3 -1
- package/dist/services/policies.js.map +1 -1
- package/dist/services/policy_engine.d.ts +6 -0
- package/dist/services/policy_engine.d.ts.map +1 -1
- package/dist/services/policy_engine.js +4 -1
- package/dist/services/policy_engine.js.map +1 -1
- package/dist/services/project_management.d.ts +15 -0
- package/dist/services/project_management.d.ts.map +1 -1
- package/dist/services/project_management.js +14 -1
- package/dist/services/project_management.js.map +1 -1
- package/dist/services/prompts.js +1 -1
- package/dist/services/public.d.ts +124 -0
- package/dist/services/public.d.ts.map +1 -1
- package/dist/services/public.js +158 -1
- package/dist/services/public.js.map +1 -1
- package/dist/services/rag.js +1 -1
- package/dist/services/realtime.d.ts +4 -26
- package/dist/services/realtime.d.ts.map +1 -1
- package/dist/services/realtime.js +7 -57
- package/dist/services/realtime.js.map +1 -1
- package/dist/services/rerank.d.ts +72 -0
- package/dist/services/rerank.d.ts.map +1 -1
- package/dist/services/rerank.js +61 -4
- package/dist/services/rerank.js.map +1 -1
- package/dist/services/responses.d.ts +30 -0
- package/dist/services/responses.d.ts.map +1 -1
- package/dist/services/responses.js +59 -1
- package/dist/services/responses.js.map +1 -1
- package/dist/services/router_settings.js +1 -1
- package/dist/services/rust_control_plane.js +1 -1
- package/dist/services/scim.d.ts +41 -1
- package/dist/services/scim.d.ts.map +1 -1
- package/dist/services/scim.js +60 -2
- package/dist/services/scim.js.map +1 -1
- package/dist/services/search.js +1 -1
- package/dist/services/search_tools.js +1 -1
- package/dist/services/settings.d.ts +88 -0
- package/dist/services/settings.d.ts.map +1 -1
- package/dist/services/settings.js +113 -1
- package/dist/services/settings.js.map +1 -1
- package/dist/services/sso_settings.d.ts +1 -1
- package/dist/services/sso_settings.d.ts.map +1 -1
- package/dist/services/sso_settings.js +1 -1
- package/dist/services/sso_settings.js.map +1 -1
- package/dist/services/tag_management.d.ts +7 -0
- package/dist/services/tag_management.d.ts.map +1 -1
- package/dist/services/tag_management.js +8 -1
- package/dist/services/tag_management.js.map +1 -1
- package/dist/services/team_management.d.ts +265 -17
- package/dist/services/team_management.d.ts.map +1 -1
- package/dist/services/team_management.js +247 -12
- package/dist/services/team_management.js.map +1 -1
- package/dist/services/tools.js +1 -1
- package/dist/services/ui_settings.js +1 -1
- package/dist/services/ui_theme_settings.js +1 -1
- package/dist/services/usage_ai.js +1 -1
- package/dist/services/vantage.js +1 -1
- package/dist/services/vector_store_management.js +1 -1
- package/dist/services/vector_stores.js +1 -1
- package/dist/services/videos.js +1 -1
- package/dist/services/web_socket.d.ts +0 -18
- package/dist/services/web_socket.d.ts.map +1 -1
- package/dist/services/web_socket.js +1 -33
- package/dist/services/web_socket.js.map +1 -1
- package/dist/services/workflow_management.js +1 -1
- package/package.json +1 -1
- package/src/services/a2a.ts +1 -1
- package/src/services/a2a_registration.ts +1 -1
- package/src/services/access_groups.ts +44 -1
- package/src/services/adaptive_router.ts +1 -1
- package/src/services/agents.ts +22 -1
- package/src/services/alerting.ts +1 -1
- package/src/services/anthropic_passthrough.ts +1 -1
- package/src/services/anthropic_skills.ts +12 -2
- package/src/services/assistants.ts +1 -1
- package/src/services/audio.ts +1 -1
- package/src/services/audit_logging.ts +4 -1
- package/src/services/auto_router.ts +303 -93
- package/src/services/batch.ts +1 -1
- package/src/services/beta_agents.ts +1 -1
- package/src/services/beta_mcp.ts +1 -1
- package/src/services/budget_management.ts +12 -4
- package/src/services/budget_spend_tracking.ts +38 -6
- package/src/services/cache_settings.ts +1 -1
- package/src/services/caching.ts +1 -1
- package/src/services/chat_completions.ts +1 -1
- package/src/services/claude_code_marketplace.ts +20 -14
- package/src/services/cloudzero.ts +1 -1
- package/src/services/completions.ts +1 -1
- package/src/services/compliance.ts +1 -1
- package/src/services/config_overrides.ts +8 -2
- package/src/services/config_yaml.ts +251 -6
- package/src/services/containers.ts +29 -11
- package/src/services/coordination_redis_settings.ts +1 -1
- package/src/services/cost_tracking.ts +198 -2
- package/src/services/credential_management.ts +9 -4
- package/src/services/customer_management.ts +49 -3
- package/src/services/email_management.ts +1 -1
- package/src/services/embeddings.ts +1 -1
- package/src/services/evals.ts +1 -1
- package/src/services/experimental.ts +1 -1
- package/src/services/fallback_management.ts +1 -1
- package/src/services/files.ts +1 -1
- package/src/services/fine_tuning.ts +1 -1
- package/src/services/gemini_agents.ts +1 -1
- package/src/services/google_genai_endpoints.ts +1 -1
- package/src/services/guardrails.ts +198 -8
- package/src/services/health.ts +3 -2
- package/src/services/images.ts +1 -1
- package/src/services/index.ts +1 -15
- package/src/services/internal_user_management.ts +565 -103
- package/src/services/invite_links.ts +1 -1
- package/src/services/jwt_mappings.ts +7 -1
- package/src/services/key_management.ts +229 -126
- package/src/services/langfuse_passthrough.ts +1 -1
- package/src/services/llm_passthrough.ts +4542 -0
- package/src/services/llm_utils.ts +1 -1
- package/src/services/logging_callbacks.ts +1 -1
- package/src/services/mcp_byok_oauth.ts +65 -45
- package/src/services/mcp_discoverable.ts +109 -1
- package/src/services/mcp_management.ts +258 -3
- package/src/services/mcp_rest.ts +30 -5
- package/src/services/memory_management.ts +4 -1
- package/src/services/misc.ts +269 -2
- package/src/services/model_management.ts +666 -21
- package/src/services/moderations.ts +1 -1
- package/src/services/ocr.ts +1 -1
- package/src/services/open_ai_pass_through.ts +81 -286
- package/src/services/organization_management.ts +44 -2
- package/src/services/plugins.ts +1 -1
- package/src/services/policies.ts +5 -1
- package/src/services/policy_engine.ts +10 -1
- package/src/services/project_management.ts +33 -1
- package/src/services/prompts.ts +1 -1
- package/src/services/public.ts +356 -1
- package/src/services/rag.ts +1 -1
- package/src/services/realtime.ts +11 -111
- package/src/services/rerank.ts +187 -7
- package/src/services/responses.ts +117 -1
- package/src/services/router_settings.ts +1 -1
- package/src/services/rust_control_plane.ts +1 -1
- package/src/services/scim.ts +134 -3
- package/src/services/search.ts +1 -1
- package/src/services/search_tools.ts +1 -1
- package/src/services/settings.ts +278 -1
- package/src/services/sso_settings.ts +2 -1
- package/src/services/tag_management.ts +15 -1
- package/src/services/team_management.ts +676 -37
- package/src/services/tools.ts +1 -1
- package/src/services/ui_settings.ts +1 -1
- package/src/services/ui_theme_settings.ts +1 -1
- package/src/services/usage_ai.ts +1 -1
- package/src/services/vantage.ts +1 -1
- package/src/services/vector_store_management.ts +1 -1
- package/src/services/vector_stores.ts +1 -1
- package/src/services/videos.ts +1 -1
- package/src/services/web_socket.ts +1 -61
- package/src/services/workflow_management.ts +1 -1
- package/dist/services/anthropic_pass_through.d.ts +0 -70
- package/dist/services/anthropic_pass_through.d.ts.map +0 -1
- package/dist/services/anthropic_pass_through.js +0 -114
- package/dist/services/anthropic_pass_through.js.map +0 -1
- package/dist/services/assembly_ai_eu_pass_through.d.ts +0 -70
- package/dist/services/assembly_ai_eu_pass_through.d.ts.map +0 -1
- package/dist/services/assembly_ai_eu_pass_through.js +0 -114
- package/dist/services/assembly_ai_eu_pass_through.js.map +0 -1
- package/dist/services/assembly_ai_pass_through.d.ts +0 -70
- package/dist/services/assembly_ai_pass_through.d.ts.map +0 -1
- package/dist/services/assembly_ai_pass_through.js +0 -114
- package/dist/services/assembly_ai_pass_through.js.map +0 -1
- package/dist/services/aws_comprehend_medical_pass_through.d.ts +0 -36
- package/dist/services/aws_comprehend_medical_pass_through.d.ts.map +0 -1
- package/dist/services/aws_comprehend_medical_pass_through.js +0 -56
- package/dist/services/aws_comprehend_medical_pass_through.js.map +0 -1
- package/dist/services/azure_ai_pass_through.d.ts +0 -70
- package/dist/services/azure_ai_pass_through.d.ts.map +0 -1
- package/dist/services/azure_ai_pass_through.js +0 -112
- package/dist/services/azure_ai_pass_through.js.map +0 -1
- package/dist/services/azure_pass_through.d.ts +0 -70
- package/dist/services/azure_pass_through.d.ts.map +0 -1
- package/dist/services/azure_pass_through.js +0 -107
- package/dist/services/azure_pass_through.js.map +0 -1
- package/dist/services/bedrock_pass_through.d.ts +0 -70
- package/dist/services/bedrock_pass_through.d.ts.map +0 -1
- package/dist/services/bedrock_pass_through.js +0 -114
- package/dist/services/bedrock_pass_through.js.map +0 -1
- package/dist/services/cohere_pass_through.d.ts +0 -70
- package/dist/services/cohere_pass_through.d.ts.map +0 -1
- package/dist/services/cohere_pass_through.js +0 -112
- package/dist/services/cohere_pass_through.js.map +0 -1
- package/dist/services/cursor_pass_through.d.ts +0 -70
- package/dist/services/cursor_pass_through.d.ts.map +0 -1
- package/dist/services/cursor_pass_through.js +0 -112
- package/dist/services/cursor_pass_through.js.map +0 -1
- package/dist/services/google_ai_studio_pass_through.d.ts +0 -70
- package/dist/services/google_ai_studio_pass_through.d.ts.map +0 -1
- package/dist/services/google_ai_studio_pass_through.js +0 -112
- package/dist/services/google_ai_studio_pass_through.js.map +0 -1
- package/dist/services/milvus_pass_through.d.ts +0 -70
- package/dist/services/milvus_pass_through.d.ts.map +0 -1
- package/dist/services/milvus_pass_through.js +0 -112
- package/dist/services/milvus_pass_through.js.map +0 -1
- package/dist/services/mistral_pass_through.d.ts +0 -70
- package/dist/services/mistral_pass_through.d.ts.map +0 -1
- package/dist/services/mistral_pass_through.js +0 -114
- package/dist/services/mistral_pass_through.js.map +0 -1
- package/dist/services/vertex_ai_pass_through.d.ts +0 -180
- package/dist/services/vertex_ai_pass_through.d.ts.map +0 -1
- package/dist/services/vertex_ai_pass_through.js +0 -334
- package/dist/services/vertex_ai_pass_through.js.map +0 -1
- package/dist/services/vllm_pass_through.d.ts +0 -70
- package/dist/services/vllm_pass_through.d.ts.map +0 -1
- package/dist/services/vllm_pass_through.js +0 -104
- package/dist/services/vllm_pass_through.js.map +0 -1
- package/dist/services/watsonx_pass_through.d.ts +0 -70
- package/dist/services/watsonx_pass_through.d.ts.map +0 -1
- package/dist/services/watsonx_pass_through.js +0 -114
- package/dist/services/watsonx_pass_through.js.map +0 -1
- package/src/services/anthropic_pass_through.ts +0 -233
- package/src/services/assembly_ai_eu_pass_through.ts +0 -237
- package/src/services/assembly_ai_pass_through.ts +0 -237
- package/src/services/aws_comprehend_medical_pass_through.ts +0 -109
- package/src/services/azure_ai_pass_through.ts +0 -231
- package/src/services/azure_pass_through.ts +0 -227
- package/src/services/bedrock_pass_through.ts +0 -229
- package/src/services/cohere_pass_through.ts +0 -227
- package/src/services/cursor_pass_through.ts +0 -227
- package/src/services/google_ai_studio_pass_through.ts +0 -227
- package/src/services/milvus_pass_through.ts +0 -227
- package/src/services/mistral_pass_through.ts +0 -229
- package/src/services/vertex_ai_pass_through.ts +0 -684
- package/src/services/vllm_pass_through.ts +0 -227
- package/src/services/watsonx_pass_through.ts +0 -229
|
@@ -43,13 +43,55 @@ declare const UnprocessableEntity_base: S.Class<UnprocessableEntity, S.TaggedStr
|
|
|
43
43
|
});
|
|
44
44
|
export declare class UnprocessableEntity extends /*@__PURE__*/ UnprocessableEntity_base {
|
|
45
45
|
}
|
|
46
|
+
export type LiteLLMObjectPermissionBaseAgentAccessGroupsList = Array<string>;
|
|
47
|
+
export declare const LiteLLMObjectPermissionBaseAgentAccessGroupsList: S.Schema<LiteLLMObjectPermissionBaseAgentAccessGroupsList>;
|
|
48
|
+
export type LiteLLMObjectPermissionBaseAgentsList = Array<string>;
|
|
49
|
+
export declare const LiteLLMObjectPermissionBaseAgentsList: S.Schema<LiteLLMObjectPermissionBaseAgentsList>;
|
|
50
|
+
export type LiteLLMObjectPermissionBaseBlockedToolsList = Array<string>;
|
|
51
|
+
export declare const LiteLLMObjectPermissionBaseBlockedToolsList: S.Schema<LiteLLMObjectPermissionBaseBlockedToolsList>;
|
|
52
|
+
export type LiteLLMObjectPermissionBaseMcpAccessGroupsList = Array<string>;
|
|
53
|
+
export declare const LiteLLMObjectPermissionBaseMcpAccessGroupsList: S.Schema<LiteLLMObjectPermissionBaseMcpAccessGroupsList>;
|
|
54
|
+
export type LiteLLMObjectPermissionBaseMcpServersList = Array<string>;
|
|
55
|
+
export declare const LiteLLMObjectPermissionBaseMcpServersList: S.Schema<LiteLLMObjectPermissionBaseMcpServersList>;
|
|
56
|
+
export type LiteLLMObjectPermissionBaseMcpToolPermissionsValueList = Array<string>;
|
|
57
|
+
export declare const LiteLLMObjectPermissionBaseMcpToolPermissionsValueList: S.Schema<LiteLLMObjectPermissionBaseMcpToolPermissionsValueList>;
|
|
58
|
+
export type LiteLLMObjectPermissionBaseMcpToolPermissionsMap = {
|
|
59
|
+
[key: string]: LiteLLMObjectPermissionBaseMcpToolPermissionsValueList | undefined;
|
|
60
|
+
};
|
|
61
|
+
export declare const LiteLLMObjectPermissionBaseMcpToolPermissionsMap: S.Schema<LiteLLMObjectPermissionBaseMcpToolPermissionsMap>;
|
|
62
|
+
export type LiteLLMObjectPermissionBaseMcpToolsetsList = Array<string>;
|
|
63
|
+
export declare const LiteLLMObjectPermissionBaseMcpToolsetsList: S.Schema<LiteLLMObjectPermissionBaseMcpToolsetsList>;
|
|
64
|
+
export type LiteLLMObjectPermissionBaseModelsList = Array<string>;
|
|
65
|
+
export declare const LiteLLMObjectPermissionBaseModelsList: S.Schema<LiteLLMObjectPermissionBaseModelsList>;
|
|
66
|
+
export type LiteLLMObjectPermissionBaseSearchToolsList = Array<string>;
|
|
67
|
+
export declare const LiteLLMObjectPermissionBaseSearchToolsList: S.Schema<LiteLLMObjectPermissionBaseSearchToolsList>;
|
|
68
|
+
export type LiteLLMObjectPermissionBaseSkillsList = Array<string>;
|
|
69
|
+
export declare const LiteLLMObjectPermissionBaseSkillsList: S.Schema<LiteLLMObjectPermissionBaseSkillsList>;
|
|
70
|
+
export type LiteLLMObjectPermissionBaseVectorStoresList = Array<string>;
|
|
71
|
+
export declare const LiteLLMObjectPermissionBaseVectorStoresList: S.Schema<LiteLLMObjectPermissionBaseVectorStoresList>;
|
|
72
|
+
export interface LiteLLMObjectPermissionBase {
|
|
73
|
+
agent_access_groups?: LiteLLMObjectPermissionBaseAgentAccessGroupsList | null;
|
|
74
|
+
agents?: LiteLLMObjectPermissionBaseAgentsList | null;
|
|
75
|
+
blocked_tools?: LiteLLMObjectPermissionBaseBlockedToolsList | null;
|
|
76
|
+
mcp_access_groups?: LiteLLMObjectPermissionBaseMcpAccessGroupsList | null;
|
|
77
|
+
mcp_servers?: LiteLLMObjectPermissionBaseMcpServersList | null;
|
|
78
|
+
mcp_tool_permissions?: LiteLLMObjectPermissionBaseMcpToolPermissionsMap | null;
|
|
79
|
+
mcp_tool_search_enabled?: boolean | null;
|
|
80
|
+
mcp_toolsets?: LiteLLMObjectPermissionBaseMcpToolsetsList | null;
|
|
81
|
+
models?: LiteLLMObjectPermissionBaseModelsList | null;
|
|
82
|
+
search_tools?: LiteLLMObjectPermissionBaseSearchToolsList | null;
|
|
83
|
+
skills?: LiteLLMObjectPermissionBaseSkillsList | null;
|
|
84
|
+
vector_stores?: LiteLLMObjectPermissionBaseVectorStoresList | null;
|
|
85
|
+
}
|
|
86
|
+
export declare const LiteLLMObjectPermissionBase: S.Schema<LiteLLMObjectPermissionBase>;
|
|
46
87
|
export type BulkUpdateKeyRequestItemTagsList = Array<string>;
|
|
47
88
|
export declare const BulkUpdateKeyRequestItemTagsList: S.Schema<BulkUpdateKeyRequestItemTagsList>;
|
|
48
|
-
/**
|
|
89
|
+
/** One /key/bulk_update item; only the fields it carries are written. */
|
|
49
90
|
export interface BulkUpdateKeyRequestItem {
|
|
50
91
|
budget_id?: string | null;
|
|
51
92
|
key: string;
|
|
52
93
|
max_budget?: number | null;
|
|
94
|
+
object_permission?: LiteLLMObjectPermissionBase | null;
|
|
53
95
|
tags?: BulkUpdateKeyRequestItemTagsList | null;
|
|
54
96
|
team_id?: string | null;
|
|
55
97
|
}
|
|
@@ -234,44 +276,6 @@ export type GenerateKeyFnKeyGeneratePostRequestModelTpmLimitMap = {
|
|
|
234
276
|
export declare const GenerateKeyFnKeyGeneratePostRequestModelTpmLimitMap: S.Schema<GenerateKeyFnKeyGeneratePostRequestModelTpmLimitMap>;
|
|
235
277
|
export type GenerateKeyFnKeyGeneratePostRequestModelsList = Array<unknown>;
|
|
236
278
|
export declare const GenerateKeyFnKeyGeneratePostRequestModelsList: S.Schema<GenerateKeyFnKeyGeneratePostRequestModelsList>;
|
|
237
|
-
export type LiteLLMObjectPermissionBaseAgentAccessGroupsList = Array<string>;
|
|
238
|
-
export declare const LiteLLMObjectPermissionBaseAgentAccessGroupsList: S.Schema<LiteLLMObjectPermissionBaseAgentAccessGroupsList>;
|
|
239
|
-
export type LiteLLMObjectPermissionBaseAgentsList = Array<string>;
|
|
240
|
-
export declare const LiteLLMObjectPermissionBaseAgentsList: S.Schema<LiteLLMObjectPermissionBaseAgentsList>;
|
|
241
|
-
export type LiteLLMObjectPermissionBaseBlockedToolsList = Array<string>;
|
|
242
|
-
export declare const LiteLLMObjectPermissionBaseBlockedToolsList: S.Schema<LiteLLMObjectPermissionBaseBlockedToolsList>;
|
|
243
|
-
export type LiteLLMObjectPermissionBaseMcpAccessGroupsList = Array<string>;
|
|
244
|
-
export declare const LiteLLMObjectPermissionBaseMcpAccessGroupsList: S.Schema<LiteLLMObjectPermissionBaseMcpAccessGroupsList>;
|
|
245
|
-
export type LiteLLMObjectPermissionBaseMcpServersList = Array<string>;
|
|
246
|
-
export declare const LiteLLMObjectPermissionBaseMcpServersList: S.Schema<LiteLLMObjectPermissionBaseMcpServersList>;
|
|
247
|
-
export type LiteLLMObjectPermissionBaseMcpToolPermissionsValueList = Array<string>;
|
|
248
|
-
export declare const LiteLLMObjectPermissionBaseMcpToolPermissionsValueList: S.Schema<LiteLLMObjectPermissionBaseMcpToolPermissionsValueList>;
|
|
249
|
-
export type LiteLLMObjectPermissionBaseMcpToolPermissionsMap = {
|
|
250
|
-
[key: string]: LiteLLMObjectPermissionBaseMcpToolPermissionsValueList | undefined;
|
|
251
|
-
};
|
|
252
|
-
export declare const LiteLLMObjectPermissionBaseMcpToolPermissionsMap: S.Schema<LiteLLMObjectPermissionBaseMcpToolPermissionsMap>;
|
|
253
|
-
export type LiteLLMObjectPermissionBaseMcpToolsetsList = Array<string>;
|
|
254
|
-
export declare const LiteLLMObjectPermissionBaseMcpToolsetsList: S.Schema<LiteLLMObjectPermissionBaseMcpToolsetsList>;
|
|
255
|
-
export type LiteLLMObjectPermissionBaseModelsList = Array<string>;
|
|
256
|
-
export declare const LiteLLMObjectPermissionBaseModelsList: S.Schema<LiteLLMObjectPermissionBaseModelsList>;
|
|
257
|
-
export type LiteLLMObjectPermissionBaseSearchToolsList = Array<string>;
|
|
258
|
-
export declare const LiteLLMObjectPermissionBaseSearchToolsList: S.Schema<LiteLLMObjectPermissionBaseSearchToolsList>;
|
|
259
|
-
export type LiteLLMObjectPermissionBaseVectorStoresList = Array<string>;
|
|
260
|
-
export declare const LiteLLMObjectPermissionBaseVectorStoresList: S.Schema<LiteLLMObjectPermissionBaseVectorStoresList>;
|
|
261
|
-
export interface LiteLLMObjectPermissionBase {
|
|
262
|
-
agent_access_groups?: LiteLLMObjectPermissionBaseAgentAccessGroupsList | null;
|
|
263
|
-
agents?: LiteLLMObjectPermissionBaseAgentsList | null;
|
|
264
|
-
blocked_tools?: LiteLLMObjectPermissionBaseBlockedToolsList | null;
|
|
265
|
-
mcp_access_groups?: LiteLLMObjectPermissionBaseMcpAccessGroupsList | null;
|
|
266
|
-
mcp_servers?: LiteLLMObjectPermissionBaseMcpServersList | null;
|
|
267
|
-
mcp_tool_permissions?: LiteLLMObjectPermissionBaseMcpToolPermissionsMap | null;
|
|
268
|
-
mcp_tool_search_enabled?: boolean | null;
|
|
269
|
-
mcp_toolsets?: LiteLLMObjectPermissionBaseMcpToolsetsList | null;
|
|
270
|
-
models?: LiteLLMObjectPermissionBaseModelsList | null;
|
|
271
|
-
search_tools?: LiteLLMObjectPermissionBaseSearchToolsList | null;
|
|
272
|
-
vector_stores?: LiteLLMObjectPermissionBaseVectorStoresList | null;
|
|
273
|
-
}
|
|
274
|
-
export declare const LiteLLMObjectPermissionBase: S.Schema<LiteLLMObjectPermissionBase>;
|
|
275
279
|
export type GenerateKeyFnKeyGeneratePostRequestPermissionsMap = {
|
|
276
280
|
[key: string]: unknown | undefined;
|
|
277
281
|
};
|
|
@@ -313,8 +317,10 @@ export interface RetryPolicy {
|
|
|
313
317
|
AuthenticationErrorRetries?: number | null;
|
|
314
318
|
BadRequestErrorRetries?: number | null;
|
|
315
319
|
ContentPolicyViolationErrorRetries?: number | null;
|
|
320
|
+
DefaultRetries?: number | null;
|
|
316
321
|
InternalServerErrorRetries?: number | null;
|
|
317
322
|
RateLimitErrorRetries?: number | null;
|
|
323
|
+
ServiceUnavailableErrorRetries?: number | null;
|
|
318
324
|
TimeoutErrorRetries?: number | null;
|
|
319
325
|
}
|
|
320
326
|
export declare const RetryPolicy: S.Schema<RetryPolicy>;
|
|
@@ -322,6 +328,10 @@ export type UpdateRouterConfigModelGroupRetryPolicyMap = {
|
|
|
322
328
|
[key: string]: RetryPolicy | undefined;
|
|
323
329
|
};
|
|
324
330
|
export declare const UpdateRouterConfigModelGroupRetryPolicyMap: S.Schema<UpdateRouterConfigModelGroupRetryPolicyMap>;
|
|
331
|
+
export type UpdateRouterConfigOptionalPreCallChecksItem = "prompt_caching" | "router_budget_limiting" | "responses_api_deployment_check" | "deployment_affinity" | "session_affinity" | "forward_client_headers_by_model_group" | "enforce_model_rate_limits" | "encrypted_content_affinity";
|
|
332
|
+
export declare const UpdateRouterConfigOptionalPreCallChecksItem: any;
|
|
333
|
+
export type UpdateRouterConfigOptionalPreCallChecksList = Array<UpdateRouterConfigOptionalPreCallChecksItem | (string & {})>;
|
|
334
|
+
export declare const UpdateRouterConfigOptionalPreCallChecksList: S.Schema<UpdateRouterConfigOptionalPreCallChecksList>;
|
|
325
335
|
export type RoutingGroupModelsList = Array<string>;
|
|
326
336
|
export declare const RoutingGroupModelsList: S.Schema<RoutingGroupModelsList>;
|
|
327
337
|
export type RoutingGroupRoutingStrategyArgsMap = {
|
|
@@ -354,6 +364,7 @@ export interface UpdateRouterConfig {
|
|
|
354
364
|
model_group_alias?: UpdateRouterConfigModelGroupAliasMap | null;
|
|
355
365
|
model_group_retry_policy?: UpdateRouterConfigModelGroupRetryPolicyMap | null;
|
|
356
366
|
num_retries?: number | null;
|
|
367
|
+
optional_pre_call_checks?: UpdateRouterConfigOptionalPreCallChecksList | null;
|
|
357
368
|
retry_after?: number | null;
|
|
358
369
|
retry_policy?: RetryPolicy | null;
|
|
359
370
|
routing_groups?: UpdateRouterConfigRoutingGroupsList | null;
|
|
@@ -361,6 +372,7 @@ export interface UpdateRouterConfig {
|
|
|
361
372
|
routing_strategy_args?: UpdateRouterConfigRoutingStrategyArgsMap | null;
|
|
362
373
|
tag_routing_prefix?: string | null;
|
|
363
374
|
timeout?: number | null;
|
|
375
|
+
weights?: unknown | null;
|
|
364
376
|
}
|
|
365
377
|
export declare const UpdateRouterConfig: S.Schema<UpdateRouterConfig>;
|
|
366
378
|
export type GenerateKeyFnKeyGeneratePostRequestRpmLimitType = "guaranteed_throughput" | "best_effort_throughput" | "dynamic";
|
|
@@ -394,6 +406,7 @@ export interface GenerateKeyFnKeyGeneratePostRequest {
|
|
|
394
406
|
disable_global_guardrails?: boolean | null;
|
|
395
407
|
duration?: string | null;
|
|
396
408
|
enable_prompt_caching?: boolean | null;
|
|
409
|
+
end_user_budget_id?: string | null;
|
|
397
410
|
enforced_params?: GenerateKeyFnKeyGeneratePostRequestEnforcedParamsList | null;
|
|
398
411
|
guardrails?: GenerateKeyFnKeyGeneratePostRequestGuardrailsList | null;
|
|
399
412
|
key?: string | null;
|
|
@@ -426,6 +439,7 @@ export interface GenerateKeyFnKeyGeneratePostRequest {
|
|
|
426
439
|
tags?: GenerateKeyFnKeyGeneratePostRequestTagsList | null;
|
|
427
440
|
team_id?: string | null;
|
|
428
441
|
throttle_on_budget_exceeded?: boolean | null;
|
|
442
|
+
tpd_limit?: number | null;
|
|
429
443
|
tpm_limit?: number | null;
|
|
430
444
|
tpm_limit_type?: GenerateKeyFnKeyGeneratePostRequestTpmLimitType | (string & {}) | null;
|
|
431
445
|
user_id?: string | null;
|
|
@@ -526,6 +540,7 @@ export interface GenerateKeyResponse {
|
|
|
526
540
|
disable_global_guardrails?: boolean | null;
|
|
527
541
|
duration?: string | null;
|
|
528
542
|
enable_prompt_caching?: boolean | null;
|
|
543
|
+
end_user_budget_id?: string | null;
|
|
529
544
|
enforced_params?: GenerateKeyResponseEnforcedParamsList | null;
|
|
530
545
|
expires?: string | null;
|
|
531
546
|
guardrails?: GenerateKeyResponseGuardrailsList | null;
|
|
@@ -558,6 +573,7 @@ export interface GenerateKeyResponse {
|
|
|
558
573
|
throttle_on_budget_exceeded?: boolean | null;
|
|
559
574
|
token?: string | null;
|
|
560
575
|
token_id?: string | null;
|
|
576
|
+
tpd_limit?: number | null;
|
|
561
577
|
tpm_limit?: number | null;
|
|
562
578
|
tpm_limit_type?: GenerateKeyResponseTpmLimitType | null;
|
|
563
579
|
updated_at?: string | null;
|
|
@@ -660,6 +676,7 @@ export interface GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequest
|
|
|
660
676
|
disable_global_guardrails?: boolean | null;
|
|
661
677
|
duration?: string | null;
|
|
662
678
|
enable_prompt_caching?: boolean | null;
|
|
679
|
+
end_user_budget_id?: string | null;
|
|
663
680
|
enforced_params?: GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequestEnforcedParamsList | null;
|
|
664
681
|
guardrails?: GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequestGuardrailsList | null;
|
|
665
682
|
key?: string | null;
|
|
@@ -692,6 +709,7 @@ export interface GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequest
|
|
|
692
709
|
tags?: GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequestTagsList | null;
|
|
693
710
|
team_id?: string | null;
|
|
694
711
|
throttle_on_budget_exceeded?: boolean | null;
|
|
712
|
+
tpd_limit?: number | null;
|
|
695
713
|
tpm_limit?: number | null;
|
|
696
714
|
tpm_limit_type?: GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequestTpmLimitType | (string & {}) | null;
|
|
697
715
|
user_id?: string | null;
|
|
@@ -702,7 +720,7 @@ export interface GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRespons
|
|
|
702
720
|
}
|
|
703
721
|
export declare const GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostResponse: S.Schema<GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostResponse>;
|
|
704
722
|
export interface GetInfoKeyFnKeyInfoRequest {
|
|
705
|
-
/** Key
|
|
723
|
+
/** Key to look up. Pass the key's sha256 hash so the raw key stays out of URLs and access logs. Example key='d5345c0ecc68ae6295c69f91926b2bd379e25481a40c34b5884d157a9f65d8fa' */
|
|
706
724
|
key?: string;
|
|
707
725
|
}
|
|
708
726
|
export declare const GetInfoKeyFnKeyInfoRequest: S.Schema<GetInfoKeyFnKeyInfoRequest>;
|
|
@@ -744,8 +762,10 @@ export interface ListKeysKeyListGetRequest {
|
|
|
744
762
|
organization_id?: string;
|
|
745
763
|
/** Filter keys by key hash */
|
|
746
764
|
key_hash?: string;
|
|
747
|
-
/** Filter keys by key alias. Exact match by default; set substring_matching=true
|
|
765
|
+
/** Filter keys by key alias. Exact match by default; set substring_matching=true for case-insensitive substring matching. */
|
|
748
766
|
key_alias?: string;
|
|
767
|
+
/** Combined search: matches keys whose token (key hash) equals the value OR whose key_alias contains it (case-insensitive). */
|
|
768
|
+
search?: string;
|
|
749
769
|
/** Return full key object */
|
|
750
770
|
return_full_object?: boolean;
|
|
751
771
|
/** Include all keys for teams that user is an admin of. */
|
|
@@ -758,7 +778,7 @@ export interface ListKeysKeyListGetRequest {
|
|
|
758
778
|
sort_order?: string;
|
|
759
779
|
/** Expand related objects (e.g. 'user') */
|
|
760
780
|
expand?: ListKeysKeyListGetRequestExpandList;
|
|
761
|
-
/** Filter by status (
|
|
781
|
+
/** Filter by status: 'active' (not blocked, not expired), 'expired' (not blocked, past expiry), 'revoked' (blocked) or 'deleted' (archived keys). Omit to return live keys regardless of status. */
|
|
762
782
|
status?: string;
|
|
763
783
|
/** Filter keys by project ID */
|
|
764
784
|
project_id?: string;
|
|
@@ -766,7 +786,7 @@ export interface ListKeysKeyListGetRequest {
|
|
|
766
786
|
access_group_id?: string;
|
|
767
787
|
/** Filter keys by agent ID */
|
|
768
788
|
agent_id?: string;
|
|
769
|
-
/** If true (proxy admins only)
|
|
789
|
+
/** If true, match key_alias (any caller) and user_id (proxy admins only) as case-insensitive substrings instead of exact values. Defaults to false: /key/list matched these exactly before substring search was added, and an exact user_id filter must never return another user's keys. */
|
|
770
790
|
substring_matching?: boolean;
|
|
771
791
|
/** Filter keys by expiration. 'expired' returns keys whose expires is in the past; 'active' returns keys that never expire or expire in the future. Omit to return keys regardless of expiration. */
|
|
772
792
|
expires?: string;
|
|
@@ -826,6 +846,8 @@ export type LiteLLMObjectPermissionTableModelsList = Array<string>;
|
|
|
826
846
|
export declare const LiteLLMObjectPermissionTableModelsList: S.Schema<LiteLLMObjectPermissionTableModelsList>;
|
|
827
847
|
export type LiteLLMObjectPermissionTableSearchToolsList = Array<string>;
|
|
828
848
|
export declare const LiteLLMObjectPermissionTableSearchToolsList: S.Schema<LiteLLMObjectPermissionTableSearchToolsList>;
|
|
849
|
+
export type LiteLLMObjectPermissionTableSkillsList = Array<string>;
|
|
850
|
+
export declare const LiteLLMObjectPermissionTableSkillsList: S.Schema<LiteLLMObjectPermissionTableSkillsList>;
|
|
829
851
|
export type LiteLLMObjectPermissionTableVectorStoresList = Array<string>;
|
|
830
852
|
export declare const LiteLLMObjectPermissionTableVectorStoresList: S.Schema<LiteLLMObjectPermissionTableVectorStoresList>;
|
|
831
853
|
/** Represents a LiteLLM_ObjectPermissionTable record */
|
|
@@ -841,6 +863,7 @@ export interface LiteLLMObjectPermissionTable {
|
|
|
841
863
|
models?: LiteLLMObjectPermissionTableModelsList | null;
|
|
842
864
|
object_permission_id: string;
|
|
843
865
|
search_tools?: LiteLLMObjectPermissionTableSearchToolsList | null;
|
|
866
|
+
skills?: LiteLLMObjectPermissionTableSkillsList | null;
|
|
844
867
|
vector_stores?: LiteLLMObjectPermissionTableVectorStoresList | null;
|
|
845
868
|
}
|
|
846
869
|
export declare const LiteLLMObjectPermissionTable: S.Schema<LiteLLMObjectPermissionTable>;
|
|
@@ -906,6 +929,10 @@ export type UserAPIKeyAuthTeamModelAliasesMap = {
|
|
|
906
929
|
[key: string]: unknown | undefined;
|
|
907
930
|
};
|
|
908
931
|
export declare const UserAPIKeyAuthTeamModelAliasesMap: S.Schema<UserAPIKeyAuthTeamModelAliasesMap>;
|
|
932
|
+
export type UserAPIKeyAuthTeamModelMaxBudgetMap = {
|
|
933
|
+
[key: string]: unknown | undefined;
|
|
934
|
+
};
|
|
935
|
+
export declare const UserAPIKeyAuthTeamModelMaxBudgetMap: S.Schema<UserAPIKeyAuthTeamModelMaxBudgetMap>;
|
|
909
936
|
export type UserAPIKeyAuthTeamModelsList = Array<unknown>;
|
|
910
937
|
export declare const UserAPIKeyAuthTeamModelsList: S.Schema<UserAPIKeyAuthTeamModelsList>;
|
|
911
938
|
export type UserAPIKeyAuthTpmLimitPerModelMap = {
|
|
@@ -944,6 +971,7 @@ export interface UserAPIKeyAuth {
|
|
|
944
971
|
end_user_model_max_budget?: UserAPIKeyAuthEndUserModelMaxBudgetMap | null;
|
|
945
972
|
end_user_object_permission?: LiteLLMObjectPermissionTable | null;
|
|
946
973
|
end_user_rpm_limit?: number | null;
|
|
974
|
+
end_user_tpd_limit?: number | null;
|
|
947
975
|
end_user_tpm_limit?: number | null;
|
|
948
976
|
expires?: string | null;
|
|
949
977
|
is_session_token?: boolean;
|
|
@@ -995,14 +1023,18 @@ export interface UserAPIKeyAuth {
|
|
|
995
1023
|
team_member_tpm_limit?: number | null;
|
|
996
1024
|
team_metadata?: UserAPIKeyAuthTeamMetadataMap | null;
|
|
997
1025
|
team_model_aliases?: UserAPIKeyAuthTeamModelAliasesMap | null;
|
|
1026
|
+
team_model_max_budget?: UserAPIKeyAuthTeamModelMaxBudgetMap | null;
|
|
998
1027
|
team_models?: UserAPIKeyAuthTeamModelsList;
|
|
999
1028
|
team_object_permission?: LiteLLMObjectPermissionTable | null;
|
|
1000
1029
|
team_object_permission_id?: string | null;
|
|
1001
1030
|
team_rpm_limit?: number | null;
|
|
1002
1031
|
team_soft_budget?: number | null;
|
|
1003
1032
|
team_spend?: number | null;
|
|
1033
|
+
team_tpd_limit?: number | null;
|
|
1004
1034
|
team_tpm_limit?: number | null;
|
|
1005
1035
|
token?: string | null;
|
|
1036
|
+
total_spend?: number;
|
|
1037
|
+
tpd_limit?: number | null;
|
|
1006
1038
|
tpm_limit?: number | null;
|
|
1007
1039
|
tpm_limit_per_model?: UserAPIKeyAuthTpmLimitPerModelMap | null;
|
|
1008
1040
|
updated_at?: string | null;
|
|
@@ -1109,6 +1141,7 @@ export interface LiteLLMDeletedVerificationToken {
|
|
|
1109
1141
|
object_permission?: LiteLLMObjectPermissionTable | null;
|
|
1110
1142
|
object_permission_id?: string | null;
|
|
1111
1143
|
org_id?: string | null;
|
|
1144
|
+
organization_id?: string | null;
|
|
1112
1145
|
permissions?: LiteLLMDeletedVerificationTokenPermissionsMap;
|
|
1113
1146
|
project_id?: string | null;
|
|
1114
1147
|
rotation_count?: number | null;
|
|
@@ -1120,6 +1153,8 @@ export interface LiteLLMDeletedVerificationToken {
|
|
|
1120
1153
|
spend?: number;
|
|
1121
1154
|
team_id?: string | null;
|
|
1122
1155
|
token?: string | null;
|
|
1156
|
+
total_spend?: number;
|
|
1157
|
+
tpd_limit?: number | null;
|
|
1123
1158
|
tpm_limit?: number | null;
|
|
1124
1159
|
updated_at?: string | null;
|
|
1125
1160
|
updated_by?: string | null;
|
|
@@ -1237,6 +1272,8 @@ export interface LiteLLMVerificationToken {
|
|
|
1237
1272
|
spend?: number;
|
|
1238
1273
|
team_id?: string | null;
|
|
1239
1274
|
token?: string | null;
|
|
1275
|
+
total_spend?: number;
|
|
1276
|
+
tpd_limit?: number | null;
|
|
1240
1277
|
tpm_limit?: number | null;
|
|
1241
1278
|
updated_at?: string | null;
|
|
1242
1279
|
updated_by?: string | null;
|
|
@@ -1379,6 +1416,7 @@ export interface RegenerateKeyRequest {
|
|
|
1379
1416
|
disable_global_guardrails?: boolean | null;
|
|
1380
1417
|
duration?: string | null;
|
|
1381
1418
|
enable_prompt_caching?: boolean | null;
|
|
1419
|
+
end_user_budget_id?: string | null;
|
|
1382
1420
|
enforced_params?: RegenerateKeyRequestEnforcedParamsList | null;
|
|
1383
1421
|
grace_period?: string | null;
|
|
1384
1422
|
guardrails?: RegenerateKeyRequestGuardrailsList | null;
|
|
@@ -1414,6 +1452,7 @@ export interface RegenerateKeyRequest {
|
|
|
1414
1452
|
tags?: RegenerateKeyRequestTagsList | null;
|
|
1415
1453
|
team_id?: string | null;
|
|
1416
1454
|
throttle_on_budget_exceeded?: boolean | null;
|
|
1455
|
+
tpd_limit?: number | null;
|
|
1417
1456
|
tpm_limit?: number | null;
|
|
1418
1457
|
tpm_limit_type?: RegenerateKeyRequestTpmLimitType | (string & {}) | null;
|
|
1419
1458
|
user_id?: string | null;
|
|
@@ -1552,6 +1591,7 @@ export interface UpdateKeyFnKeyUpdatePostRequest {
|
|
|
1552
1591
|
disable_global_guardrails?: boolean | null;
|
|
1553
1592
|
duration?: string | null;
|
|
1554
1593
|
enable_prompt_caching?: boolean | null;
|
|
1594
|
+
end_user_budget_id?: string | null;
|
|
1555
1595
|
enforced_params?: UpdateKeyFnKeyUpdatePostRequestEnforcedParamsList | null;
|
|
1556
1596
|
guardrails?: UpdateKeyFnKeyUpdatePostRequestGuardrailsList | null;
|
|
1557
1597
|
key?: string | null;
|
|
@@ -1568,11 +1608,14 @@ export interface UpdateKeyFnKeyUpdatePostRequest {
|
|
|
1568
1608
|
organization_id?: string | null;
|
|
1569
1609
|
permissions?: UpdateKeyFnKeyUpdatePostRequestPermissionsMap | null;
|
|
1570
1610
|
policies?: UpdateKeyFnKeyUpdatePostRequestPoliciesList | null;
|
|
1611
|
+
/** Omit to retain the project, or send null to detach. Assigning a different project is not supported. */
|
|
1612
|
+
project_id?: string | null;
|
|
1571
1613
|
prompts?: UpdateKeyFnKeyUpdatePostRequestPromptsList | null;
|
|
1572
1614
|
rotation_interval?: string | null;
|
|
1573
1615
|
router_settings?: UpdateRouterConfig | null;
|
|
1574
1616
|
rpm_limit?: number | null;
|
|
1575
1617
|
rpm_limit_type?: UpdateKeyFnKeyUpdatePostRequestRpmLimitType | (string & {}) | null;
|
|
1618
|
+
soft_budget?: number | null;
|
|
1576
1619
|
spend?: number | null;
|
|
1577
1620
|
tag_rpm_limit?: UpdateKeyFnKeyUpdatePostRequestTagRpmLimitMap | null;
|
|
1578
1621
|
tags?: UpdateKeyFnKeyUpdatePostRequestTagsList | null;
|
|
@@ -1580,6 +1623,7 @@ export interface UpdateKeyFnKeyUpdatePostRequest {
|
|
|
1580
1623
|
temp_budget_expiry?: string | null;
|
|
1581
1624
|
temp_budget_increase?: number | null;
|
|
1582
1625
|
throttle_on_budget_exceeded?: boolean | null;
|
|
1626
|
+
tpd_limit?: number | null;
|
|
1583
1627
|
tpm_limit?: number | null;
|
|
1584
1628
|
tpm_limit_type?: UpdateKeyFnKeyUpdatePostRequestTpmLimitType | (string & {}) | null;
|
|
1585
1629
|
user_id?: string | null;
|
|
@@ -1590,7 +1634,7 @@ export interface UpdateKeyFnKeyUpdatePostResponse {
|
|
|
1590
1634
|
}
|
|
1591
1635
|
export declare const UpdateKeyFnKeyUpdatePostResponse: S.Schema<UpdateKeyFnKeyUpdatePostResponse>;
|
|
1592
1636
|
export type BulkUpdateKeysKeyBulkUpdatePostError = BadRequest | Forbidden | UnprocessableEntity | LitellmOpError;
|
|
1593
|
-
/** Bulk Update Keys Bulk update multiple keys at once. This endpoint allows updating multiple keys in a single request. Each key update is processed independently - if some updates fail, others will still succeed. Parameters: - keys: List[BulkUpdateKeyRequestItem] - List of key update requests, each containing: - key: str - The key identifier (token) to update - budget_id: Optional[str] - Budget ID associated with the key - max_budget: Optional[float] - Max budget for key - team_id: Optional[str] - Team ID associated with key - tags: Optional[List[str]] - Tags for organizing keys Returns: - total_requested: int - Total number of keys requested for update - successful_updates: List[SuccessfulKeyUpdate] - List of successfully updated keys with their updated info - failed_updates: List[FailedKeyUpdate] - List of failed updates with key_info and failed_reason Example request: ```bash curl --location 'http://0.0.0.0:4000/key/bulk_update' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "keys": [ { "key": "sk-1234", "max_budget": 100.0, "team_id": "team-123", "tags": ["production", "api"] }, { "key": "sk-5678", "budget_id": "budget-456", "tags": ["staging"] } ] }' ``` */
|
|
1637
|
+
/** Bulk Update Keys Bulk update multiple keys at once. This endpoint allows updating multiple keys in a single request. Each key update is processed independently - if some updates fail, others will still succeed. Parameters: - keys: List[BulkUpdateKeyRequestItem] - List of key update requests, each containing: - key: str - The key identifier (token) to update - budget_id: Optional[str] - Budget ID associated with the key - max_budget: Optional[float] - Max budget for key - team_id: Optional[str] - Team ID associated with key - tags: Optional[List[str]] - Tags for organizing keys - object_permission: Optional[LiteLLM_ObjectPermissionBase] - key-specific object permission, as on /key/update Only the fields an item carries are written: a field left out keeps its current value, and a field sent explicitly, null included, is applied exactly as /key/update applies it. Returns: - total_requested: int - Total number of keys requested for update - successful_updates: List[SuccessfulKeyUpdate] - List of successfully updated keys with their updated info - failed_updates: List[FailedKeyUpdate] - List of failed updates with key_info and failed_reason Example request: ```bash curl --location 'http://0.0.0.0:4000/key/bulk_update' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "keys": [ { "key": "sk-1234", "max_budget": 100.0, "team_id": "team-123", "tags": ["production", "api"] }, { "key": "sk-5678", "budget_id": "budget-456", "tags": ["staging"] } ] }' ``` */
|
|
1594
1638
|
export declare const bulkUpdateKeysKeyBulkUpdatePost: API.OperationMethod<BulkUpdateKeysKeyBulkUpdatePostRequest, BulkUpdateKeyResponse, BulkUpdateKeysKeyBulkUpdatePostError, LitellmOpContext>;
|
|
1595
1639
|
export type BulkUpdateTeamKeysTeamKeyBulkUpdatePostError = UnprocessableEntity | LitellmOpError;
|
|
1596
1640
|
/** Bulk Update Team Keys Apply one update payload to many keys inside a single team. Pass `team_id` plus either `key_ids` or `all_keys_in_team=True`. The `update_fields` payload is broadcast to every selected key. Per-key failures are returned in `failed_updates` rather than aborting the batch. Callable by proxy admins, or by team admins with `KEY_UPDATE` permission. */
|
|
@@ -1599,19 +1643,19 @@ export type DeleteKeyFnKeyDeletePostError = BadRequest | UnprocessableEntity | L
|
|
|
1599
1643
|
/** Delete Key Fn Delete a key from the key management system. Parameters:: - keys (List[str]): A list of keys or hashed keys to delete. Example {"keys": ["sk-QWrxEynunsNpV1zT48HIrw", "837e17519f44683334df5291321d97b8bf1098cd490e49e215f6fea935aa28be"]} - key_aliases (List[str]): A list of key aliases to delete. Can be passed instead of `keys`.Example {"key_aliases": ["alias1", "alias2"]} Returns: - deleted_keys (List[str]): A list of deleted keys. Example {"deleted_keys": ["sk-QWrxEynunsNpV1zT48HIrw", "837e17519f44683334df5291321d97b8bf1098cd490e49e215f6fea935aa28be"]} Example: ```bash curl --location 'http://0.0.0.0:4000/key/delete' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "keys": ["sk-QWrxEynunsNpV1zT48HIrw"] }' ``` Raises: HTTPException: If an error occurs during key deletion. */
|
|
1600
1644
|
export declare const deleteKeyFnKeyDeletePost: API.OperationMethod<DeleteKeyFnKeyDeletePostRequest, DeleteKeyFnKeyDeletePostResponse, DeleteKeyFnKeyDeletePostError, LitellmOpContext>;
|
|
1601
1645
|
export type GenerateKeyFnKeyGeneratePostError = BadRequest | Forbidden | UnprocessableEntity | LitellmOpError;
|
|
1602
|
-
/** Generate Key Fn Generate an API key based on the provided data. Docs: https://docs.litellm.ai/docs/proxy/virtual_keys Parameters: - duration: Optional[str] - Specify the length of time the token is valid for. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - key_alias: Optional[str] - User defined key alias - key: Optional[str] - User defined key value. Must start with 'sk-' and be at least 16 characters long. If not set, a 16-digit unique sk-key is created for you. - team_id: Optional[str] - The team id of the key - user_id: Optional[str] - The user id of the key - agent_id: Optional[str] - The agent id associated with the key. - organization_id: Optional[str] - The organization id of the key. If not set, and team_id is set, the organization id will be the same as the team id. If conflict, an error will be raised. - project_id: Optional[str] - The project id of the key. When set, models and max_budget are validated against the project's limits. - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - models: Optional[list] - Model_name's a user is allowed to call. (if empty, key is allowed to call all models) - aliases: Optional[dict] - Any alias mappings, on top of anything in the config.yaml model list. - https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---upgradedowngrade-models - config: Optional[dict] - any key-specific configs, overrides config in config.yaml - spend: Optional[int] - Amount spent by key. Default is 0. Will be updated by proxy whenever key is used. https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---tracking-spend - send_invite_email: Optional[bool] - Whether to send an invite email to the user_id, with the generate key - max_budget: Optional[float] - Specify max budget for a given key. - budget_duration: Optional[str] - Budget is reset at the end of specified duration. If not set, budget is never reset. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - max_parallel_requests: Optional[int] - Rate limit a user based on the number of parallel requests. Raises 429 error, if user's parallel requests > x. - metadata: Optional[dict] - Metadata for key, store information for key. Example metadata = {"team": "core-infra", "app": "app2", "email": "ishaan@berri.ai" } - guardrails: Optional[List[str]] - List of active guardrails for the key - policies: Optional[List[str]] - List of policy names to apply to the key. Policies define guardrails, conditions, and inheritance rules. - disable_global_guardrails: Optional[bool] - Whether to disable global guardrails for the key. - throttle_on_budget_exceeded: Optional[bool] - When the key exceeds its max_budget, throttle its tpm/rpm to the global budget_exceeded_throttle_percentage instead of blocking the key entirely. - enable_prompt_caching: Optional[bool] - Auto-inject prompt caching breakpoints (Anthropic cache_control markers) on requests made with this key. Anthropic and
|
|
1646
|
+
/** Generate Key Fn Generate an API key based on the provided data. Docs: https://docs.litellm.ai/docs/proxy/virtual_keys Parameters: - duration: Optional[str] - Specify the length of time the token is valid for. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - key_alias: Optional[str] - User defined key alias - key: Optional[str] - User defined key value. Must start with 'sk-' and be at least 16 characters long. If not set, a 16-digit unique sk-key is created for you. - team_id: Optional[str] - The team id of the key - user_id: Optional[str] - The user id of the key - agent_id: Optional[str] - The agent id associated with the key. - organization_id: Optional[str] - The organization id of the key. If not set, and team_id is set, the organization id will be the same as the team id. If conflict, an error will be raised. - project_id: Optional[str] - The project id of the key. When set, models and max_budget are validated against the project's limits. - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - end_user_budget_id: Optional[str] - Proxy admin only. Budget id applied to end users first seen through this key that carry no budget of their own. Takes precedence over `litellm_settings.max_end_user_budget_id`. - models: Optional[list] - Model_name's a user is allowed to call. (if empty, key is allowed to call all models) - aliases: Optional[dict] - Any alias mappings, on top of anything in the config.yaml model list. - https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---upgradedowngrade-models - config: Optional[dict] - any key-specific configs, overrides config in config.yaml - spend: Optional[int] - Amount spent by key. Default is 0. Will be updated by proxy whenever key is used. https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---tracking-spend - send_invite_email: Optional[bool] - Whether to send an invite email to the user_id, with the generate key - max_budget: Optional[float] - Specify max budget for a given key. - budget_duration: Optional[str] - Budget is reset at the end of specified duration. If not set, budget is never reset. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - max_parallel_requests: Optional[int] - Rate limit a user based on the number of parallel requests. Raises 429 error, if user's parallel requests > x. - metadata: Optional[dict] - Metadata for key, store information for key. Example metadata = {"team": "core-infra", "app": "app2", "email": "ishaan@berri.ai" } - guardrails: Optional[List[str]] - List of active guardrails for the key - policies: Optional[List[str]] - List of policy names to apply to the key. Policies define guardrails, conditions, and inheritance rules. - disable_global_guardrails: Optional[bool] - Whether to disable global guardrails for the key. - throttle_on_budget_exceeded: Optional[bool] - When the key exceeds its max_budget, throttle its tpm/rpm to the global budget_exceeded_throttle_percentage instead of blocking the key entirely. - enable_prompt_caching: Optional[bool] - Auto-inject prompt caching breakpoints (Anthropic cache_control markers) on requests made with this key. Supported Claude models on Anthropic, Bedrock, Vertex AI, and Azure AI only. - permissions: Optional[dict] - key-specific permissions. Currently just used for turning off pii masking (if connected). Example - {"pii": false} - model_max_budget: Optional[Dict[str, BudgetConfig]] - Model-specific budgets {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}}}. IF null or {} then no model specific budget. - budget_fallbacks: Optional[Dict[str, List[str]]] - Per-model fallback chain tried in order when that model's own `model_max_budget` is exceeded, e.g. {"gpt-4o": ["gpt-4o-mini"]}. - model_rpm_limit: Optional[dict] - key-specific model rpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific rpm limit. - model_tpm_limit: Optional[dict] - key-specific model tpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific tpm limit. - default_estimated_output_tokens: Optional[int] - Proxy admin only. Expected output tokens reserved for TPM limiting when a request omits max_tokens. Positive integer. Falls back to the team setting, then to the built-in estimate. - default_estimated_output_tokens_per_model: Optional[dict] - Proxy admin only. Per-model override of the above. Example - {"gpt-4": 4096, "gpt-3.5-turbo": 1024}. Takes precedence over the key-wide value. - mcp_rpm_limit: Optional[dict] - key-specific per-MCP-server rpm limit, keyed by MCP server name (alias if set, else the configured name). Example - {"github": 100, "slack": 200}. IF null or {} then no MCP-specific rpm limit. - tag_rpm_limit: Optional[dict] - key-specific per-request-tag rpm limit, keyed by request tag. Example - {"cell-1": 1000, "cell-2": 500}. Each tag gets an independent counter; requests whose tag is absent fall back to the key-level rpm limit. - tpm_limit_type: Optional[str] - Type of tpm limit. Options: "best_effort_throughput" (no error if we're overallocating tpm), "guaranteed_throughput" (raise an error if we're overallocating tpm), "dynamic" (dynamically exceed limit when no 429 errors). Defaults to "best_effort_throughput". - rpm_limit_type: Optional[str] - Type of rpm limit. Options: "best_effort_throughput" (no error if we're overallocating rpm), "guaranteed_throughput" (raise an error if we're overallocating rpm), "dynamic" (dynamically exceed limit when no 429 errors). Defaults to "best_effort_throughput". - allowed_cache_controls: Optional[list] - List of allowed cache control values. Example - ["no-cache", "no-store"]. See all values - https://docs.litellm.ai/docs/proxy/caching#turn-on--off-caching-per-request - blocked: Optional[bool] - Whether the key is blocked. - rpm_limit: Optional[int] - Specify rpm limit for a given key (Requests per minute) - tpm_limit: Optional[int] - Specify tpm limit for a given key (Tokens per minute) - tpd_limit: Optional[int] - Specify tpd limit for a given key (Tokens per day). Charged by batch submissions instead of tpm_limit/rpm_limit. - soft_budget: Optional[float] - Specify soft budget for a given key. Will trigger a slack alert when this soft budget is reached. - tags: Optional[List[str]] - Tags for [tracking spend](https://litellm.vercel.app/docs/proxy/enterprise#tracking-spend-for-custom-tags) and/or doing [tag-based routing](https://litellm.vercel.app/docs/proxy/tag_routing). - prompts: Optional[List[str]] - List of prompts that the key is allowed to use. - enforced_params: Optional[List[str]] - List of enforced params for the key (Enterprise only). [Docs](https://docs.litellm.ai/docs/proxy/enterprise#enforce-required-params-for-llm-requests) - prompts: Optional[List[str]] - List of prompts that the key is allowed to use. - allowed_routes: Optional[list] - List of allowed routes for the key. Store the actual route or store a wildcard pattern for a set of routes. Example - ["/chat/completions", "/embeddings", "/keys/*"] - allowed_passthrough_routes: Optional[list] - List of allowed pass through endpoints for the key. Store the actual endpoint or store a wildcard pattern for a set of endpoints. Example - ["/my-custom-endpoint"]. Use this instead of allowed_routes, if you just want to specify which pass through endpoints the key can access, without specifying the routes. If allowed_routes is specified, allowed_pass_through_endpoints is ignored. - object_permission: Optional[LiteLLM_ObjectPermissionBase] - key-specific object permission. Example - {"vector_stores": ["vector_store_1", "vector_store_2"], "agents": ["agent_1", "agent_2"], "agent_access_groups": ["dev_group"]}. IF null or {} then no object permission. - key_type: Optional[str] - Type of key that determines default allowed routes. Options: "llm_api" (can call LLM API routes), "management" (can call management routes), "read_only" (can only call info/read routes), "default" (uses default allowed routes). Defaults to "default". - prompts: Optional[List[str]] - List of allowed prompts for the key. If specified, the key will only be able to use these specific prompts. - auto_rotate: Optional[bool] - Whether this key should be automatically rotated (regenerated) - rotation_interval: Optional[str] - How often to auto-rotate this key (e.g., '30s', '30m', '30h', '30d'). Required if auto_rotate=True. - allowed_vector_store_indexes: Optional[List[dict]] - List of allowed vector store indexes for the key. Example - [{"index_name": "my-index", "index_permissions": ["write", "read"]}]. If specified, the key will only be able to use these specific vector store indexes. Create index, using `/v1/indexes` endpoint. - router_settings: Optional[UpdateRouterConfig] - key-specific router settings. Example - {"model_group_retry_policy": {"gpt-4": {"RateLimitErrorRetries": 5}}}. IF null or {} then no router settings. - access_group_ids: Optional[List[str]] - List of access group IDs to associate with the key. Access groups define which models a key can access. Example - ["access_group_1", "access_group_2"]. - budget_limits: Optional[list] - List of concurrent budget windows for the key. Each window specifies a budget_limit, time_period, and optional budget_duration. Example - [{"budget_limit": 10.0, "time_period": "1d"}, {"budget_limit": 50.0, "time_period": "7d"}]. Examples: 1. Allow users to turn on/off pii masking ```bash curl --location 'http://0.0.0.0:4000/key/generate' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "permissions": {"allow_pii_controls": true} }' ``` Returns: - key: (str) The generated api key - expires: (datetime) Datetime object for when key expires. - user_id: (str) Unique user id - used for tracking spend across multiple keys for same user id. */
|
|
1603
1647
|
export declare const generateKeyFnKeyGeneratePost: API.OperationMethod<GenerateKeyFnKeyGeneratePostRequest, GenerateKeyResponse, GenerateKeyFnKeyGeneratePostError, LitellmOpContext>;
|
|
1604
1648
|
export type GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostError = UnprocessableEntity | LitellmOpError;
|
|
1605
|
-
/** Generate Service Account Key Fn Generate a Service Account API key based on the provided data. This key does not belong to any user. It belongs to the team. Why use a service account key? - Prevent key from being deleted when user is deleted. - Apply team limits, not team member limits to key. Docs: https://docs.litellm.ai/docs/proxy/virtual_keys Parameters: - duration: Optional[str] - Specify the length of time the token is valid for. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - key_alias: Optional[str] - User defined key alias - key: Optional[str] - User defined key value. Must start with 'sk-' and be at least 16 characters long. If not set, a 16-digit unique sk-key is created for you. - team_id: Optional[str] - The team id of the key - user_id: Optional[str] - [NON-FUNCTIONAL] THIS WILL BE IGNORED. The user id of the key - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - models: Optional[list] - Model_name's a user is allowed to call. (if empty, key is allowed to call all models) - aliases: Optional[dict] - Any alias mappings, on top of anything in the config.yaml model list. - https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---upgradedowngrade-models - config: Optional[dict] - any key-specific configs, overrides config in config.yaml - spend: Optional[int] - Amount spent by key. Default is 0. Will be updated by proxy whenever key is used. https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---tracking-spend - send_invite_email: Optional[bool] - Whether to send an invite email to the user_id, with the generate key - max_budget: Optional[float] - Specify max budget for a given key. - budget_duration: Optional[str] - Budget is reset at the end of specified duration. If not set, budget is never reset. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - max_parallel_requests: Optional[int] - Rate limit a user based on the number of parallel requests. Raises 429 error, if user's parallel requests > x. - metadata: Optional[dict] - Metadata for key, store information for key. Example metadata = {"team": "core-infra", "app": "app2", "email": "ishaan@berri.ai" } - guardrails: Optional[List[str]] - List of active guardrails for the key - permissions: Optional[dict] - key-specific permissions. Currently just used for turning off pii masking (if connected). Example - {"pii": false} - model_max_budget: Optional[Dict[str, BudgetConfig]] - Model-specific budgets {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}}}. IF null or {} then no model specific budget. - budget_fallbacks: Optional[Dict[str, List[str]]] - Per-model fallback chain tried in order when that model's own `model_max_budget` is exceeded, e.g. {"gpt-4o": ["gpt-4o-mini"]}. - model_rpm_limit: Optional[dict] - key-specific model rpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific rpm limit. - model_tpm_limit: Optional[dict] - key-specific model tpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific tpm limit. - default_estimated_output_tokens: Optional[int] - Proxy admin only. Expected output tokens reserved for TPM limiting when a request omits max_tokens. Positive integer. Falls back to the team setting, then to the built-in estimate. - default_estimated_output_tokens_per_model: Optional[dict] - Proxy admin only. Per-model override of the above. Example - {"gpt-4": 4096, "gpt-3.5-turbo": 1024}. Takes precedence over the key-wide value. - mcp_rpm_limit: Optional[dict] - key-specific per-MCP-server rpm limit, keyed by MCP server name (alias if set, else the configured name). Example - {"github": 100, "slack": 200}. IF null or {} then no MCP-specific rpm limit. - tpm_limit_type: Optional[str] - TPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - rpm_limit_type: Optional[str] - RPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - allowed_cache_controls: Optional[list] - List of allowed cache control values. Example - ["no-cache", "no-store"]. See all values - https://docs.litellm.ai/docs/proxy/caching#turn-on--off-caching-per-request - blocked: Optional[bool] - Whether the key is blocked. - rpm_limit: Optional[int] - Specify rpm limit for a given key (Requests per minute) - tpm_limit: Optional[int] - Specify tpm limit for a given key (Tokens per minute) - soft_budget: Optional[float] - Specify soft budget for a given key. Will trigger a slack alert when this soft budget is reached. - tags: Optional[List[str]] - Tags for [tracking spend](https://litellm.vercel.app/docs/proxy/enterprise#tracking-spend-for-custom-tags) and/or doing [tag-based routing](https://litellm.vercel.app/docs/proxy/tag_routing). - enforced_params: Optional[List[str]] - List of enforced params for the key (Enterprise only). [Docs](https://docs.litellm.ai/docs/proxy/enterprise#enforce-required-params-for-llm-requests) - allowed_routes: Optional[list] - List of allowed routes for the key. Store the actual route or store a wildcard pattern for a set of routes. Example - ["/chat/completions", "/embeddings", "/keys/*"] - object_permission: Optional[LiteLLM_ObjectPermissionBase] - key-specific object permission. Example - {"vector_stores": ["vector_store_1", "vector_store_2"], "agents": ["agent_1", "agent_2"], "agent_access_groups": ["dev_group"]}. IF null or {} then no object permission. Examples: - allowed_vector_store_indexes: Optional[List[dict]] - List of allowed vector store indexes for the key. Example - [{"index_name": "my-index", "index_permissions": ["write", "read"]}]. If specified, the key will only be able to use these specific vector store indexes. Create index, using `/v1/indexes` endpoint. 1. Allow users to turn on/off pii masking ```bash curl --location 'http://0.0.0.0:4000/key/generate' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "permissions": {"allow_pii_controls": true} }' ``` Returns: - key: (str) The generated api key - expires: (datetime) Datetime object for when key expires. - user_id: (str) Unique user id - used for tracking spend across multiple keys for same user id. */
|
|
1649
|
+
/** Generate Service Account Key Fn Generate a Service Account API key based on the provided data. This key does not belong to any user. It belongs to the team. Why use a service account key? - Prevent key from being deleted when user is deleted. - Apply team limits, not team member limits to key. Docs: https://docs.litellm.ai/docs/proxy/virtual_keys Parameters: - duration: Optional[str] - Specify the length of time the token is valid for. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - key_alias: Optional[str] - User defined key alias - key: Optional[str] - User defined key value. Must start with 'sk-' and be at least 16 characters long. If not set, a 16-digit unique sk-key is created for you. - team_id: Optional[str] - The team id of the key - user_id: Optional[str] - [NON-FUNCTIONAL] THIS WILL BE IGNORED. The user id of the key - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - end_user_budget_id: Optional[str] - Proxy admin only. Budget id applied to end users first seen through this key that carry no budget of their own. Omit to keep the current value, pass an empty string to clear it. - models: Optional[list] - Model_name's a user is allowed to call. (if empty, key is allowed to call all models) - aliases: Optional[dict] - Any alias mappings, on top of anything in the config.yaml model list. - https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---upgradedowngrade-models - config: Optional[dict] - any key-specific configs, overrides config in config.yaml - spend: Optional[int] - Amount spent by key. Default is 0. Will be updated by proxy whenever key is used. https://docs.litellm.ai/docs/proxy/virtual_keys#managing-auth---tracking-spend - send_invite_email: Optional[bool] - Whether to send an invite email to the user_id, with the generate key - max_budget: Optional[float] - Specify max budget for a given key. - budget_duration: Optional[str] - Budget is reset at the end of specified duration. If not set, budget is never reset. You can set duration as seconds ("30s"), minutes ("30m"), hours ("30h"), days ("30d"). - max_parallel_requests: Optional[int] - Rate limit a user based on the number of parallel requests. Raises 429 error, if user's parallel requests > x. - metadata: Optional[dict] - Metadata for key, store information for key. Example metadata = {"team": "core-infra", "app": "app2", "email": "ishaan@berri.ai" } - guardrails: Optional[List[str]] - List of active guardrails for the key - permissions: Optional[dict] - key-specific permissions. Currently just used for turning off pii masking (if connected). Example - {"pii": false} - model_max_budget: Optional[Dict[str, BudgetConfig]] - Model-specific budgets {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}}}. IF null or {} then no model specific budget. - budget_fallbacks: Optional[Dict[str, List[str]]] - Per-model fallback chain tried in order when that model's own `model_max_budget` is exceeded, e.g. {"gpt-4o": ["gpt-4o-mini"]}. - model_rpm_limit: Optional[dict] - key-specific model rpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific rpm limit. - model_tpm_limit: Optional[dict] - key-specific model tpm limit. Example - {"text-davinci-002": 1000, "gpt-3.5-turbo": 1000}. IF null or {} then no model specific tpm limit. - default_estimated_output_tokens: Optional[int] - Proxy admin only. Expected output tokens reserved for TPM limiting when a request omits max_tokens. Positive integer. Falls back to the team setting, then to the built-in estimate. - default_estimated_output_tokens_per_model: Optional[dict] - Proxy admin only. Per-model override of the above. Example - {"gpt-4": 4096, "gpt-3.5-turbo": 1024}. Takes precedence over the key-wide value. - mcp_rpm_limit: Optional[dict] - key-specific per-MCP-server rpm limit, keyed by MCP server name (alias if set, else the configured name). Example - {"github": 100, "slack": 200}. IF null or {} then no MCP-specific rpm limit. - tpm_limit_type: Optional[str] - TPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - rpm_limit_type: Optional[str] - RPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - allowed_cache_controls: Optional[list] - List of allowed cache control values. Example - ["no-cache", "no-store"]. See all values - https://docs.litellm.ai/docs/proxy/caching#turn-on--off-caching-per-request - blocked: Optional[bool] - Whether the key is blocked. - rpm_limit: Optional[int] - Specify rpm limit for a given key (Requests per minute) - tpm_limit: Optional[int] - Specify tpm limit for a given key (Tokens per minute) - tpd_limit: Optional[int] - Specify tpd limit for a given key (Tokens per day). Charged by batch submissions instead of tpm_limit/rpm_limit. - soft_budget: Optional[float] - Specify soft budget for a given key. Will trigger a slack alert when this soft budget is reached. - tags: Optional[List[str]] - Tags for [tracking spend](https://litellm.vercel.app/docs/proxy/enterprise#tracking-spend-for-custom-tags) and/or doing [tag-based routing](https://litellm.vercel.app/docs/proxy/tag_routing). - enforced_params: Optional[List[str]] - List of enforced params for the key (Enterprise only). [Docs](https://docs.litellm.ai/docs/proxy/enterprise#enforce-required-params-for-llm-requests) - allowed_routes: Optional[list] - List of allowed routes for the key. Store the actual route or store a wildcard pattern for a set of routes. Example - ["/chat/completions", "/embeddings", "/keys/*"] - object_permission: Optional[LiteLLM_ObjectPermissionBase] - key-specific object permission. Example - {"vector_stores": ["vector_store_1", "vector_store_2"], "agents": ["agent_1", "agent_2"], "agent_access_groups": ["dev_group"]}. IF null or {} then no object permission. Examples: - allowed_vector_store_indexes: Optional[List[dict]] - List of allowed vector store indexes for the key. Example - [{"index_name": "my-index", "index_permissions": ["write", "read"]}]. If specified, the key will only be able to use these specific vector store indexes. Create index, using `/v1/indexes` endpoint. 1. Allow users to turn on/off pii masking ```bash curl --location 'http://0.0.0.0:4000/key/generate' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "permissions": {"allow_pii_controls": true} }' ``` Returns: - key: (str) The generated api key - expires: (datetime) Datetime object for when key expires. - user_id: (str) Unique user id - used for tracking spend across multiple keys for same user id. */
|
|
1606
1650
|
export declare const generateServiceAccountKeyFnKeyServiceAccountGeneratePost: API.OperationMethod<GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostRequest, GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostResponse, GenerateServiceAccountKeyFnKeyServiceAccountGeneratePostError, LitellmOpContext>;
|
|
1607
1651
|
export type GetInfoKeyFnKeyInfoError = UnprocessableEntity | LitellmOpError;
|
|
1608
|
-
/** Info Key Fn Retrieve information about a key. Parameters: - key: str | None (query parameter) - The key to look up. Accepts the plaintext key or its hash. Defaults to the key in the Authorization header. Returns: - key: str - The key that was looked up, echoed back as it was passed in - info: dict - The key's row, minus the hashed token - key_alias: str | None - User-friendly key alias - spend: float - Amount spent by the key. When budget_duration is set this covers only the current budget window, not the key's lifetime - max_budget: float | None - Max budget for the key, enforced against spend - budget_duration: str | None - Budget reset period ("30d", "1h", etc.) - budget_reset_at: datetime | None - When the current budget window ends and spend is next reset to 0, not when it was last reset. Reset times snap to standard boundaries in the configured timezone (30d and 1mo land on the 1st of the month, 7d on Monday, 1h on the hour), so subtracting budget_duration from it does not give the window's start - model_max_budget: dict - Per-model budgets, e.g. {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}} - model_max_budget_usage: dict | None - Current-window spend per model, present only when the key has per-model budgets - models: list - Model_name's the key is allowed to call - tpm_limit / rpm_limit: int | None - Tokens and requests per minute limits - metadata: dict - Metadata for the key, e.g. {"team": "core-infra"} - blocked: bool | None - Whether the key is blocked - expires: datetime | None - When the key stops authenticating requests - last_active: datetime | None - When the key was last used - object_permission: dict | None - Resolved vector store / MCP permissions when the key has an object_permission_id Example Curl: ``` curl -X GET "http://0.0.0.0:4000/key/info?key=
|
|
1652
|
+
/** Info Key Fn Retrieve information about a key. Parameters: - key: str | None (query parameter) - The key to look up. Accepts the plaintext key or its hash; prefer the hash, since a query parameter is recorded verbatim by any HTTP access log in front of the proxy. Defaults to the key in the Authorization header. Returns: - key: str - The key that was looked up, echoed back as it was passed in - info: dict - The key's row, minus the hashed token. Deleted keys are served from the LiteLLM_DeletedVerificationToken archive and carry deleted_at / deleted_by - status: "active" | "expired" | "revoked" | "deleted" - Derived from blocked, expires and whether the row came from the archive - key_alias: str | None - User-friendly key alias - spend: float - Amount spent by the key. When budget_duration is set this covers only the current budget window, not the key's lifetime - max_budget: float | None - Max budget for the key, enforced against spend - budget_duration: str | None - Budget reset period ("30d", "1h", etc.) - budget_reset_at: datetime | None - When the current budget window ends and spend is next reset to 0, not when it was last reset. Reset times snap to standard boundaries in the configured timezone (30d and 1mo land on the 1st of the month, 7d on Monday, 1h on the hour), so subtracting budget_duration from it does not give the window's start - model_max_budget: dict - Per-model budgets, e.g. {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}} - model_max_budget_usage: dict | None - Current-window spend per model, present only when the key has per-model budgets - budget_limits: list | None - Concurrent budget windows, exactly as stored - budget_limits_usage: dict | None - Current-window spend per budget window, e.g. {"1h": {"current_spend": 0.0009}}, present only when the key has budget windows (read from the same cross-pod spend counter the budget enforcement uses) - models: list - Model_name's the key is allowed to call - tpm_limit / rpm_limit: int | None - Tokens and requests per minute limits - metadata: dict - Metadata for the key, e.g. {"team": "core-infra"} - blocked: bool | None - Whether the key is blocked - expires: datetime | None - When the key stops authenticating requests - last_active: datetime | None - When the key was last used - object_permission: dict | None - Resolved vector store / MCP permissions when the key has an object_permission_id Example Curl: ``` curl -X GET "http://0.0.0.0:4000/key/info?key=d5345c0ecc68ae6295c69f91926b2bd379e25481a40c34b5884d157a9f65d8fa" -H "Authorization: Bearer sk-1234" ``` Example Curl - if no key is passed, it will use the Key Passed in Authorization Header ``` curl -X GET "http://0.0.0.0:4000/key/info" -H "Authorization: Bearer sk-test-example-key-123" ``` */
|
|
1609
1653
|
export declare const getInfoKeyFnKeyInfo: API.OperationMethod<GetInfoKeyFnKeyInfoRequest, GetInfoKeyFnKeyInfoResponse, GetInfoKeyFnKeyInfoError, LitellmOpContext>;
|
|
1610
1654
|
export type GetKeyAliasesKeyAliasError = UnprocessableEntity | LitellmOpError;
|
|
1611
1655
|
/** Key Aliases Lists key aliases with pagination and optional search. Non-admin users only see aliases for keys they own or keys belonging to their teams. Returns: { "aliases": List[str], "total_count": int, "current_page": int, "total_pages": int, "size": int, } */
|
|
1612
1656
|
export declare const getKeyAliasesKeyAlias: API.OperationMethod<GetKeyAliasesKeyAliasRequest, GetKeyAliasesKeyAliasResponse, GetKeyAliasesKeyAliasError, LitellmOpContext>;
|
|
1613
1657
|
export type ListKeysKeyListGetError = BadRequest | UnprocessableEntity | LitellmOpError;
|
|
1614
|
-
/** List Keys List all keys for a given user / team / organization. Parameters: expand: Optional[List[str]] - Expand related objects (e.g. 'user' to include user information) status: Optional[str] - Filter by status
|
|
1658
|
+
/** List Keys List all keys for a given user / team / organization. Parameters: expand: Optional[List[str]] - Expand related objects (e.g. 'user' to include user information) status: Optional[str] - Filter by status: "active", "expired", "revoked" (blocked) or "deleted". "deleted" reads the LiteLLM_DeletedVerificationToken archive; the other values partition the live key table, so every live key matches exactly one of them. Returns: { "keys": List[str] or List[UserAPIKeyAuth], "total_count": int, "current_page": int, "total_pages": int, } When expand includes "user", each key object will include a "user" field with the associated user object. Note: When expand=user is specified, full key objects are returned regardless of the return_full_object parameter. */
|
|
1615
1659
|
export declare const listKeysKeyListGet: API.OperationMethod<ListKeysKeyListGetRequest, KeyListResponseObject, ListKeysKeyListGetError, LitellmOpContext>;
|
|
1616
1660
|
export type PostBlockKeyKeyBlockError = UnprocessableEntity | LitellmOpError;
|
|
1617
1661
|
/** Block Key Block an Virtual key from making any requests. Parameters: - key: str - The key to block. Can be either the unhashed key (sk-...) or the hashed key value Example: ```bash curl --location 'http://0.0.0.0:4000/key/block' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "key": "sk-Fn8Ej39NxjAXrvpUGKghGw" }' ``` Note: This is an admin-only endpoint. Only proxy admins, team admins, or org admins can block keys. */
|
|
@@ -1635,6 +1679,6 @@ export type UnblockKeyKeyUnblockPostError = UnprocessableEntity | LitellmOpError
|
|
|
1635
1679
|
/** Unblock Key Unblock a Virtual key to allow it to make requests again. Parameters: - key: str - The key to unblock. Can be either the unhashed key (sk-...) or the hashed key value Example: ```bash curl --location 'http://0.0.0.0:4000/key/unblock' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "key": "sk-Fn8Ej39NxjAXrvpUGKghGw" }' ``` Note: This is an admin-only endpoint. Only proxy admins, team admins, or org admins can unblock keys. */
|
|
1636
1680
|
export declare const unblockKeyKeyUnblockPost: API.OperationMethod<UnblockKeyKeyUnblockPostRequest, UnblockKeyKeyUnblockPostResponse, UnblockKeyKeyUnblockPostError, LitellmOpContext>;
|
|
1637
1681
|
export type UpdateKeyFnKeyUpdatePostError = BadRequest | Forbidden | NotFound | UnprocessableEntity | LitellmOpError;
|
|
1638
|
-
/** Update Key Fn Update an existing API key's parameters. Parameters: - key: Optional[str] - The key to update. Either key or key_alias must be provided. - key_alias: Optional[str] - User-friendly key alias. If key is omitted, also identifies the key to update (must match exactly one key, same as /key/delete's key_aliases) - user_id: Optional[str] - User ID associated with key - team_id: Optional[str] - Team ID associated with key - agent_id: Optional[str] - The agent id associated with the key. - organization_id: Optional[str] - The organization id of the key. - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - models: Optional[list] - Model_name's a user is allowed to call - tags: Optional[List[str]] - Tags for organizing keys (Enterprise only) - prompts: Optional[List[str]] - List of prompts that the key is allowed to use. - enforced_params: Optional[List[str]] - List of enforced params for the key (Enterprise only). [Docs](https://docs.litellm.ai/docs/proxy/enterprise#enforce-required-params-for-llm-requests) - spend: Optional[float] - Amount spent by key - max_budget: Optional[float] - Max budget for key - model_max_budget: Optional[Dict[str, BudgetConfig]] - Model-specific budgets {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}} - budget_fallbacks: Optional[Dict[str, List[str]]] - Per-model fallback chain tried in order when that model's own `model_max_budget` is exceeded, e.g. {"gpt-4o": ["gpt-4o-mini"]}. - budget_duration: Optional[str] - Budget reset period ("30d", "1h", etc.) - soft_budget: Optional[float] -
|
|
1682
|
+
/** Update Key Fn Update an existing API key's parameters. The body is a merge patch: a field left out keeps its stored value, and on the key's own columns an explicit null clears it. The metadata-backed fields below are the exception, merging into the stored metadata instead: passing one as null leaves it unchanged, while `metadata` itself replaces the stored metadata wholesale. Parameters: - key: Optional[str] - The key to update. Either key or key_alias must be provided. - key_alias: Optional[str] - User-friendly key alias. If key is omitted, also identifies the key to update (must match exactly one key, same as /key/delete's key_aliases) - user_id: Optional[str] - User ID associated with key - team_id: Optional[str] - Team ID associated with key - agent_id: Optional[str] - The agent id associated with the key. - project_id: Optional[str] - Omit to retain the project, or send null to detach. A different project ID is rejected. - organization_id: Optional[str] - The organization id of the key. - budget_id: Optional[str] - The budget id associated with the key. Created by calling `/budget/new`. - end_user_budget_id: Optional[str] - Proxy admin only. Budget id applied to end users first seen through this key that carry no budget of their own. Omit to keep the current value, pass an empty string to clear it. - models: Optional[list] - Model_name's a user is allowed to call - tags: Optional[List[str]] - Tags for organizing keys (Enterprise only) - prompts: Optional[List[str]] - List of prompts that the key is allowed to use. - enforced_params: Optional[List[str]] - List of enforced params for the key (Enterprise only). [Docs](https://docs.litellm.ai/docs/proxy/enterprise#enforce-required-params-for-llm-requests) - spend: Optional[float] - Amount spent by key - max_budget: Optional[float] - Max budget for key - model_max_budget: Optional[Dict[str, BudgetConfig]] - Model-specific budgets {"gpt-4": {"budget_limit": 0.0005, "time_period": "30d"}} - budget_fallbacks: Optional[Dict[str, List[str]]] - Per-model fallback chain tried in order when that model's own `model_max_budget` is exceeded, e.g. {"gpt-4o": ["gpt-4o-mini"]}. - budget_duration: Optional[str] - Budget reset period ("30d", "1h", etc.) - soft_budget: Optional[float] - Soft budget limit (warning vs. hard stop). Will trigger a slack alert when this soft budget is reached. Set to null to remove the soft budget. - max_parallel_requests: Optional[int] - Rate limit for parallel requests - metadata: Optional[dict] - Metadata for key. Example {"team": "core-infra", "app": "app2"} - tpm_limit: Optional[int] - Tokens per minute limit - rpm_limit: Optional[int] - Requests per minute limit - tpd_limit: Optional[int] - Tokens per day limit, charged by batch submissions instead of tpm_limit/rpm_limit - model_rpm_limit: Optional[dict] - Model-specific RPM limits {"gpt-4": 100, "claude-v1": 200} - mcp_rpm_limit: Optional[dict] - Per-MCP-server RPM limits, keyed by MCP server name {"github": 100, "slack": 200} - tag_rpm_limit: Optional[dict] - Per-request-tag RPM limits, keyed by request tag {"cell-1": 1000, "cell-2": 500}. Each tag gets an independent counter; absent tags fall back to the key-level rpm limit. - model_tpm_limit: Optional[dict] - Model-specific TPM limits {"gpt-4": 100000, "claude-v1": 200000} - default_estimated_output_tokens: Optional[int] - Proxy admin only. Expected output tokens reserved for TPM limiting when a request omits max_tokens. Positive integer. - default_estimated_output_tokens_per_model: Optional[dict] - Proxy admin only. Per-model override of the above {"gpt-4": 4096, "gpt-3.5-turbo": 1024} - tpm_limit_type: Optional[str] - TPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - rpm_limit_type: Optional[str] - RPM rate limit type - "best_effort_throughput", "guaranteed_throughput", or "dynamic" - allowed_cache_controls: Optional[list] - List of allowed cache control values - duration: Optional[str] - Key validity duration ("30d", "1h", etc.), null to never expire, or "-1" to never expire (deprecated, use null) - permissions: Optional[dict] - Key-specific permissions - send_invite_email: Optional[bool] - Send invite email to user_id - guardrails: Optional[List[str]] - List of active guardrails for the key - policies: Optional[List[str]] - List of policy names to apply to the key. Policies define guardrails, conditions, and inheritance rules. - disable_global_guardrails: Optional[bool] - Whether to disable global guardrails for the key. - throttle_on_budget_exceeded: Optional[bool] - When the key exceeds its max_budget, throttle its tpm/rpm to the global budget_exceeded_throttle_percentage instead of blocking the key entirely. - enable_prompt_caching: Optional[bool] - Auto-inject prompt caching breakpoints (Anthropic cache_control markers) on requests made with this key. Supported Claude models on Anthropic, Bedrock, Vertex AI, and Azure AI only. - prompts: Optional[List[str]] - List of prompts that the key is allowed to use. - blocked: Optional[bool] - Whether the key is blocked - aliases: Optional[dict] - Model aliases for the key - [Docs](https://litellm.vercel.app/docs/proxy/virtual_keys#model-aliases) - config: Optional[dict] - [DEPRECATED PARAM] Key-specific config. - temp_budget_increase: Optional[float] - Temporary budget increase for the key (Enterprise only). - temp_budget_expiry: Optional[str] - Expiry time for the temporary budget increase (Enterprise only). - allowed_routes: Optional[list] - List of allowed routes for the key. Store the actual route or store a wildcard pattern for a set of routes. Example - ["/chat/completions", "/embeddings", "/keys/*"] - allowed_passthrough_routes: Optional[list] - List of allowed pass through routes for the key. Store the actual route or store a wildcard pattern for a set of routes. Example - ["/my-custom-endpoint"]. Use this instead of allowed_routes, if you just want to specify which pass through routes the key can access, without specifying the routes. If allowed_routes is specified, allowed_passthrough_routes is ignored. - prompts: Optional[List[str]] - List of allowed prompts for the key. If specified, the key will only be able to use these specific prompts. - object_permission: Optional[LiteLLM_ObjectPermissionBase] - key-specific object permission. Example - {"vector_stores": ["vector_store_1", "vector_store_2"], "agents": ["agent_1", "agent_2"], "agent_access_groups": ["dev_group"]}. IF null or {} then no object permission. - auto_rotate: Optional[bool] - Whether this key should be automatically rotated - rotation_interval: Optional[str] - How often to rotate this key (e.g., '30d', '90d'). Required if auto_rotate=True - allowed_vector_store_indexes: Optional[List[dict]] - List of allowed vector store indexes for the key. Example - [{"index_name": "my-index", "index_permissions": ["write", "read"]}]. If specified, the key will only be able to use these specific vector store indexes. Create index, using `/v1/indexes` endpoint. - router_settings: Optional[UpdateRouterConfig] - key-specific router settings. Example - {"model_group_retry_policy": {"gpt-4": {"RateLimitErrorRetries": 5}}}. IF null or {} then no router settings. - access_group_ids: Optional[List[str]] - List of access group IDs to associate with the key. Access groups define which models a key can access. Example - ["access_group_1", "access_group_2"]. - budget_limits: Optional[list] - List of concurrent budget windows for the key. Each window specifies a budget_limit, time_period, and optional budget_duration. Example - [{"budget_limit": 10.0, "time_period": "1d"}, {"budget_limit": 50.0, "time_period": "7d"}]. Example: ```bash curl --location 'http://0.0.0.0:4000/key/update' --header 'Authorization: Bearer sk-1234' --header 'Content-Type: application/json' --data '{ "key": "sk-1234", "key_alias": "my-key", "user_id": "user-1234", "team_id": "team-1234", "max_budget": 100, "metadata": {"any_key": "any-val"}, }' ``` */
|
|
1639
1683
|
export declare const updateKeyFnKeyUpdatePost: API.OperationMethod<UpdateKeyFnKeyUpdatePostRequest, UpdateKeyFnKeyUpdatePostResponse, UpdateKeyFnKeyUpdatePostError, LitellmOpContext>;
|
|
1640
1684
|
//# sourceMappingURL=key_management.d.ts.map
|