@homeflare/distilled-litellm 0.2.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (563) hide show
  1. package/LICENSE +201 -0
  2. package/README.md +77 -0
  3. package/dist/credentials.d.ts +27 -0
  4. package/dist/credentials.d.ts.map +1 -0
  5. package/dist/credentials.js +67 -0
  6. package/dist/credentials.js.map +1 -0
  7. package/dist/errors.d.ts +92 -0
  8. package/dist/errors.d.ts.map +1 -0
  9. package/dist/errors.js +72 -0
  10. package/dist/errors.js.map +1 -0
  11. package/dist/index.d.ts +30 -0
  12. package/dist/index.d.ts.map +1 -0
  13. package/dist/index.js +30 -0
  14. package/dist/index.js.map +1 -0
  15. package/dist/protocol.d.ts +18 -0
  16. package/dist/protocol.d.ts.map +1 -0
  17. package/dist/protocol.js +135 -0
  18. package/dist/protocol.js.map +1 -0
  19. package/dist/retry.d.ts +50 -0
  20. package/dist/retry.d.ts.map +1 -0
  21. package/dist/retry.js +43 -0
  22. package/dist/retry.js.map +1 -0
  23. package/dist/services/a2a.d.ts +70 -0
  24. package/dist/services/a2a.d.ts.map +1 -0
  25. package/dist/services/a2a.js +128 -0
  26. package/dist/services/a2a.js.map +1 -0
  27. package/dist/services/a2a_registration.d.ts +43 -0
  28. package/dist/services/a2a_registration.d.ts.map +1 -0
  29. package/dist/services/a2a_registration.js +40 -0
  30. package/dist/services/a2a_registration.js.map +1 -0
  31. package/dist/services/access_groups.d.ts +192 -0
  32. package/dist/services/access_groups.d.ts.map +1 -0
  33. package/dist/services/access_groups.js +283 -0
  34. package/dist/services/access_groups.js.map +1 -0
  35. package/dist/services/adaptive_router.d.ts +15 -0
  36. package/dist/services/adaptive_router.d.ts.map +1 -0
  37. package/dist/services/adaptive_router.js +25 -0
  38. package/dist/services/adaptive_router.js.map +1 -0
  39. package/dist/services/agents.d.ts +537 -0
  40. package/dist/services/agents.d.ts.map +1 -0
  41. package/dist/services/agents.js +481 -0
  42. package/dist/services/agents.js.map +1 -0
  43. package/dist/services/alerting.d.ts +15 -0
  44. package/dist/services/alerting.d.ts.map +1 -0
  45. package/dist/services/alerting.js +25 -0
  46. package/dist/services/alerting.js.map +1 -0
  47. package/dist/services/anthropic_pass_through.d.ts +70 -0
  48. package/dist/services/anthropic_pass_through.d.ts.map +1 -0
  49. package/dist/services/anthropic_pass_through.js +114 -0
  50. package/dist/services/anthropic_pass_through.js.map +1 -0
  51. package/dist/services/anthropic_passthrough.d.ts +35 -0
  52. package/dist/services/anthropic_passthrough.d.ts.map +1 -0
  53. package/dist/services/anthropic_passthrough.js +59 -0
  54. package/dist/services/anthropic_passthrough.js.map +1 -0
  55. package/dist/services/anthropic_skills.d.ts +74 -0
  56. package/dist/services/anthropic_skills.d.ts.map +1 -0
  57. package/dist/services/anthropic_skills.js +94 -0
  58. package/dist/services/anthropic_skills.js.map +1 -0
  59. package/dist/services/assembly_ai_eu_pass_through.d.ts +70 -0
  60. package/dist/services/assembly_ai_eu_pass_through.d.ts.map +1 -0
  61. package/dist/services/assembly_ai_eu_pass_through.js +114 -0
  62. package/dist/services/assembly_ai_eu_pass_through.js.map +1 -0
  63. package/dist/services/assembly_ai_pass_through.d.ts +70 -0
  64. package/dist/services/assembly_ai_pass_through.d.ts.map +1 -0
  65. package/dist/services/assembly_ai_pass_through.js +114 -0
  66. package/dist/services/assembly_ai_pass_through.js.map +1 -0
  67. package/dist/services/assistants.d.ts +185 -0
  68. package/dist/services/assistants.d.ts.map +1 -0
  69. package/dist/services/assistants.js +331 -0
  70. package/dist/services/assistants.js.map +1 -0
  71. package/dist/services/audio.d.ts +57 -0
  72. package/dist/services/audio.d.ts.map +1 -0
  73. package/dist/services/audio.js +96 -0
  74. package/dist/services/audio.js.map +1 -0
  75. package/dist/services/audit_logging.d.ts +98 -0
  76. package/dist/services/audit_logging.d.ts.map +1 -0
  77. package/dist/services/audit_logging.js +83 -0
  78. package/dist/services/audit_logging.js.map +1 -0
  79. package/dist/services/auto_router.d.ts +286 -0
  80. package/dist/services/auto_router.d.ts.map +1 -0
  81. package/dist/services/auto_router.js +248 -0
  82. package/dist/services/auto_router.js.map +1 -0
  83. package/dist/services/aws_comprehend_medical_pass_through.d.ts +36 -0
  84. package/dist/services/aws_comprehend_medical_pass_through.d.ts.map +1 -0
  85. package/dist/services/aws_comprehend_medical_pass_through.js +56 -0
  86. package/dist/services/aws_comprehend_medical_pass_through.js.map +1 -0
  87. package/dist/services/azure_ai_pass_through.d.ts +70 -0
  88. package/dist/services/azure_ai_pass_through.d.ts.map +1 -0
  89. package/dist/services/azure_ai_pass_through.js +112 -0
  90. package/dist/services/azure_ai_pass_through.js.map +1 -0
  91. package/dist/services/azure_pass_through.d.ts +70 -0
  92. package/dist/services/azure_pass_through.d.ts.map +1 -0
  93. package/dist/services/azure_pass_through.js +107 -0
  94. package/dist/services/azure_pass_through.js.map +1 -0
  95. package/dist/services/batch.d.ts +165 -0
  96. package/dist/services/batch.d.ts.map +1 -0
  97. package/dist/services/batch.js +265 -0
  98. package/dist/services/batch.js.map +1 -0
  99. package/dist/services/bedrock_pass_through.d.ts +70 -0
  100. package/dist/services/bedrock_pass_through.d.ts.map +1 -0
  101. package/dist/services/bedrock_pass_through.js +114 -0
  102. package/dist/services/bedrock_pass_through.js.map +1 -0
  103. package/dist/services/beta_agents.d.ts +199 -0
  104. package/dist/services/beta_agents.d.ts.map +1 -0
  105. package/dist/services/beta_agents.js +151 -0
  106. package/dist/services/beta_agents.js.map +1 -0
  107. package/dist/services/beta_mcp.d.ts +37 -0
  108. package/dist/services/beta_mcp.d.ts.map +1 -0
  109. package/dist/services/beta_mcp.js +40 -0
  110. package/dist/services/beta_mcp.js.map +1 -0
  111. package/dist/services/budget_management.d.ts +185 -0
  112. package/dist/services/budget_management.d.ts.map +1 -0
  113. package/dist/services/budget_management.js +196 -0
  114. package/dist/services/budget_management.js.map +1 -0
  115. package/dist/services/budget_spend_tracking.d.ts +967 -0
  116. package/dist/services/budget_spend_tracking.d.ts.map +1 -0
  117. package/dist/services/budget_spend_tracking.js +1103 -0
  118. package/dist/services/budget_spend_tracking.js.map +1 -0
  119. package/dist/services/cache_settings.d.ts +96 -0
  120. package/dist/services/cache_settings.d.ts.map +1 -0
  121. package/dist/services/cache_settings.js +95 -0
  122. package/dist/services/cache_settings.js.map +1 -0
  123. package/dist/services/caching.d.ts +54 -0
  124. package/dist/services/caching.d.ts.map +1 -0
  125. package/dist/services/caching.js +78 -0
  126. package/dist/services/caching.js.map +1 -0
  127. package/dist/services/chat_completions.d.ts +656 -0
  128. package/dist/services/chat_completions.d.ts.map +1 -0
  129. package/dist/services/chat_completions.js +590 -0
  130. package/dist/services/chat_completions.js.map +1 -0
  131. package/dist/services/claude_code_marketplace.d.ts +212 -0
  132. package/dist/services/claude_code_marketplace.d.ts.map +1 -0
  133. package/dist/services/claude_code_marketplace.js +248 -0
  134. package/dist/services/claude_code_marketplace.js.map +1 -0
  135. package/dist/services/cloudzero.d.ts +117 -0
  136. package/dist/services/cloudzero.d.ts.map +1 -0
  137. package/dist/services/cloudzero.js +130 -0
  138. package/dist/services/cloudzero.js.map +1 -0
  139. package/dist/services/cohere_pass_through.d.ts +70 -0
  140. package/dist/services/cohere_pass_through.d.ts.map +1 -0
  141. package/dist/services/cohere_pass_through.js +112 -0
  142. package/dist/services/cohere_pass_through.js.map +1 -0
  143. package/dist/services/completions.d.ts +59 -0
  144. package/dist/services/completions.d.ts.map +1 -0
  145. package/dist/services/completions.js +98 -0
  146. package/dist/services/completions.js.map +1 -0
  147. package/dist/services/compliance.d.ts +66 -0
  148. package/dist/services/compliance.d.ts.map +1 -0
  149. package/dist/services/compliance.js +74 -0
  150. package/dist/services/compliance.js.map +1 -0
  151. package/dist/services/config_overrides.d.ts +157 -0
  152. package/dist/services/config_overrides.d.ts.map +1 -0
  153. package/dist/services/config_overrides.js +207 -0
  154. package/dist/services/config_overrides.js.map +1 -0
  155. package/dist/services/config_yaml.d.ts +647 -0
  156. package/dist/services/config_yaml.d.ts.map +1 -0
  157. package/dist/services/config_yaml.js +504 -0
  158. package/dist/services/config_yaml.js.map +1 -0
  159. package/dist/services/containers.d.ts +215 -0
  160. package/dist/services/containers.d.ts.map +1 -0
  161. package/dist/services/containers.js +416 -0
  162. package/dist/services/containers.js.map +1 -0
  163. package/dist/services/coordination_redis_settings.d.ts +93 -0
  164. package/dist/services/coordination_redis_settings.d.ts.map +1 -0
  165. package/dist/services/coordination_redis_settings.js +104 -0
  166. package/dist/services/coordination_redis_settings.js.map +1 -0
  167. package/dist/services/cost_tracking.d.ts +140 -0
  168. package/dist/services/cost_tracking.d.ts.map +1 -0
  169. package/dist/services/cost_tracking.js +181 -0
  170. package/dist/services/cost_tracking.js.map +1 -0
  171. package/dist/services/credential_management.d.ts +133 -0
  172. package/dist/services/credential_management.d.ts.map +1 -0
  173. package/dist/services/credential_management.js +198 -0
  174. package/dist/services/credential_management.js.map +1 -0
  175. package/dist/services/cursor_pass_through.d.ts +70 -0
  176. package/dist/services/cursor_pass_through.d.ts.map +1 -0
  177. package/dist/services/cursor_pass_through.js +112 -0
  178. package/dist/services/cursor_pass_through.js.map +1 -0
  179. package/dist/services/customer_management.d.ts +545 -0
  180. package/dist/services/customer_management.d.ts.map +1 -0
  181. package/dist/services/customer_management.js +567 -0
  182. package/dist/services/customer_management.js.map +1 -0
  183. package/dist/services/email_management.d.ts +57 -0
  184. package/dist/services/email_management.d.ts.map +1 -0
  185. package/dist/services/email_management.js +79 -0
  186. package/dist/services/email_management.js.map +1 -0
  187. package/dist/services/embeddings.d.ts +155 -0
  188. package/dist/services/embeddings.d.ts.map +1 -0
  189. package/dist/services/embeddings.js +168 -0
  190. package/dist/services/embeddings.js.map +1 -0
  191. package/dist/services/evals.d.ts +245 -0
  192. package/dist/services/evals.d.ts.map +1 -0
  193. package/dist/services/evals.js +292 -0
  194. package/dist/services/evals.js.map +1 -0
  195. package/dist/services/experimental.d.ts +166 -0
  196. package/dist/services/experimental.d.ts.map +1 -0
  197. package/dist/services/experimental.js +265 -0
  198. package/dist/services/experimental.js.map +1 -0
  199. package/dist/services/fallback_management.d.ts +91 -0
  200. package/dist/services/fallback_management.d.ts.map +1 -0
  201. package/dist/services/fallback_management.js +86 -0
  202. package/dist/services/fallback_management.js.map +1 -0
  203. package/dist/services/files.d.ts +219 -0
  204. package/dist/services/files.d.ts.map +1 -0
  205. package/dist/services/files.js +358 -0
  206. package/dist/services/files.js.map +1 -0
  207. package/dist/services/fine_tuning.d.ts +155 -0
  208. package/dist/services/fine_tuning.d.ts.map +1 -0
  209. package/dist/services/fine_tuning.js +232 -0
  210. package/dist/services/fine_tuning.js.map +1 -0
  211. package/dist/services/gemini_agents.d.ts +68 -0
  212. package/dist/services/gemini_agents.d.ts.map +1 -0
  213. package/dist/services/gemini_agents.js +110 -0
  214. package/dist/services/gemini_agents.js.map +1 -0
  215. package/dist/services/google_ai_studio_pass_through.d.ts +70 -0
  216. package/dist/services/google_ai_studio_pass_through.d.ts.map +1 -0
  217. package/dist/services/google_ai_studio_pass_through.js +112 -0
  218. package/dist/services/google_ai_studio_pass_through.js.map +1 -0
  219. package/dist/services/google_genai_endpoints.d.ts +174 -0
  220. package/dist/services/google_genai_endpoints.d.ts.map +1 -0
  221. package/dist/services/google_genai_endpoints.js +340 -0
  222. package/dist/services/google_genai_endpoints.js.map +1 -0
  223. package/dist/services/guardrails.d.ts +1221 -0
  224. package/dist/services/guardrails.d.ts.map +1 -0
  225. package/dist/services/guardrails.js +1087 -0
  226. package/dist/services/guardrails.js.map +1 -0
  227. package/dist/services/health.d.ts +212 -0
  228. package/dist/services/health.d.ts.map +1 -0
  229. package/dist/services/health.js +309 -0
  230. package/dist/services/health.js.map +1 -0
  231. package/dist/services/images.d.ts +117 -0
  232. package/dist/services/images.d.ts.map +1 -0
  233. package/dist/services/images.js +184 -0
  234. package/dist/services/images.js.map +1 -0
  235. package/dist/services/index.d.ts +106 -0
  236. package/dist/services/index.d.ts.map +1 -0
  237. package/dist/services/index.js +107 -0
  238. package/dist/services/index.js.map +1 -0
  239. package/dist/services/internal_user_management.d.ts +1062 -0
  240. package/dist/services/internal_user_management.d.ts.map +1 -0
  241. package/dist/services/internal_user_management.js +850 -0
  242. package/dist/services/internal_user_management.js.map +1 -0
  243. package/dist/services/invite_links.d.ts +56 -0
  244. package/dist/services/invite_links.d.ts.map +1 -0
  245. package/dist/services/invite_links.js +82 -0
  246. package/dist/services/invite_links.js.map +1 -0
  247. package/dist/services/jwt_mappings.d.ts +79 -0
  248. package/dist/services/jwt_mappings.d.ts.map +1 -0
  249. package/dist/services/jwt_mappings.js +116 -0
  250. package/dist/services/jwt_mappings.js.map +1 -0
  251. package/dist/services/key_management.d.ts +1640 -0
  252. package/dist/services/key_management.d.ts.map +1 -0
  253. package/dist/services/key_management.js +1344 -0
  254. package/dist/services/key_management.js.map +1 -0
  255. package/dist/services/langfuse_passthrough.d.ts +70 -0
  256. package/dist/services/langfuse_passthrough.d.ts.map +1 -0
  257. package/dist/services/langfuse_passthrough.js +114 -0
  258. package/dist/services/langfuse_passthrough.js.map +1 -0
  259. package/dist/services/llm_utils.d.ts +101 -0
  260. package/dist/services/llm_utils.d.ts.map +1 -0
  261. package/dist/services/llm_utils.js +110 -0
  262. package/dist/services/llm_utils.js.map +1 -0
  263. package/dist/services/logging_callbacks.d.ts +33 -0
  264. package/dist/services/logging_callbacks.d.ts.map +1 -0
  265. package/dist/services/logging_callbacks.js +46 -0
  266. package/dist/services/logging_callbacks.js.map +1 -0
  267. package/dist/services/mcp_byok_oauth.d.ts +112 -0
  268. package/dist/services/mcp_byok_oauth.d.ts.map +1 -0
  269. package/dist/services/mcp_byok_oauth.js +226 -0
  270. package/dist/services/mcp_byok_oauth.js.map +1 -0
  271. package/dist/services/mcp_discoverable.d.ts +197 -0
  272. package/dist/services/mcp_discoverable.d.ts.map +1 -0
  273. package/dist/services/mcp_discoverable.js +304 -0
  274. package/dist/services/mcp_discoverable.js.map +1 -0
  275. package/dist/services/mcp_management.d.ts +956 -0
  276. package/dist/services/mcp_management.d.ts.map +1 -0
  277. package/dist/services/mcp_management.js +1139 -0
  278. package/dist/services/mcp_management.js.map +1 -0
  279. package/dist/services/mcp_rest.d.ts +281 -0
  280. package/dist/services/mcp_rest.d.ts.map +1 -0
  281. package/dist/services/mcp_rest.js +269 -0
  282. package/dist/services/mcp_rest.js.map +1 -0
  283. package/dist/services/memory_management.d.ts +93 -0
  284. package/dist/services/memory_management.d.ts.map +1 -0
  285. package/dist/services/memory_management.js +117 -0
  286. package/dist/services/memory_management.js.map +1 -0
  287. package/dist/services/milvus_pass_through.d.ts +70 -0
  288. package/dist/services/milvus_pass_through.d.ts.map +1 -0
  289. package/dist/services/milvus_pass_through.js +112 -0
  290. package/dist/services/milvus_pass_through.js.map +1 -0
  291. package/dist/services/misc.d.ts +676 -0
  292. package/dist/services/misc.d.ts.map +1 -0
  293. package/dist/services/misc.js +954 -0
  294. package/dist/services/misc.js.map +1 -0
  295. package/dist/services/mistral_pass_through.d.ts +70 -0
  296. package/dist/services/mistral_pass_through.d.ts.map +1 -0
  297. package/dist/services/mistral_pass_through.js +114 -0
  298. package/dist/services/mistral_pass_through.js.map +1 -0
  299. package/dist/services/model_management.d.ts +1649 -0
  300. package/dist/services/model_management.d.ts.map +1 -0
  301. package/dist/services/model_management.js +1805 -0
  302. package/dist/services/model_management.js.map +1 -0
  303. package/dist/services/moderations.d.ts +25 -0
  304. package/dist/services/moderations.d.ts.map +1 -0
  305. package/dist/services/moderations.js +39 -0
  306. package/dist/services/moderations.js.map +1 -0
  307. package/dist/services/ocr.d.ts +25 -0
  308. package/dist/services/ocr.d.ts.map +1 -0
  309. package/dist/services/ocr.js +39 -0
  310. package/dist/services/ocr.js.map +1 -0
  311. package/dist/services/open_ai_pass_through.d.ts +125 -0
  312. package/dist/services/open_ai_pass_through.d.ts.map +1 -0
  313. package/dist/services/open_ai_pass_through.js +232 -0
  314. package/dist/services/open_ai_pass_through.js.map +1 -0
  315. package/dist/services/organization_management.d.ts +664 -0
  316. package/dist/services/organization_management.d.ts.map +1 -0
  317. package/dist/services/organization_management.js +644 -0
  318. package/dist/services/organization_management.js.map +1 -0
  319. package/dist/services/plugins.d.ts +106 -0
  320. package/dist/services/plugins.d.ts.map +1 -0
  321. package/dist/services/plugins.js +180 -0
  322. package/dist/services/plugins.js.map +1 -0
  323. package/dist/services/policies.d.ts +618 -0
  324. package/dist/services/policies.d.ts.map +1 -0
  325. package/dist/services/policies.js +616 -0
  326. package/dist/services/policies.js.map +1 -0
  327. package/dist/services/policy_engine.d.ts +483 -0
  328. package/dist/services/policy_engine.d.ts.map +1 -0
  329. package/dist/services/policy_engine.js +501 -0
  330. package/dist/services/policy_engine.js.map +1 -0
  331. package/dist/services/project_management.d.ts +356 -0
  332. package/dist/services/project_management.d.ts.map +1 -0
  333. package/dist/services/project_management.js +318 -0
  334. package/dist/services/project_management.js.map +1 -0
  335. package/dist/services/prompts.d.ts +193 -0
  336. package/dist/services/prompts.d.ts.map +1 -0
  337. package/dist/services/prompts.js +260 -0
  338. package/dist/services/prompts.js.map +1 -0
  339. package/dist/services/public.d.ts +300 -0
  340. package/dist/services/public.d.ts.map +1 -0
  341. package/dist/services/public.js +362 -0
  342. package/dist/services/public.js.map +1 -0
  343. package/dist/services/rag.d.ts +45 -0
  344. package/dist/services/rag.d.ts.map +1 -0
  345. package/dist/services/rag.js +71 -0
  346. package/dist/services/rag.js.map +1 -0
  347. package/dist/services/realtime.d.ts +91 -0
  348. package/dist/services/realtime.d.ts.map +1 -0
  349. package/dist/services/realtime.js +164 -0
  350. package/dist/services/realtime.js.map +1 -0
  351. package/dist/services/rerank.d.ts +35 -0
  352. package/dist/services/rerank.d.ts.map +1 -0
  353. package/dist/services/rerank.js +55 -0
  354. package/dist/services/rerank.js.map +1 -0
  355. package/dist/services/responses.d.ts +237 -0
  356. package/dist/services/responses.d.ts.map +1 -0
  357. package/dist/services/responses.js +445 -0
  358. package/dist/services/responses.js.map +1 -0
  359. package/dist/services/router_settings.d.ts +67 -0
  360. package/dist/services/router_settings.d.ts.map +1 -0
  361. package/dist/services/router_settings.js +63 -0
  362. package/dist/services/router_settings.js.map +1 -0
  363. package/dist/services/rust_control_plane.d.ts +52 -0
  364. package/dist/services/rust_control_plane.d.ts.map +1 -0
  365. package/dist/services/rust_control_plane.js +54 -0
  366. package/dist/services/rust_control_plane.js.map +1 -0
  367. package/dist/services/scim.d.ts +412 -0
  368. package/dist/services/scim.d.ts.map +1 -0
  369. package/dist/services/scim.js +519 -0
  370. package/dist/services/scim.js.map +1 -0
  371. package/dist/services/search.d.ts +79 -0
  372. package/dist/services/search.d.ts.map +1 -0
  373. package/dist/services/search.js +122 -0
  374. package/dist/services/search.js.map +1 -0
  375. package/dist/services/search_tools.d.ts +140 -0
  376. package/dist/services/search_tools.d.ts.map +1 -0
  377. package/dist/services/search_tools.js +201 -0
  378. package/dist/services/search_tools.js.map +1 -0
  379. package/dist/services/settings.d.ts +53 -0
  380. package/dist/services/settings.d.ts.map +1 -0
  381. package/dist/services/settings.js +67 -0
  382. package/dist/services/settings.js.map +1 -0
  383. package/dist/services/sso_settings.d.ts +239 -0
  384. package/dist/services/sso_settings.d.ts.map +1 -0
  385. package/dist/services/sso_settings.js +214 -0
  386. package/dist/services/sso_settings.js.map +1 -0
  387. package/dist/services/tag_management.d.ts +390 -0
  388. package/dist/services/tag_management.d.ts.map +1 -0
  389. package/dist/services/tag_management.js +407 -0
  390. package/dist/services/tag_management.js.map +1 -0
  391. package/dist/services/team_management.d.ts +1502 -0
  392. package/dist/services/team_management.d.ts.map +1 -0
  393. package/dist/services/team_management.js +1429 -0
  394. package/dist/services/team_management.js.map +1 -0
  395. package/dist/services/tools.d.ts +227 -0
  396. package/dist/services/tools.d.ts.map +1 -0
  397. package/dist/services/tools.js +266 -0
  398. package/dist/services/tools.js.map +1 -0
  399. package/dist/services/ui_settings.d.ts +90 -0
  400. package/dist/services/ui_settings.d.ts.map +1 -0
  401. package/dist/services/ui_settings.js +96 -0
  402. package/dist/services/ui_settings.js.map +1 -0
  403. package/dist/services/ui_theme_settings.d.ts +62 -0
  404. package/dist/services/ui_theme_settings.d.ts.map +1 -0
  405. package/dist/services/ui_theme_settings.js +79 -0
  406. package/dist/services/ui_theme_settings.js.map +1 -0
  407. package/dist/services/usage_ai.d.ts +39 -0
  408. package/dist/services/usage_ai.d.ts.map +1 -0
  409. package/dist/services/usage_ai.js +40 -0
  410. package/dist/services/usage_ai.js.map +1 -0
  411. package/dist/services/vantage.d.ts +108 -0
  412. package/dist/services/vantage.d.ts.map +1 -0
  413. package/dist/services/vantage.js +124 -0
  414. package/dist/services/vantage.js.map +1 -0
  415. package/dist/services/vector_store_management.d.ts +161 -0
  416. package/dist/services/vector_store_management.d.ts.map +1 -0
  417. package/dist/services/vector_store_management.js +191 -0
  418. package/dist/services/vector_store_management.js.map +1 -0
  419. package/dist/services/vector_stores.d.ts +342 -0
  420. package/dist/services/vector_stores.d.ts.map +1 -0
  421. package/dist/services/vector_stores.js +639 -0
  422. package/dist/services/vector_stores.js.map +1 -0
  423. package/dist/services/vertex_ai_pass_through.d.ts +180 -0
  424. package/dist/services/vertex_ai_pass_through.d.ts.map +1 -0
  425. package/dist/services/vertex_ai_pass_through.js +334 -0
  426. package/dist/services/vertex_ai_pass_through.js.map +1 -0
  427. package/dist/services/videos.d.ts +207 -0
  428. package/dist/services/videos.d.ts.map +1 -0
  429. package/dist/services/videos.js +373 -0
  430. package/dist/services/videos.js.map +1 -0
  431. package/dist/services/vllm_pass_through.d.ts +70 -0
  432. package/dist/services/vllm_pass_through.d.ts.map +1 -0
  433. package/dist/services/vllm_pass_through.js +104 -0
  434. package/dist/services/vllm_pass_through.js.map +1 -0
  435. package/dist/services/watsonx_pass_through.d.ts +70 -0
  436. package/dist/services/watsonx_pass_through.d.ts.map +1 -0
  437. package/dist/services/watsonx_pass_through.js +114 -0
  438. package/dist/services/watsonx_pass_through.js.map +1 -0
  439. package/dist/services/web_socket.d.ts +77 -0
  440. package/dist/services/web_socket.d.ts.map +1 -0
  441. package/dist/services/web_socket.js +135 -0
  442. package/dist/services/web_socket.js.map +1 -0
  443. package/dist/services/workflow_management.d.ts +140 -0
  444. package/dist/services/workflow_management.d.ts.map +1 -0
  445. package/dist/services/workflow_management.js +220 -0
  446. package/dist/services/workflow_management.js.map +1 -0
  447. package/dist/traits.d.ts +14 -0
  448. package/dist/traits.d.ts.map +1 -0
  449. package/dist/traits.js +15 -0
  450. package/dist/traits.js.map +1 -0
  451. package/package.json +75 -0
  452. package/src/credentials.ts +97 -0
  453. package/src/errors.ts +108 -0
  454. package/src/index.ts +33 -0
  455. package/src/protocol.ts +160 -0
  456. package/src/retry.ts +74 -0
  457. package/src/services/a2a.ts +250 -0
  458. package/src/services/a2a_registration.ts +94 -0
  459. package/src/services/access_groups.ts +716 -0
  460. package/src/services/adaptive_router.ts +49 -0
  461. package/src/services/agents.ts +1314 -0
  462. package/src/services/alerting.ts +49 -0
  463. package/src/services/anthropic_pass_through.ts +233 -0
  464. package/src/services/anthropic_passthrough.ts +123 -0
  465. package/src/services/anthropic_skills.ts +200 -0
  466. package/src/services/assembly_ai_eu_pass_through.ts +237 -0
  467. package/src/services/assembly_ai_pass_through.ts +237 -0
  468. package/src/services/assistants.ts +685 -0
  469. package/src/services/audio.ts +189 -0
  470. package/src/services/audit_logging.ts +194 -0
  471. package/src/services/auto_router.ts +625 -0
  472. package/src/services/aws_comprehend_medical_pass_through.ts +109 -0
  473. package/src/services/azure_ai_pass_through.ts +231 -0
  474. package/src/services/azure_pass_through.ts +227 -0
  475. package/src/services/batch.ts +553 -0
  476. package/src/services/bedrock_pass_through.ts +229 -0
  477. package/src/services/beta_agents.ts +444 -0
  478. package/src/services/beta_mcp.ts +105 -0
  479. package/src/services/budget_management.ts +465 -0
  480. package/src/services/budget_spend_tracking.ts +2670 -0
  481. package/src/services/cache_settings.ts +234 -0
  482. package/src/services/caching.ts +175 -0
  483. package/src/services/chat_completions.ts +1788 -0
  484. package/src/services/claude_code_marketplace.ts +569 -0
  485. package/src/services/cloudzero.ts +297 -0
  486. package/src/services/cohere_pass_through.ts +227 -0
  487. package/src/services/completions.ts +194 -0
  488. package/src/services/compliance.ts +175 -0
  489. package/src/services/config_overrides.ts +466 -0
  490. package/src/services/config_yaml.ts +1537 -0
  491. package/src/services/containers.ts +848 -0
  492. package/src/services/coordination_redis_settings.ts +251 -0
  493. package/src/services/cost_tracking.ts +407 -0
  494. package/src/services/credential_management.ts +436 -0
  495. package/src/services/cursor_pass_through.ts +227 -0
  496. package/src/services/customer_management.ts +1475 -0
  497. package/src/services/email_management.ts +171 -0
  498. package/src/services/embeddings.ts +424 -0
  499. package/src/services/evals.ts +672 -0
  500. package/src/services/experimental.ts +572 -0
  501. package/src/services/fallback_management.ts +224 -0
  502. package/src/services/files.ts +732 -0
  503. package/src/services/fine_tuning.ts +547 -0
  504. package/src/services/gemini_agents.ts +227 -0
  505. package/src/services/google_ai_studio_pass_through.ts +227 -0
  506. package/src/services/google_genai_endpoints.ts +688 -0
  507. package/src/services/guardrails.ts +3006 -0
  508. package/src/services/health.ts +713 -0
  509. package/src/services/images.ts +430 -0
  510. package/src/services/index.ts +106 -0
  511. package/src/services/internal_user_management.ts +2574 -0
  512. package/src/services/invite_links.ts +167 -0
  513. package/src/services/jwt_mappings.ts +238 -0
  514. package/src/services/key_management.ts +4162 -0
  515. package/src/services/langfuse_passthrough.ts +231 -0
  516. package/src/services/llm_utils.ts +429 -0
  517. package/src/services/logging_callbacks.ts +104 -0
  518. package/src/services/mcp_byok_oauth.ts +468 -0
  519. package/src/services/mcp_discoverable.ts +626 -0
  520. package/src/services/mcp_management.ts +2915 -0
  521. package/src/services/mcp_rest.ts +779 -0
  522. package/src/services/memory_management.ts +248 -0
  523. package/src/services/milvus_pass_through.ts +227 -0
  524. package/src/services/misc.ts +2128 -0
  525. package/src/services/mistral_pass_through.ts +229 -0
  526. package/src/services/model_management.ts +4501 -0
  527. package/src/services/moderations.ts +80 -0
  528. package/src/services/ocr.ts +78 -0
  529. package/src/services/open_ai_pass_through.ts +462 -0
  530. package/src/services/organization_management.ts +1749 -0
  531. package/src/services/plugins.ts +367 -0
  532. package/src/services/policies.ts +1676 -0
  533. package/src/services/policy_engine.ts +1318 -0
  534. package/src/services/project_management.ts +938 -0
  535. package/src/services/prompts.ts +592 -0
  536. package/src/services/public.ts +880 -0
  537. package/src/services/rag.ts +148 -0
  538. package/src/services/realtime.ts +350 -0
  539. package/src/services/rerank.ts +111 -0
  540. package/src/services/responses.ts +909 -0
  541. package/src/services/router_settings.ts +168 -0
  542. package/src/services/rust_control_plane.ts +123 -0
  543. package/src/services/scim.ts +1243 -0
  544. package/src/services/search.ts +260 -0
  545. package/src/services/search_tools.ts +439 -0
  546. package/src/services/settings.ts +146 -0
  547. package/src/services/sso_settings.ts +626 -0
  548. package/src/services/tag_management.ts +1004 -0
  549. package/src/services/team_management.ts +3990 -0
  550. package/src/services/tools.ts +628 -0
  551. package/src/services/ui_settings.ts +237 -0
  552. package/src/services/ui_theme_settings.ts +173 -0
  553. package/src/services/usage_ai.ts +86 -0
  554. package/src/services/vantage.ts +283 -0
  555. package/src/services/vector_store_management.ts +454 -0
  556. package/src/services/vector_stores.ts +1321 -0
  557. package/src/services/vertex_ai_pass_through.ts +684 -0
  558. package/src/services/videos.ts +762 -0
  559. package/src/services/vllm_pass_through.ts +227 -0
  560. package/src/services/watsonx_pass_through.ts +229 -0
  561. package/src/services/web_socket.ts +254 -0
  562. package/src/services/workflow_management.ts +482 -0
  563. package/src/traits.ts +44 -0
@@ -0,0 +1,1649 @@
1
+ import * as S from "@distilled.cloud/core/schema";
2
+ import * as API from "@distilled.cloud/core/api";
3
+ import { type LitellmOpError, type LitellmOpContext } from "../protocol.ts";
4
+ export type { LitellmOpError, LitellmOpContext };
5
+ declare const BadRequest_base: S.Class<BadRequest, S.TaggedStruct<"BadRequest", {
6
+ readonly code: any;
7
+ readonly message: any;
8
+ }>, import("effect/Cause").YieldableError> & (new (...args: any[]) => {
9
+ "@distilled.cloud/error/categories": {
10
+ BadRequestError: true;
11
+ };
12
+ });
13
+ export declare class BadRequest extends /*@__PURE__*/ BadRequest_base {
14
+ }
15
+ declare const UnprocessableEntity_base: S.Class<UnprocessableEntity, S.TaggedStruct<"UnprocessableEntity", {
16
+ readonly code: any;
17
+ readonly message: any;
18
+ }>, import("effect/Cause").YieldableError> & (new (...args: any[]) => {
19
+ "@distilled.cloud/error/categories": {
20
+ BadRequestError: true;
21
+ };
22
+ });
23
+ export declare class UnprocessableEntity extends /*@__PURE__*/ UnprocessableEntity_base {
24
+ }
25
+ export type LiteLLMParamsAdaptiveRouterConfigMap = {
26
+ [key: string]: unknown | undefined;
27
+ };
28
+ export declare const LiteLLMParamsAdaptiveRouterConfigMap: S.Schema<LiteLLMParamsAdaptiveRouterConfigMap>;
29
+ export type LiteLLMParamsBedrockTagsList = Array<unknown>;
30
+ export declare const LiteLLMParamsBedrockTagsList: S.Schema<LiteLLMParamsBedrockTagsList>;
31
+ export type LiteLLMParamsComplexityRouterConfigMap = {
32
+ [key: string]: unknown | undefined;
33
+ };
34
+ export declare const LiteLLMParamsComplexityRouterConfigMap: S.Schema<LiteLLMParamsComplexityRouterConfigMap>;
35
+ export interface ConfigurableClientsideParamsCustomAuthInput {
36
+ api_base: string;
37
+ }
38
+ export declare const ConfigurableClientsideParamsCustomAuthInput: S.Schema<ConfigurableClientsideParamsCustomAuthInput>;
39
+ export type LiteLLMParamsConfigurableClientsideAuthParamsItem = string | ConfigurableClientsideParamsCustomAuthInput;
40
+ export declare const LiteLLMParamsConfigurableClientsideAuthParamsItem: S.Schema<LiteLLMParamsConfigurableClientsideAuthParamsItem>;
41
+ export type LiteLLMParamsConfigurableClientsideAuthParamsList = Array<LiteLLMParamsConfigurableClientsideAuthParamsItem>;
42
+ export declare const LiteLLMParamsConfigurableClientsideAuthParamsList: S.Schema<LiteLLMParamsConfigurableClientsideAuthParamsList>;
43
+ export type LiteLLMParamsMilvusPartitionNamesList = Array<string>;
44
+ export declare const LiteLLMParamsMilvusPartitionNamesList: S.Schema<LiteLLMParamsMilvusPartitionNamesList>;
45
+ export type ChoicesFinishReason = "stop" | "content_filter" | "function_call" | "tool_calls" | "length" | "guardrail_intervened" | "eos" | "finish_reason_unspecified" | "malformed_function_call";
46
+ export declare const ChoicesFinishReason: any;
47
+ export type ChatCompletionTokenLogprobBytesList = Array<number>;
48
+ export declare const ChatCompletionTokenLogprobBytesList: S.Schema<ChatCompletionTokenLogprobBytesList>;
49
+ export type TopLogprobBytesList = Array<number>;
50
+ export declare const TopLogprobBytesList: S.Schema<TopLogprobBytesList>;
51
+ export interface TopLogprob {
52
+ bytes?: TopLogprobBytesList | null;
53
+ logprob: number;
54
+ token: string;
55
+ }
56
+ export declare const TopLogprob: S.Schema<TopLogprob>;
57
+ export type ChatCompletionTokenLogprobTopLogprobsList = Array<TopLogprob>;
58
+ export declare const ChatCompletionTokenLogprobTopLogprobsList: S.Schema<ChatCompletionTokenLogprobTopLogprobsList>;
59
+ export interface ChatCompletionTokenLogprob {
60
+ bytes?: ChatCompletionTokenLogprobBytesList | null;
61
+ logprob: number;
62
+ token: string;
63
+ top_logprobs: ChatCompletionTokenLogprobTopLogprobsList;
64
+ }
65
+ export declare const ChatCompletionTokenLogprob: S.Schema<ChatCompletionTokenLogprob>;
66
+ export type ChoiceLogprobsContentList = Array<ChatCompletionTokenLogprob>;
67
+ export declare const ChoiceLogprobsContentList: S.Schema<ChoiceLogprobsContentList>;
68
+ export interface ChoiceLogprobs {
69
+ content?: ChoiceLogprobsContentList | null;
70
+ }
71
+ export declare const ChoiceLogprobs: S.Schema<ChoiceLogprobs>;
72
+ export type ChoicesLogprobs = ChoiceLogprobs | unknown;
73
+ export declare const ChoicesLogprobs: S.Schema<ChoicesLogprobs>;
74
+ export interface ChatCompletionAnnotationURLCitation {
75
+ end_index?: number;
76
+ start_index?: number;
77
+ title?: string;
78
+ url?: string;
79
+ }
80
+ export declare const ChatCompletionAnnotationURLCitation: S.Schema<ChatCompletionAnnotationURLCitation>;
81
+ export interface ChatCompletionAnnotation {
82
+ type?: string;
83
+ url_citation?: ChatCompletionAnnotationURLCitation;
84
+ }
85
+ export declare const ChatCompletionAnnotation: S.Schema<ChatCompletionAnnotation>;
86
+ export type MessageAnnotationsList = Array<ChatCompletionAnnotation>;
87
+ export declare const MessageAnnotationsList: S.Schema<MessageAnnotationsList>;
88
+ export interface ChatCompletionAudioResponse {
89
+ data: string;
90
+ expires_at: number;
91
+ id: string;
92
+ transcript: string;
93
+ }
94
+ export declare const ChatCompletionAudioResponse: S.Schema<ChatCompletionAudioResponse>;
95
+ export interface FunctionCall {
96
+ arguments: string;
97
+ name?: string | null;
98
+ }
99
+ export declare const FunctionCall: S.Schema<FunctionCall>;
100
+ export interface ImageURLObject {
101
+ detail?: string | null;
102
+ url: string;
103
+ }
104
+ export declare const ImageURLObject: S.Schema<ImageURLObject>;
105
+ export interface ImageURLListItem {
106
+ image_url: ImageURLObject;
107
+ index: number;
108
+ type: string;
109
+ }
110
+ export declare const ImageURLListItem: S.Schema<ImageURLListItem>;
111
+ export type MessageImagesList = Array<ImageURLListItem>;
112
+ export declare const MessageImagesList: S.Schema<MessageImagesList>;
113
+ export type MessageProviderSpecificFieldsMap = {
114
+ [key: string]: unknown | undefined;
115
+ };
116
+ export declare const MessageProviderSpecificFieldsMap: S.Schema<MessageProviderSpecificFieldsMap>;
117
+ export interface ChatCompletionReasoningSummaryTextBlock {
118
+ text?: string;
119
+ type: string;
120
+ }
121
+ export declare const ChatCompletionReasoningSummaryTextBlock: S.Schema<ChatCompletionReasoningSummaryTextBlock>;
122
+ export type ChatCompletionReasoningItemSummaryList = Array<ChatCompletionReasoningSummaryTextBlock>;
123
+ export declare const ChatCompletionReasoningItemSummaryList: S.Schema<ChatCompletionReasoningItemSummaryList>;
124
+ /** Represents an OpenAI Responses API reasoning item for round-tripping in conversation history. */
125
+ export interface ChatCompletionReasoningItem {
126
+ encrypted_content?: string | null;
127
+ id?: string;
128
+ summary?: ChatCompletionReasoningItemSummaryList;
129
+ type: string;
130
+ }
131
+ export declare const ChatCompletionReasoningItem: S.Schema<ChatCompletionReasoningItem>;
132
+ export type MessageReasoningItemsList = Array<ChatCompletionReasoningItem>;
133
+ export declare const MessageReasoningItemsList: S.Schema<MessageReasoningItemsList>;
134
+ export type MessageRole = "assistant" | "user" | "system" | "tool" | "function";
135
+ export declare const MessageRole: any;
136
+ export type ChatCompletionThinkingBlockCacheControlCase0Map = {
137
+ [key: string]: unknown | undefined;
138
+ };
139
+ export declare const ChatCompletionThinkingBlockCacheControlCase0Map: S.Schema<ChatCompletionThinkingBlockCacheControlCase0Map>;
140
+ export type ChatCompletionCachedContentTtl = "5m" | "1h";
141
+ export declare const ChatCompletionCachedContentTtl: any;
142
+ export interface ChatCompletionCachedContent {
143
+ ttl?: ChatCompletionCachedContentTtl | (string & {});
144
+ type: string;
145
+ }
146
+ export declare const ChatCompletionCachedContent: S.Schema<ChatCompletionCachedContent>;
147
+ export type ChatCompletionThinkingBlockCacheControl = ChatCompletionThinkingBlockCacheControlCase0Map | ChatCompletionCachedContent;
148
+ export declare const ChatCompletionThinkingBlockCacheControl: S.Schema<ChatCompletionThinkingBlockCacheControl>;
149
+ export interface ChatCompletionThinkingBlock {
150
+ cache_control?: ChatCompletionThinkingBlockCacheControl | null;
151
+ signature?: string | null;
152
+ thinking?: string;
153
+ type: string;
154
+ }
155
+ export declare const ChatCompletionThinkingBlock: S.Schema<ChatCompletionThinkingBlock>;
156
+ export type ChatCompletionRedactedThinkingBlockCacheControlCase0Map = {
157
+ [key: string]: unknown | undefined;
158
+ };
159
+ export declare const ChatCompletionRedactedThinkingBlockCacheControlCase0Map: S.Schema<ChatCompletionRedactedThinkingBlockCacheControlCase0Map>;
160
+ export type ChatCompletionRedactedThinkingBlockCacheControl = ChatCompletionRedactedThinkingBlockCacheControlCase0Map | ChatCompletionCachedContent;
161
+ export declare const ChatCompletionRedactedThinkingBlockCacheControl: S.Schema<ChatCompletionRedactedThinkingBlockCacheControl>;
162
+ export interface ChatCompletionRedactedThinkingBlock {
163
+ cache_control?: ChatCompletionRedactedThinkingBlockCacheControl | null;
164
+ data?: string;
165
+ type: string;
166
+ }
167
+ export declare const ChatCompletionRedactedThinkingBlock: S.Schema<ChatCompletionRedactedThinkingBlock>;
168
+ export type MessageThinkingBlocksItem = ChatCompletionThinkingBlock | ChatCompletionRedactedThinkingBlock;
169
+ export declare const MessageThinkingBlocksItem: S.Schema<MessageThinkingBlocksItem>;
170
+ export type MessageThinkingBlocksList = Array<MessageThinkingBlocksItem>;
171
+ export declare const MessageThinkingBlocksList: S.Schema<MessageThinkingBlocksList>;
172
+ export type ChatCompletionMessageToolCall = {
173
+ [key: string]: unknown | undefined;
174
+ };
175
+ export declare const ChatCompletionMessageToolCall: S.Schema<ChatCompletionMessageToolCall>;
176
+ export interface ChatCompletionCustomToolCallPayload {
177
+ input: string;
178
+ name: string;
179
+ }
180
+ export declare const ChatCompletionCustomToolCallPayload: S.Schema<ChatCompletionCustomToolCallPayload>;
181
+ export interface ChatCompletionMessageCustomToolCall {
182
+ custom: ChatCompletionCustomToolCallPayload;
183
+ id: string;
184
+ type?: string;
185
+ }
186
+ export declare const ChatCompletionMessageCustomToolCall: S.Schema<ChatCompletionMessageCustomToolCall>;
187
+ export type MessageToolCallsItem = ChatCompletionMessageToolCall | ChatCompletionMessageCustomToolCall;
188
+ export declare const MessageToolCallsItem: S.Schema<MessageToolCallsItem>;
189
+ export type MessageToolCallsList = Array<MessageToolCallsItem>;
190
+ export declare const MessageToolCallsList: S.Schema<MessageToolCallsList>;
191
+ export interface Message {
192
+ annotations?: MessageAnnotationsList | null;
193
+ audio?: ChatCompletionAudioResponse | null;
194
+ content: string | null;
195
+ function_call: FunctionCall | null;
196
+ images?: MessageImagesList | null;
197
+ provider_specific_fields?: MessageProviderSpecificFieldsMap | null;
198
+ reasoning_content?: string | null;
199
+ reasoning_items?: MessageReasoningItemsList | null;
200
+ role: MessageRole | (string & {});
201
+ thinking_blocks?: MessageThinkingBlocksList | null;
202
+ tool_calls: MessageToolCallsList | null;
203
+ }
204
+ export declare const Message: S.Schema<Message>;
205
+ export type ChoicesProviderSpecificFieldsMap = {
206
+ [key: string]: unknown | undefined;
207
+ };
208
+ export declare const ChoicesProviderSpecificFieldsMap: S.Schema<ChoicesProviderSpecificFieldsMap>;
209
+ export interface Choices {
210
+ finish_reason: ChoicesFinishReason | (string & {});
211
+ index: number;
212
+ logprobs?: ChoicesLogprobs | null;
213
+ message: Message;
214
+ provider_specific_fields?: ChoicesProviderSpecificFieldsMap | null;
215
+ }
216
+ export declare const Choices: S.Schema<Choices>;
217
+ export type ModelResponseChoicesList = Array<Choices>;
218
+ export declare const ModelResponseChoicesList: S.Schema<ModelResponseChoicesList>;
219
+ export interface ModelResponse {
220
+ choices: ModelResponseChoicesList;
221
+ created: number;
222
+ id: string;
223
+ model?: string | null;
224
+ object: string;
225
+ system_fingerprint?: string | null;
226
+ }
227
+ export declare const ModelResponse: S.Schema<ModelResponse>;
228
+ export type LiteLLMParamsMockResponse = string | ModelResponse | unknown;
229
+ export declare const LiteLLMParamsMockResponse: S.Schema<LiteLLMParamsMockResponse>;
230
+ export type LiteLLMParamsModelInfoMap = {
231
+ [key: string]: unknown | undefined;
232
+ };
233
+ export declare const LiteLLMParamsModelInfoMap: S.Schema<LiteLLMParamsModelInfoMap>;
234
+ export type LiteLLMParamsQualityRouterConfigMap = {
235
+ [key: string]: unknown | undefined;
236
+ };
237
+ export declare const LiteLLMParamsQualityRouterConfigMap: S.Schema<LiteLLMParamsQualityRouterConfigMap>;
238
+ export type LiteLLMParamsSearchContextCostPerQueryMap = {
239
+ [key: string]: unknown | undefined;
240
+ };
241
+ export declare const LiteLLMParamsSearchContextCostPerQueryMap: S.Schema<LiteLLMParamsSearchContextCostPerQueryMap>;
242
+ export type LiteLLMParamsStreamTimeout = number | string;
243
+ export declare const LiteLLMParamsStreamTimeout: S.Schema<LiteLLMParamsStreamTimeout>;
244
+ export type LiteLLMParamsTagRegexList = Array<string>;
245
+ export declare const LiteLLMParamsTagRegexList: S.Schema<LiteLLMParamsTagRegexList>;
246
+ export type LiteLLMParamsTagsList = Array<string>;
247
+ export declare const LiteLLMParamsTagsList: S.Schema<LiteLLMParamsTagsList>;
248
+ export type LiteLLMParamsTieredPricingItemMap = {
249
+ [key: string]: unknown | undefined;
250
+ };
251
+ export declare const LiteLLMParamsTieredPricingItemMap: S.Schema<LiteLLMParamsTieredPricingItemMap>;
252
+ export type LiteLLMParamsTieredPricingList = Array<LiteLLMParamsTieredPricingItemMap>;
253
+ export declare const LiteLLMParamsTieredPricingList: S.Schema<LiteLLMParamsTieredPricingList>;
254
+ export type LiteLLMParamsTimeout = number | string;
255
+ export declare const LiteLLMParamsTimeout: S.Schema<LiteLLMParamsTimeout>;
256
+ export type LiteLLMParamsVertexCredentialsCase1Map = {
257
+ [key: string]: unknown | undefined;
258
+ };
259
+ export declare const LiteLLMParamsVertexCredentialsCase1Map: S.Schema<LiteLLMParamsVertexCredentialsCase1Map>;
260
+ export type LiteLLMParamsVertexCredentials = string | LiteLLMParamsVertexCredentialsCase1Map;
261
+ export declare const LiteLLMParamsVertexCredentials: S.Schema<LiteLLMParamsVertexCredentials>;
262
+ /** LiteLLM Params with 'model' requirement - used for completions */
263
+ export interface LiteLLMParams {
264
+ adaptive_router_config?: LiteLLMParamsAdaptiveRouterConfigMap | null;
265
+ adaptive_router_default_model?: string | null;
266
+ allow_client_keepalive_override?: boolean | null;
267
+ annotation_cost_per_page?: number | null;
268
+ api_base?: string | null;
269
+ api_key?: string | null;
270
+ api_version?: string | null;
271
+ auto_router_config?: string | null;
272
+ auto_router_config_path?: string | null;
273
+ auto_router_default_model?: string | null;
274
+ auto_router_embedding_model?: string | null;
275
+ auto_router_max_input_chars?: number | null;
276
+ aws_access_key_id?: string | null;
277
+ aws_batch_role_arn?: string | null;
278
+ aws_bedrock_project_id?: string | null;
279
+ aws_bedrock_runtime_endpoint?: string | null;
280
+ aws_external_id?: string | null;
281
+ aws_profile_name?: string | null;
282
+ aws_region_name?: string | null;
283
+ aws_role_name?: string | null;
284
+ aws_secret_access_key?: string | null;
285
+ aws_session_name?: string | null;
286
+ aws_session_token?: string | null;
287
+ aws_sts_endpoint?: string | null;
288
+ aws_web_identity_token?: string | null;
289
+ azure_ad_token?: string | null;
290
+ bedrock_tags?: LiteLLMParamsBedrockTagsList | null;
291
+ budget_duration?: string | null;
292
+ cache_creation_input_audio_token_cost?: number | null;
293
+ cache_creation_input_token_cost?: number | null;
294
+ cache_creation_input_token_cost_above_1hr?: number | null;
295
+ cache_creation_input_token_cost_above_200k_tokens?: number | null;
296
+ cache_creation_input_token_cost_above_272k_tokens?: number | null;
297
+ cache_creation_input_token_cost_above_272k_tokens_flex?: number | null;
298
+ cache_creation_input_token_cost_above_272k_tokens_priority?: number | null;
299
+ cache_creation_input_token_cost_flex?: number | null;
300
+ cache_creation_input_token_cost_priority?: number | null;
301
+ cache_creation_input_token_cost_ultrafast?: number | null;
302
+ cache_read_input_audio_token_cost?: number | null;
303
+ cache_read_input_token_cost?: number | null;
304
+ cache_read_input_token_cost_above_200k_tokens?: number | null;
305
+ cache_read_input_token_cost_above_200k_tokens_priority?: number | null;
306
+ cache_read_input_token_cost_above_272k_tokens?: number | null;
307
+ cache_read_input_token_cost_above_272k_tokens_flex?: number | null;
308
+ cache_read_input_token_cost_above_272k_tokens_priority?: number | null;
309
+ cache_read_input_token_cost_above_512k_tokens?: number | null;
310
+ cache_read_input_token_cost_flex?: number | null;
311
+ cache_read_input_token_cost_priority?: number | null;
312
+ cache_read_input_token_cost_ultrafast?: number | null;
313
+ citation_cost_per_token?: number | null;
314
+ complexity_router_config?: LiteLLMParamsComplexityRouterConfigMap | null;
315
+ complexity_router_default_model?: string | null;
316
+ configurable_clientside_auth_params?: LiteLLMParamsConfigurableClientsideAuthParamsList | null;
317
+ custom_llm_provider?: string | null;
318
+ default_api_key_rpm_limit?: number | null;
319
+ default_api_key_tpm_limit?: number | null;
320
+ gcs_bucket_name?: string | null;
321
+ google_maps_grounding_cost_per_query?: number | null;
322
+ input_cost_per_audio_per_second?: number | null;
323
+ input_cost_per_audio_per_second_above_128k_tokens?: number | null;
324
+ input_cost_per_audio_token?: number | null;
325
+ input_cost_per_character?: number | null;
326
+ input_cost_per_character_above_128k_tokens?: number | null;
327
+ input_cost_per_image?: number | null;
328
+ input_cost_per_image_above_128k_tokens?: number | null;
329
+ input_cost_per_image_token?: number | null;
330
+ input_cost_per_pixel?: number | null;
331
+ input_cost_per_query?: number | null;
332
+ input_cost_per_second?: number | null;
333
+ input_cost_per_token?: number | null;
334
+ input_cost_per_token_above_128k_tokens?: number | null;
335
+ input_cost_per_token_above_200k_tokens?: number | null;
336
+ input_cost_per_token_above_200k_tokens_priority?: number | null;
337
+ input_cost_per_token_above_272k_tokens?: number | null;
338
+ input_cost_per_token_above_272k_tokens_flex?: number | null;
339
+ input_cost_per_token_above_272k_tokens_priority?: number | null;
340
+ input_cost_per_token_above_512k_tokens?: number | null;
341
+ input_cost_per_token_batches?: number | null;
342
+ input_cost_per_token_cache_hit?: number | null;
343
+ input_cost_per_token_flex?: number | null;
344
+ input_cost_per_token_priority?: number | null;
345
+ input_cost_per_token_ultrafast?: number | null;
346
+ input_cost_per_video_per_second?: number | null;
347
+ input_cost_per_video_per_second_above_128k_tokens?: number | null;
348
+ input_cost_per_video_per_second_above_15s_interval?: number | null;
349
+ input_cost_per_video_per_second_above_8s_interval?: number | null;
350
+ input_cost_per_video_token?: number | null;
351
+ itpm?: number | null;
352
+ keepalive_seconds?: number | null;
353
+ litellm_credential_name?: string | null;
354
+ litellm_trace_id?: string | null;
355
+ max_budget?: number | null;
356
+ max_file_size_mb?: number | null;
357
+ max_retries?: number | null;
358
+ merge_reasoning_content_in_choices?: boolean | null;
359
+ milvus_db_name?: string | null;
360
+ milvus_partition_names?: LiteLLMParamsMilvusPartitionNamesList | null;
361
+ milvus_text_field?: string | null;
362
+ mock_response?: LiteLLMParamsMockResponse | null;
363
+ model: string;
364
+ model_info?: LiteLLMParamsModelInfoMap | null;
365
+ ocr_cost_per_credit?: number | null;
366
+ ocr_cost_per_page?: number | null;
367
+ organization?: string | null;
368
+ otpm?: number | null;
369
+ output_cost_per_audio_per_second?: number | null;
370
+ output_cost_per_audio_token?: number | null;
371
+ output_cost_per_character?: number | null;
372
+ output_cost_per_character_above_128k_tokens?: number | null;
373
+ output_cost_per_image?: number | null;
374
+ output_cost_per_image_token?: number | null;
375
+ output_cost_per_pixel?: number | null;
376
+ output_cost_per_reasoning_token?: number | null;
377
+ output_cost_per_reasoning_token_flex?: number | null;
378
+ output_cost_per_reasoning_token_priority?: number | null;
379
+ output_cost_per_second?: number | null;
380
+ output_cost_per_second_1080p?: number | null;
381
+ output_cost_per_second_480p?: number | null;
382
+ output_cost_per_second_4k?: number | null;
383
+ output_cost_per_token?: number | null;
384
+ output_cost_per_token_above_128k_tokens?: number | null;
385
+ output_cost_per_token_above_200k_tokens?: number | null;
386
+ output_cost_per_token_above_200k_tokens_priority?: number | null;
387
+ output_cost_per_token_above_272k_tokens?: number | null;
388
+ output_cost_per_token_above_272k_tokens_flex?: number | null;
389
+ output_cost_per_token_above_272k_tokens_priority?: number | null;
390
+ output_cost_per_token_above_512k_tokens?: number | null;
391
+ output_cost_per_token_batches?: number | null;
392
+ output_cost_per_token_flex?: number | null;
393
+ output_cost_per_token_priority?: number | null;
394
+ output_cost_per_token_ultrafast?: number | null;
395
+ output_cost_per_video_per_second?: number | null;
396
+ output_cost_per_video_token?: number | null;
397
+ output_vector_size?: number | null;
398
+ quality_router_config?: LiteLLMParamsQualityRouterConfigMap | null;
399
+ quality_router_default_model?: string | null;
400
+ region_name?: string | null;
401
+ regional_endpoint_uplift_multiplier?: number | null;
402
+ regional_processing_uplift_multiplier_eu?: number | null;
403
+ regional_processing_uplift_multiplier_us?: number | null;
404
+ rpm?: number | null;
405
+ s3_bucket_name?: string | null;
406
+ s3_encryption_key_id?: string | null;
407
+ s3_output_bucket_name?: string | null;
408
+ s3_region_name?: string | null;
409
+ search_context_cost_per_query?: LiteLLMParamsSearchContextCostPerQueryMap | null;
410
+ stream_timeout?: LiteLLMParamsStreamTimeout | null;
411
+ tag_regex?: LiteLLMParamsTagRegexList | null;
412
+ tags?: LiteLLMParamsTagsList | null;
413
+ tiered_pricing?: LiteLLMParamsTieredPricingList | null;
414
+ timeout?: LiteLLMParamsTimeout | null;
415
+ tpm?: number | null;
416
+ use_chat_completions_api?: boolean | null;
417
+ use_in_pass_through?: boolean | null;
418
+ use_litellm_proxy?: boolean | null;
419
+ /** Use stored xAI OAuth credentials when no xAI API key is configured. */
420
+ use_xai_oauth?: boolean | null;
421
+ valkey_embedding_field?: string | null;
422
+ valkey_host?: string | null;
423
+ valkey_password?: string | null;
424
+ valkey_port?: number | null;
425
+ valkey_ssl?: boolean | null;
426
+ valkey_text_field?: string | null;
427
+ vector_store_id?: string | null;
428
+ vertex_credentials?: LiteLLMParamsVertexCredentials | null;
429
+ vertex_location?: string | null;
430
+ vertex_project?: string | null;
431
+ watsonx_region_name?: string | null;
432
+ }
433
+ export declare const LiteLLMParams: S.Schema<LiteLLMParams>;
434
+ export type LitellmTypesRouterModelInfoTier = "free" | "paid";
435
+ export declare const LitellmTypesRouterModelInfoTier: any;
436
+ export type LitellmTypesRouterModelInfoTieredPricingItemMap = {
437
+ [key: string]: unknown | undefined;
438
+ };
439
+ export declare const LitellmTypesRouterModelInfoTieredPricingItemMap: S.Schema<LitellmTypesRouterModelInfoTieredPricingItemMap>;
440
+ export type LitellmTypesRouterModelInfoTieredPricingList = Array<LitellmTypesRouterModelInfoTieredPricingItemMap>;
441
+ export declare const LitellmTypesRouterModelInfoTieredPricingList: S.Schema<LitellmTypesRouterModelInfoTieredPricingList>;
442
+ export interface LitellmTypesRouterModelInfo {
443
+ allow_fail_open?: boolean | null;
444
+ base_model?: string | null;
445
+ blocked?: boolean | null;
446
+ cache_creation_input_token_cost?: number | null;
447
+ cache_read_input_token_cost?: number | null;
448
+ cost_per_ptu_per_hour?: number | null;
449
+ created_at?: string | null;
450
+ created_by?: string | null;
451
+ db_model?: boolean;
452
+ enable_tag_filtering?: boolean | null;
453
+ id: string | null;
454
+ input_cost_per_character?: number | null;
455
+ input_cost_per_token?: number | null;
456
+ output_cost_per_character?: number | null;
457
+ output_cost_per_token?: number | null;
458
+ ptu_count?: number | null;
459
+ ptu_effective_from?: string | null;
460
+ ptu_effective_to?: string | null;
461
+ team_id?: string | null;
462
+ team_public_model_name?: string | null;
463
+ tier?: LitellmTypesRouterModelInfoTier | (string & {}) | null;
464
+ tiered_pricing?: LitellmTypesRouterModelInfoTieredPricingList | null;
465
+ updated_at?: string | null;
466
+ updated_by?: string | null;
467
+ }
468
+ export declare const LitellmTypesRouterModelInfo: S.Schema<LitellmTypesRouterModelInfo>;
469
+ export interface AddNewModelModelNewPostRequest {
470
+ litellm_params: LiteLLMParams;
471
+ model_info: LitellmTypesRouterModelInfo;
472
+ model_name: string;
473
+ }
474
+ export declare const AddNewModelModelNewPostRequest: S.Schema<AddNewModelModelNewPostRequest>;
475
+ export interface AddNewModelModelNewPostResponse {
476
+ body: unknown;
477
+ }
478
+ export declare const AddNewModelModelNewPostResponse: S.Schema<AddNewModelModelNewPostResponse>;
479
+ export interface CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest {
480
+ }
481
+ export declare const CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest: S.Schema<CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest>;
482
+ export interface CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse {
483
+ body: unknown;
484
+ }
485
+ export declare const CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse: S.Schema<CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse>;
486
+ export interface CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest {
487
+ }
488
+ export declare const CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest: S.Schema<CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest>;
489
+ export interface CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse {
490
+ body: unknown;
491
+ }
492
+ export declare const CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse: S.Schema<CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse>;
493
+ export type CreateModelGroupAccessGroupNewPostRequestModelIdsList = Array<string>;
494
+ export declare const CreateModelGroupAccessGroupNewPostRequestModelIdsList: S.Schema<CreateModelGroupAccessGroupNewPostRequestModelIdsList>;
495
+ export type CreateModelGroupAccessGroupNewPostRequestModelNamesList = Array<string>;
496
+ export declare const CreateModelGroupAccessGroupNewPostRequestModelNamesList: S.Schema<CreateModelGroupAccessGroupNewPostRequestModelNamesList>;
497
+ export interface CreateModelGroupAccessGroupNewPostRequest {
498
+ access_group: string;
499
+ model_ids?: CreateModelGroupAccessGroupNewPostRequestModelIdsList | null;
500
+ model_names?: CreateModelGroupAccessGroupNewPostRequestModelNamesList | null;
501
+ }
502
+ export declare const CreateModelGroupAccessGroupNewPostRequest: S.Schema<CreateModelGroupAccessGroupNewPostRequest>;
503
+ export type NewModelGroupResponseModelIdsList = Array<string>;
504
+ export declare const NewModelGroupResponseModelIdsList: S.Schema<NewModelGroupResponseModelIdsList>;
505
+ export type NewModelGroupResponseModelNamesList = Array<string>;
506
+ export declare const NewModelGroupResponseModelNamesList: S.Schema<NewModelGroupResponseModelNamesList>;
507
+ export interface NewModelGroupResponse {
508
+ access_group: string;
509
+ model_ids?: NewModelGroupResponseModelIdsList | null;
510
+ model_names?: NewModelGroupResponseModelNamesList | null;
511
+ models_updated: number;
512
+ }
513
+ export declare const NewModelGroupResponse: S.Schema<NewModelGroupResponse>;
514
+ export interface DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest {
515
+ access_group: string;
516
+ }
517
+ export declare const DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest: S.Schema<DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest>;
518
+ export interface DeleteModelGroupResponse {
519
+ access_group: string;
520
+ message: string;
521
+ models_updated: number;
522
+ }
523
+ export declare const DeleteModelGroupResponse: S.Schema<DeleteModelGroupResponse>;
524
+ export interface DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest {
525
+ access_group: string;
526
+ }
527
+ export declare const DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest: S.Schema<DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest>;
528
+ export interface DeleteAccessGroupBudgetResponse {
529
+ access_group: string;
530
+ budget_deleted: boolean;
531
+ message: string;
532
+ }
533
+ export declare const DeleteAccessGroupBudgetResponse: S.Schema<DeleteAccessGroupBudgetResponse>;
534
+ export interface DeleteModelModelDeletePostRequest {
535
+ id: string;
536
+ }
537
+ export declare const DeleteModelModelDeletePostRequest: S.Schema<DeleteModelModelDeletePostRequest>;
538
+ export interface DeleteModelModelDeletePostResponse {
539
+ body: unknown;
540
+ }
541
+ export declare const DeleteModelModelDeletePostResponse: S.Schema<DeleteModelModelDeletePostResponse>;
542
+ export interface GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest {
543
+ access_group: string;
544
+ }
545
+ export declare const GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest: S.Schema<GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest>;
546
+ export interface AccessGroupBudget {
547
+ budget_duration?: string | null;
548
+ budget_id: string;
549
+ budget_reset_at?: string | null;
550
+ max_budget?: number | null;
551
+ soft_budget?: number | null;
552
+ }
553
+ export declare const AccessGroupBudget: S.Schema<AccessGroupBudget>;
554
+ export interface AccessGroupBudgetResponse {
555
+ access_group: string;
556
+ budget?: AccessGroupBudget | null;
557
+ spend: number;
558
+ }
559
+ export declare const AccessGroupBudgetResponse: S.Schema<AccessGroupBudgetResponse>;
560
+ export interface GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest {
561
+ access_group: string;
562
+ }
563
+ export declare const GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest: S.Schema<GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest>;
564
+ export type AccessGroupInfoModelNamesList = Array<string>;
565
+ export declare const AccessGroupInfoModelNamesList: S.Schema<AccessGroupInfoModelNamesList>;
566
+ export interface AccessGroupInfo {
567
+ access_group: string;
568
+ budget?: AccessGroupBudget | null;
569
+ deployment_count: number;
570
+ model_names: AccessGroupInfoModelNamesList;
571
+ spend?: number | null;
572
+ }
573
+ export declare const AccessGroupInfo: S.Schema<AccessGroupInfo>;
574
+ export interface GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest {
575
+ }
576
+ export declare const GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest: S.Schema<GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest>;
577
+ export interface GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse {
578
+ body: unknown;
579
+ }
580
+ export declare const GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse: S.Schema<GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse>;
581
+ /** Which calibration examples, and for BUSINESS which tier criteria, the built-in classifier rubric carries. */
582
+ export type ClassificationRubric = "legacy" | "agentic" | "chat" | "business";
583
+ export declare const ClassificationRubric: any;
584
+ export interface GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest {
585
+ context_window_size?: number;
586
+ tier_labels?: string;
587
+ classification_rubric?: ClassificationRubric | (string & {});
588
+ }
589
+ export declare const GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest: S.Schema<GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest>;
590
+ /** The built-in system prompt an auto-router's LLM classifier uses when none is configured. Served so the dashboard's prompt editor prefills the rubric the proxy actually sends, rather than a copy in the frontend that drifts the moment the rubric is edited. */
591
+ export interface AutoRouterClassifierDefaultPromptResponse {
592
+ system_prompt: string;
593
+ }
594
+ export declare const AutoRouterClassifierDefaultPromptResponse: S.Schema<AutoRouterClassifierDefaultPromptResponse>;
595
+ export interface GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest {
596
+ }
597
+ export declare const GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest: S.Schema<GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest>;
598
+ export interface GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse {
599
+ body: unknown;
600
+ }
601
+ export declare const GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse: S.Schema<GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse>;
602
+ export interface GetModelCostMapSourceModelCostMapSourceGetRequest {
603
+ }
604
+ export declare const GetModelCostMapSourceModelCostMapSourceGetRequest: S.Schema<GetModelCostMapSourceModelCostMapSourceGetRequest>;
605
+ export interface GetModelCostMapSourceModelCostMapSourceGetResponse {
606
+ body: unknown;
607
+ }
608
+ export declare const GetModelCostMapSourceModelCostMapSourceGetResponse: S.Schema<GetModelCostMapSourceModelCostMapSourceGetResponse>;
609
+ export interface GetModelDeprecationsModelDeprecationRequest {
610
+ warn_within_days?: number;
611
+ }
612
+ export declare const GetModelDeprecationsModelDeprecationRequest: S.Schema<GetModelDeprecationsModelDeprecationRequest>;
613
+ /** 'deprecated' if the date has passed, 'imminent' if it falls within warn_within_days, 'upcoming' otherwise. */
614
+ export type ModelDeprecationInfoStatus = "upcoming" | "imminent" | "deprecated";
615
+ export declare const ModelDeprecationInfoStatus: any;
616
+ export interface ModelDeprecationInfo {
617
+ /** Days remaining until the deprecation date. Negative if the model is already deprecated. */
618
+ days_until_deprecation: number;
619
+ /** The date (UTC) when the model becomes deprecated. */
620
+ deprecation_date: string;
621
+ /** The underlying litellm model string the deprecation date is sourced from. */
622
+ litellm_model?: string | null;
623
+ /** The provider this model belongs to. */
624
+ litellm_provider?: string | null;
625
+ /** The public name of the model on the proxy (model_group). */
626
+ model_name: string;
627
+ /** 'deprecated' if the date has passed, 'imminent' if it falls within warn_within_days, 'upcoming' otherwise. */
628
+ status: ModelDeprecationInfoStatus;
629
+ }
630
+ export declare const ModelDeprecationInfo: S.Schema<ModelDeprecationInfo>;
631
+ /** Models whose deprecation date has already passed. */
632
+ export type ModelDeprecationResponseDeprecatedList = Array<ModelDeprecationInfo>;
633
+ export declare const ModelDeprecationResponseDeprecatedList: S.Schema<ModelDeprecationResponseDeprecatedList>;
634
+ /** Models whose deprecation date is within warn_within_days from today and require immediate migration planning. */
635
+ export type ModelDeprecationResponseImminentList = Array<ModelDeprecationInfo>;
636
+ export declare const ModelDeprecationResponseImminentList: S.Schema<ModelDeprecationResponseImminentList>;
637
+ /** Models with a future deprecation date outside the warn window. */
638
+ export type ModelDeprecationResponseUpcomingList = Array<ModelDeprecationInfo>;
639
+ export declare const ModelDeprecationResponseUpcomingList: S.Schema<ModelDeprecationResponseUpcomingList>;
640
+ export interface ModelDeprecationResponse {
641
+ /** UTC timestamp when the deprecation snapshot was generated. */
642
+ checked_at: string;
643
+ /** Models whose deprecation date has already passed. */
644
+ deprecated?: ModelDeprecationResponseDeprecatedList;
645
+ /** Models whose deprecation date is within warn_within_days from today and require immediate migration planning. */
646
+ imminent?: ModelDeprecationResponseImminentList;
647
+ /** Models with a future deprecation date outside the warn window. */
648
+ upcoming?: ModelDeprecationResponseUpcomingList;
649
+ /** The window (in days) used to bucket 'imminent' models. */
650
+ warn_within_days: number;
651
+ }
652
+ export declare const ModelDeprecationResponse: S.Schema<ModelDeprecationResponse>;
653
+ export interface GetModelDeprecationsV1ModelDeprecationRequest {
654
+ warn_within_days?: number;
655
+ }
656
+ export declare const GetModelDeprecationsV1ModelDeprecationRequest: S.Schema<GetModelDeprecationsV1ModelDeprecationRequest>;
657
+ export interface GetModelGroupInfoModelGroupInfoRequest {
658
+ model_group?: string;
659
+ }
660
+ export declare const GetModelGroupInfoModelGroupInfoRequest: S.Schema<GetModelGroupInfoModelGroupInfoRequest>;
661
+ export interface GetModelGroupInfoModelGroupInfoResponse {
662
+ body: unknown;
663
+ }
664
+ export declare const GetModelGroupInfoModelGroupInfoResponse: S.Schema<GetModelGroupInfoModelGroupInfoResponse>;
665
+ export interface GetModelInfoModelsModelIdRequest {
666
+ model_id: string;
667
+ team_id?: string;
668
+ healthy_only?: boolean;
669
+ }
670
+ export declare const GetModelInfoModelsModelIdRequest: S.Schema<GetModelInfoModelsModelIdRequest>;
671
+ export interface GetModelInfoModelsModelIdResponse {
672
+ body: unknown;
673
+ }
674
+ export declare const GetModelInfoModelsModelIdResponse: S.Schema<GetModelInfoModelsModelIdResponse>;
675
+ export interface GetModelInfoV1ModelInfoRequest {
676
+ litellm_model_id?: string;
677
+ /** When true, filter to deployments the caller can use via direct access or team membership. */
678
+ include_team_models?: boolean;
679
+ /** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
680
+ teamId?: string;
681
+ healthy_only?: boolean;
682
+ }
683
+ export declare const GetModelInfoV1ModelInfoRequest: S.Schema<GetModelInfoV1ModelInfoRequest>;
684
+ export interface GetModelInfoV1ModelInfoResponse {
685
+ body: unknown;
686
+ }
687
+ export declare const GetModelInfoV1ModelInfoResponse: S.Schema<GetModelInfoV1ModelInfoResponse>;
688
+ export interface GetModelInfoV1ModelsModelIdRequest {
689
+ model_id: string;
690
+ team_id?: string;
691
+ healthy_only?: boolean;
692
+ }
693
+ export declare const GetModelInfoV1ModelsModelIdRequest: S.Schema<GetModelInfoV1ModelsModelIdRequest>;
694
+ export interface GetModelInfoV1ModelsModelIdResponse {
695
+ body: unknown;
696
+ }
697
+ export declare const GetModelInfoV1ModelsModelIdResponse: S.Schema<GetModelInfoV1ModelsModelIdResponse>;
698
+ export interface GetModelInfoV1V1ModelInfoRequest {
699
+ litellm_model_id?: string;
700
+ /** When true, filter to deployments the caller can use via direct access or team membership. */
701
+ include_team_models?: boolean;
702
+ /** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
703
+ teamId?: string;
704
+ healthy_only?: boolean;
705
+ }
706
+ export declare const GetModelInfoV1V1ModelInfoRequest: S.Schema<GetModelInfoV1V1ModelInfoRequest>;
707
+ export interface GetModelInfoV1V1ModelInfoResponse {
708
+ body: unknown;
709
+ }
710
+ export declare const GetModelInfoV1V1ModelInfoResponse: S.Schema<GetModelInfoV1V1ModelInfoResponse>;
711
+ export interface GetModelInfoV2V2ModelInfoRequest {
712
+ /** Specify the model name (optional) */
713
+ model?: string;
714
+ /** Only return models added by this user */
715
+ user_models_only?: boolean;
716
+ /** Return all models across all teams user is in. */
717
+ include_team_models?: boolean;
718
+ debug?: boolean;
719
+ /** Page number */
720
+ page?: number;
721
+ /** Page size */
722
+ size?: number;
723
+ /** Search model names (case-insensitive partial match) */
724
+ search?: string;
725
+ /** Search for a specific model by its unique ID */
726
+ modelId?: string;
727
+ /** Filter models by team ID. Returns models with direct_access=True or teamId in access_via_team_ids */
728
+ teamId?: string;
729
+ /** Field to sort by. Options: model_name, created_at, updated_at, costs, status */
730
+ sortBy?: string;
731
+ /** Sort order. Options: asc, desc */
732
+ sortOrder?: string;
733
+ /** Omit auto-router deployments (litellm model prefixed `auto_router/`). They select among deployments rather than being deployments themselves, so a caller rendering a deployment list can leave them out. Defaults to false, so existing callers are unaffected */
734
+ exclude_auto_routers?: boolean;
735
+ }
736
+ export declare const GetModelInfoV2V2ModelInfoRequest: S.Schema<GetModelInfoV2V2ModelInfoRequest>;
737
+ export interface GetModelInfoV2V2ModelInfoResponse {
738
+ body: unknown;
739
+ }
740
+ export declare const GetModelInfoV2V2ModelInfoResponse: S.Schema<GetModelInfoV2V2ModelInfoResponse>;
741
+ export interface GetModelMetricsExceptionsModelMetricsExceptionRequest {
742
+ _selected_model_group?: string;
743
+ startTime?: string;
744
+ endTime?: string;
745
+ api_key?: string;
746
+ customer?: string;
747
+ }
748
+ export declare const GetModelMetricsExceptionsModelMetricsExceptionRequest: S.Schema<GetModelMetricsExceptionsModelMetricsExceptionRequest>;
749
+ export interface GetModelMetricsExceptionsModelMetricsExceptionResponse {
750
+ body: unknown;
751
+ }
752
+ export declare const GetModelMetricsExceptionsModelMetricsExceptionResponse: S.Schema<GetModelMetricsExceptionsModelMetricsExceptionResponse>;
753
+ export interface GetModelMetricsModelMetricsRequest {
754
+ _selected_model_group?: string;
755
+ startTime?: string;
756
+ endTime?: string;
757
+ api_key?: string;
758
+ customer?: string;
759
+ }
760
+ export declare const GetModelMetricsModelMetricsRequest: S.Schema<GetModelMetricsModelMetricsRequest>;
761
+ export interface GetModelMetricsModelMetricsResponse {
762
+ body: unknown;
763
+ }
764
+ export declare const GetModelMetricsModelMetricsResponse: S.Schema<GetModelMetricsModelMetricsResponse>;
765
+ export interface GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest {
766
+ _selected_model_group?: string;
767
+ startTime?: string;
768
+ endTime?: string;
769
+ api_key?: string;
770
+ customer?: string;
771
+ }
772
+ export declare const GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest: S.Schema<GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest>;
773
+ export interface GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse {
774
+ body: unknown;
775
+ }
776
+ export declare const GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse: S.Schema<GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse>;
777
+ export interface GetModelSettingsModelSettingsRequest {
778
+ }
779
+ export declare const GetModelSettingsModelSettingsRequest: S.Schema<GetModelSettingsModelSettingsRequest>;
780
+ export interface GetModelSettingsModelSettingsResponse {
781
+ body: unknown;
782
+ }
783
+ export declare const GetModelSettingsModelSettingsResponse: S.Schema<GetModelSettingsModelSettingsResponse>;
784
+ export interface GetModelStreamingMetricsModelStreamingMetricsRequest {
785
+ _selected_model_group?: string;
786
+ startTime?: string;
787
+ endTime?: string;
788
+ }
789
+ export declare const GetModelStreamingMetricsModelStreamingMetricsRequest: S.Schema<GetModelStreamingMetricsModelStreamingMetricsRequest>;
790
+ export interface GetModelStreamingMetricsModelStreamingMetricsResponse {
791
+ body: unknown;
792
+ }
793
+ export declare const GetModelStreamingMetricsModelStreamingMetricsResponse: S.Schema<GetModelStreamingMetricsModelStreamingMetricsResponse>;
794
+ export interface ListAccessGroupsAccessGroupListGetRequest {
795
+ }
796
+ export declare const ListAccessGroupsAccessGroupListGetRequest: S.Schema<ListAccessGroupsAccessGroupListGetRequest>;
797
+ export type ListAccessGroupsResponseAccessGroupsList = Array<AccessGroupInfo>;
798
+ export declare const ListAccessGroupsResponseAccessGroupsList: S.Schema<ListAccessGroupsResponseAccessGroupsList>;
799
+ export interface ListAccessGroupsResponse {
800
+ access_groups: ListAccessGroupsResponseAccessGroupsList;
801
+ }
802
+ export declare const ListAccessGroupsResponse: S.Schema<ListAccessGroupsResponse>;
803
+ export interface ModelListModelsGetRequest {
804
+ return_wildcard_routes?: boolean;
805
+ team_id?: string;
806
+ include_model_access_groups?: boolean;
807
+ only_model_access_groups?: boolean;
808
+ include_metadata?: boolean;
809
+ fallback_type?: string;
810
+ scope?: string;
811
+ healthy_only?: boolean;
812
+ }
813
+ export declare const ModelListModelsGetRequest: S.Schema<ModelListModelsGetRequest>;
814
+ export interface ModelListModelsGetResponse {
815
+ body: unknown;
816
+ }
817
+ export declare const ModelListModelsGetResponse: S.Schema<ModelListModelsGetResponse>;
818
+ export interface ModelListV1ModelsGetRequest {
819
+ return_wildcard_routes?: boolean;
820
+ team_id?: string;
821
+ include_model_access_groups?: boolean;
822
+ only_model_access_groups?: boolean;
823
+ include_metadata?: boolean;
824
+ fallback_type?: string;
825
+ scope?: string;
826
+ healthy_only?: boolean;
827
+ }
828
+ export declare const ModelListV1ModelsGetRequest: S.Schema<ModelListV1ModelsGetRequest>;
829
+ export interface ModelListV1ModelsGetResponse {
830
+ body: unknown;
831
+ }
832
+ export declare const ModelListV1ModelsGetResponse: S.Schema<ModelListV1ModelsGetResponse>;
833
+ export type UpdateLiteLLMParamsAdaptiveRouterConfigMap = {
834
+ [key: string]: unknown | undefined;
835
+ };
836
+ export declare const UpdateLiteLLMParamsAdaptiveRouterConfigMap: S.Schema<UpdateLiteLLMParamsAdaptiveRouterConfigMap>;
837
+ export type UpdateLiteLLMParamsBedrockTagsList = Array<unknown>;
838
+ export declare const UpdateLiteLLMParamsBedrockTagsList: S.Schema<UpdateLiteLLMParamsBedrockTagsList>;
839
+ export type UpdateLiteLLMParamsComplexityRouterConfigMap = {
840
+ [key: string]: unknown | undefined;
841
+ };
842
+ export declare const UpdateLiteLLMParamsComplexityRouterConfigMap: S.Schema<UpdateLiteLLMParamsComplexityRouterConfigMap>;
843
+ export type UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem = string | ConfigurableClientsideParamsCustomAuthInput;
844
+ export declare const UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem: S.Schema<UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem>;
845
+ export type UpdateLiteLLMParamsConfigurableClientsideAuthParamsList = Array<UpdateLiteLLMParamsConfigurableClientsideAuthParamsItem>;
846
+ export declare const UpdateLiteLLMParamsConfigurableClientsideAuthParamsList: S.Schema<UpdateLiteLLMParamsConfigurableClientsideAuthParamsList>;
847
+ export type UpdateLiteLLMParamsMilvusPartitionNamesList = Array<string>;
848
+ export declare const UpdateLiteLLMParamsMilvusPartitionNamesList: S.Schema<UpdateLiteLLMParamsMilvusPartitionNamesList>;
849
+ export type UpdateLiteLLMParamsMockResponse = string | ModelResponse | unknown;
850
+ export declare const UpdateLiteLLMParamsMockResponse: S.Schema<UpdateLiteLLMParamsMockResponse>;
851
+ export type UpdateLiteLLMParamsModelInfoMap = {
852
+ [key: string]: unknown | undefined;
853
+ };
854
+ export declare const UpdateLiteLLMParamsModelInfoMap: S.Schema<UpdateLiteLLMParamsModelInfoMap>;
855
+ export type UpdateLiteLLMParamsQualityRouterConfigMap = {
856
+ [key: string]: unknown | undefined;
857
+ };
858
+ export declare const UpdateLiteLLMParamsQualityRouterConfigMap: S.Schema<UpdateLiteLLMParamsQualityRouterConfigMap>;
859
+ export type UpdateLiteLLMParamsSearchContextCostPerQueryMap = {
860
+ [key: string]: unknown | undefined;
861
+ };
862
+ export declare const UpdateLiteLLMParamsSearchContextCostPerQueryMap: S.Schema<UpdateLiteLLMParamsSearchContextCostPerQueryMap>;
863
+ export type UpdateLiteLLMParamsStreamTimeout = number | string;
864
+ export declare const UpdateLiteLLMParamsStreamTimeout: S.Schema<UpdateLiteLLMParamsStreamTimeout>;
865
+ export type UpdateLiteLLMParamsTagRegexList = Array<string>;
866
+ export declare const UpdateLiteLLMParamsTagRegexList: S.Schema<UpdateLiteLLMParamsTagRegexList>;
867
+ export type UpdateLiteLLMParamsTagsList = Array<string>;
868
+ export declare const UpdateLiteLLMParamsTagsList: S.Schema<UpdateLiteLLMParamsTagsList>;
869
+ export type UpdateLiteLLMParamsTieredPricingItemMap = {
870
+ [key: string]: unknown | undefined;
871
+ };
872
+ export declare const UpdateLiteLLMParamsTieredPricingItemMap: S.Schema<UpdateLiteLLMParamsTieredPricingItemMap>;
873
+ export type UpdateLiteLLMParamsTieredPricingList = Array<UpdateLiteLLMParamsTieredPricingItemMap>;
874
+ export declare const UpdateLiteLLMParamsTieredPricingList: S.Schema<UpdateLiteLLMParamsTieredPricingList>;
875
+ export type UpdateLiteLLMParamsTimeout = number | string;
876
+ export declare const UpdateLiteLLMParamsTimeout: S.Schema<UpdateLiteLLMParamsTimeout>;
877
+ export type UpdateLiteLLMParamsVertexCredentialsCase1Map = {
878
+ [key: string]: unknown | undefined;
879
+ };
880
+ export declare const UpdateLiteLLMParamsVertexCredentialsCase1Map: S.Schema<UpdateLiteLLMParamsVertexCredentialsCase1Map>;
881
+ export type UpdateLiteLLMParamsVertexCredentials = string | UpdateLiteLLMParamsVertexCredentialsCase1Map;
882
+ export declare const UpdateLiteLLMParamsVertexCredentials: S.Schema<UpdateLiteLLMParamsVertexCredentials>;
883
+ export interface UpdateLiteLLMParams {
884
+ adaptive_router_config?: UpdateLiteLLMParamsAdaptiveRouterConfigMap | null;
885
+ adaptive_router_default_model?: string | null;
886
+ allow_client_keepalive_override?: boolean | null;
887
+ annotation_cost_per_page?: number | null;
888
+ api_base?: string | null;
889
+ api_key?: string | null;
890
+ api_version?: string | null;
891
+ auto_router_config?: string | null;
892
+ auto_router_config_path?: string | null;
893
+ auto_router_default_model?: string | null;
894
+ auto_router_embedding_model?: string | null;
895
+ auto_router_max_input_chars?: number | null;
896
+ aws_access_key_id?: string | null;
897
+ aws_batch_role_arn?: string | null;
898
+ aws_bedrock_project_id?: string | null;
899
+ aws_bedrock_runtime_endpoint?: string | null;
900
+ aws_external_id?: string | null;
901
+ aws_profile_name?: string | null;
902
+ aws_region_name?: string | null;
903
+ aws_role_name?: string | null;
904
+ aws_secret_access_key?: string | null;
905
+ aws_session_name?: string | null;
906
+ aws_session_token?: string | null;
907
+ aws_sts_endpoint?: string | null;
908
+ aws_web_identity_token?: string | null;
909
+ azure_ad_token?: string | null;
910
+ bedrock_tags?: UpdateLiteLLMParamsBedrockTagsList | null;
911
+ budget_duration?: string | null;
912
+ cache_creation_input_audio_token_cost?: number | null;
913
+ cache_creation_input_token_cost?: number | null;
914
+ cache_creation_input_token_cost_above_1hr?: number | null;
915
+ cache_creation_input_token_cost_above_200k_tokens?: number | null;
916
+ cache_creation_input_token_cost_above_272k_tokens?: number | null;
917
+ cache_creation_input_token_cost_above_272k_tokens_flex?: number | null;
918
+ cache_creation_input_token_cost_above_272k_tokens_priority?: number | null;
919
+ cache_creation_input_token_cost_flex?: number | null;
920
+ cache_creation_input_token_cost_priority?: number | null;
921
+ cache_creation_input_token_cost_ultrafast?: number | null;
922
+ cache_read_input_audio_token_cost?: number | null;
923
+ cache_read_input_token_cost?: number | null;
924
+ cache_read_input_token_cost_above_200k_tokens?: number | null;
925
+ cache_read_input_token_cost_above_200k_tokens_priority?: number | null;
926
+ cache_read_input_token_cost_above_272k_tokens?: number | null;
927
+ cache_read_input_token_cost_above_272k_tokens_flex?: number | null;
928
+ cache_read_input_token_cost_above_272k_tokens_priority?: number | null;
929
+ cache_read_input_token_cost_above_512k_tokens?: number | null;
930
+ cache_read_input_token_cost_flex?: number | null;
931
+ cache_read_input_token_cost_priority?: number | null;
932
+ cache_read_input_token_cost_ultrafast?: number | null;
933
+ citation_cost_per_token?: number | null;
934
+ complexity_router_config?: UpdateLiteLLMParamsComplexityRouterConfigMap | null;
935
+ complexity_router_default_model?: string | null;
936
+ configurable_clientside_auth_params?: UpdateLiteLLMParamsConfigurableClientsideAuthParamsList | null;
937
+ custom_llm_provider?: string | null;
938
+ default_api_key_rpm_limit?: number | null;
939
+ default_api_key_tpm_limit?: number | null;
940
+ gcs_bucket_name?: string | null;
941
+ google_maps_grounding_cost_per_query?: number | null;
942
+ input_cost_per_audio_per_second?: number | null;
943
+ input_cost_per_audio_per_second_above_128k_tokens?: number | null;
944
+ input_cost_per_audio_token?: number | null;
945
+ input_cost_per_character?: number | null;
946
+ input_cost_per_character_above_128k_tokens?: number | null;
947
+ input_cost_per_image?: number | null;
948
+ input_cost_per_image_above_128k_tokens?: number | null;
949
+ input_cost_per_image_token?: number | null;
950
+ input_cost_per_pixel?: number | null;
951
+ input_cost_per_query?: number | null;
952
+ input_cost_per_second?: number | null;
953
+ input_cost_per_token?: number | null;
954
+ input_cost_per_token_above_128k_tokens?: number | null;
955
+ input_cost_per_token_above_200k_tokens?: number | null;
956
+ input_cost_per_token_above_200k_tokens_priority?: number | null;
957
+ input_cost_per_token_above_272k_tokens?: number | null;
958
+ input_cost_per_token_above_272k_tokens_flex?: number | null;
959
+ input_cost_per_token_above_272k_tokens_priority?: number | null;
960
+ input_cost_per_token_above_512k_tokens?: number | null;
961
+ input_cost_per_token_batches?: number | null;
962
+ input_cost_per_token_cache_hit?: number | null;
963
+ input_cost_per_token_flex?: number | null;
964
+ input_cost_per_token_priority?: number | null;
965
+ input_cost_per_token_ultrafast?: number | null;
966
+ input_cost_per_video_per_second?: number | null;
967
+ input_cost_per_video_per_second_above_128k_tokens?: number | null;
968
+ input_cost_per_video_per_second_above_15s_interval?: number | null;
969
+ input_cost_per_video_per_second_above_8s_interval?: number | null;
970
+ input_cost_per_video_token?: number | null;
971
+ itpm?: number | null;
972
+ keepalive_seconds?: number | null;
973
+ litellm_credential_name?: string | null;
974
+ litellm_trace_id?: string | null;
975
+ max_budget?: number | null;
976
+ max_file_size_mb?: number | null;
977
+ max_retries?: number | null;
978
+ merge_reasoning_content_in_choices?: boolean | null;
979
+ milvus_db_name?: string | null;
980
+ milvus_partition_names?: UpdateLiteLLMParamsMilvusPartitionNamesList | null;
981
+ milvus_text_field?: string | null;
982
+ mock_response?: UpdateLiteLLMParamsMockResponse | null;
983
+ model?: string | null;
984
+ model_info?: UpdateLiteLLMParamsModelInfoMap | null;
985
+ ocr_cost_per_credit?: number | null;
986
+ ocr_cost_per_page?: number | null;
987
+ organization?: string | null;
988
+ otpm?: number | null;
989
+ output_cost_per_audio_per_second?: number | null;
990
+ output_cost_per_audio_token?: number | null;
991
+ output_cost_per_character?: number | null;
992
+ output_cost_per_character_above_128k_tokens?: number | null;
993
+ output_cost_per_image?: number | null;
994
+ output_cost_per_image_token?: number | null;
995
+ output_cost_per_pixel?: number | null;
996
+ output_cost_per_reasoning_token?: number | null;
997
+ output_cost_per_reasoning_token_flex?: number | null;
998
+ output_cost_per_reasoning_token_priority?: number | null;
999
+ output_cost_per_second?: number | null;
1000
+ output_cost_per_second_1080p?: number | null;
1001
+ output_cost_per_second_480p?: number | null;
1002
+ output_cost_per_second_4k?: number | null;
1003
+ output_cost_per_token?: number | null;
1004
+ output_cost_per_token_above_128k_tokens?: number | null;
1005
+ output_cost_per_token_above_200k_tokens?: number | null;
1006
+ output_cost_per_token_above_200k_tokens_priority?: number | null;
1007
+ output_cost_per_token_above_272k_tokens?: number | null;
1008
+ output_cost_per_token_above_272k_tokens_flex?: number | null;
1009
+ output_cost_per_token_above_272k_tokens_priority?: number | null;
1010
+ output_cost_per_token_above_512k_tokens?: number | null;
1011
+ output_cost_per_token_batches?: number | null;
1012
+ output_cost_per_token_flex?: number | null;
1013
+ output_cost_per_token_priority?: number | null;
1014
+ output_cost_per_token_ultrafast?: number | null;
1015
+ output_cost_per_video_per_second?: number | null;
1016
+ output_cost_per_video_token?: number | null;
1017
+ output_vector_size?: number | null;
1018
+ quality_router_config?: UpdateLiteLLMParamsQualityRouterConfigMap | null;
1019
+ quality_router_default_model?: string | null;
1020
+ region_name?: string | null;
1021
+ regional_endpoint_uplift_multiplier?: number | null;
1022
+ regional_processing_uplift_multiplier_eu?: number | null;
1023
+ regional_processing_uplift_multiplier_us?: number | null;
1024
+ rpm?: number | null;
1025
+ s3_bucket_name?: string | null;
1026
+ s3_encryption_key_id?: string | null;
1027
+ s3_output_bucket_name?: string | null;
1028
+ s3_region_name?: string | null;
1029
+ search_context_cost_per_query?: UpdateLiteLLMParamsSearchContextCostPerQueryMap | null;
1030
+ stream_timeout?: UpdateLiteLLMParamsStreamTimeout | null;
1031
+ tag_regex?: UpdateLiteLLMParamsTagRegexList | null;
1032
+ tags?: UpdateLiteLLMParamsTagsList | null;
1033
+ tiered_pricing?: UpdateLiteLLMParamsTieredPricingList | null;
1034
+ timeout?: UpdateLiteLLMParamsTimeout | null;
1035
+ tpm?: number | null;
1036
+ use_chat_completions_api?: boolean | null;
1037
+ use_in_pass_through?: boolean | null;
1038
+ use_litellm_proxy?: boolean | null;
1039
+ /** Use stored xAI OAuth credentials when no xAI API key is configured. */
1040
+ use_xai_oauth?: boolean | null;
1041
+ valkey_embedding_field?: string | null;
1042
+ valkey_host?: string | null;
1043
+ valkey_password?: string | null;
1044
+ valkey_port?: number | null;
1045
+ valkey_ssl?: boolean | null;
1046
+ valkey_text_field?: string | null;
1047
+ vector_store_id?: string | null;
1048
+ vertex_credentials?: UpdateLiteLLMParamsVertexCredentials | null;
1049
+ vertex_location?: string | null;
1050
+ vertex_project?: string | null;
1051
+ watsonx_region_name?: string | null;
1052
+ }
1053
+ export declare const UpdateLiteLLMParams: S.Schema<UpdateLiteLLMParams>;
1054
+ export interface PatchModelModelModelIdUpdatePatchRequest {
1055
+ model_id: string;
1056
+ blocked?: boolean | null;
1057
+ litellm_params?: UpdateLiteLLMParams | null;
1058
+ model_info?: LitellmTypesRouterModelInfo | null;
1059
+ model_name?: string | null;
1060
+ }
1061
+ export declare const PatchModelModelModelIdUpdatePatchRequest: S.Schema<PatchModelModelModelIdUpdatePatchRequest>;
1062
+ export interface PatchModelModelModelIdUpdatePatchResponse {
1063
+ body: unknown;
1064
+ }
1065
+ export declare const PatchModelModelModelIdUpdatePatchResponse: S.Schema<PatchModelModelModelIdUpdatePatchResponse>;
1066
+ export interface PostBlockModelModelBlockRequest {
1067
+ model_id: string;
1068
+ }
1069
+ export declare const PostBlockModelModelBlockRequest: S.Schema<PostBlockModelModelBlockRequest>;
1070
+ export type LiteLLMProxyModelTableLitellmParamsMap = {
1071
+ [key: string]: unknown | undefined;
1072
+ };
1073
+ export declare const LiteLLMProxyModelTableLitellmParamsMap: S.Schema<LiteLLMProxyModelTableLitellmParamsMap>;
1074
+ export type LiteLLMProxyModelTableModelInfoMap = {
1075
+ [key: string]: unknown | undefined;
1076
+ };
1077
+ export declare const LiteLLMProxyModelTableModelInfoMap: S.Schema<LiteLLMProxyModelTableModelInfoMap>;
1078
+ export interface LiteLLMProxyModelTable {
1079
+ blocked?: boolean;
1080
+ created_at?: string | null;
1081
+ created_by?: string | null;
1082
+ litellm_params: LiteLLMProxyModelTableLitellmParamsMap;
1083
+ model_id: string;
1084
+ model_info?: LiteLLMProxyModelTableModelInfoMap | null;
1085
+ model_name: string;
1086
+ updated_at?: string | null;
1087
+ updated_by?: string | null;
1088
+ }
1089
+ export declare const LiteLLMProxyModelTable: S.Schema<LiteLLMProxyModelTable>;
1090
+ export interface PostBlockModelModelBlockResponse {
1091
+ body: LiteLLMProxyModelTable | null;
1092
+ }
1093
+ export declare const PostBlockModelModelBlockResponse: S.Schema<PostBlockModelModelBlockResponse>;
1094
+ /** An operator-defined tier: the name the LLM classifier must return and its rubric description. */
1095
+ export interface TierDefinition {
1096
+ /** What belongs in this tier; rendered as this tier's bullet in the classifier rubric. Required unless the name is a built-in tier (SIMPLE/MEDIUM/COMPLEX/REASONING), which inherits the built-in criteria when omitted */
1097
+ description?: string | null;
1098
+ /** Tier name; becomes a value the LLM classifier can return and a key of `tiers` */
1099
+ name: string;
1100
+ }
1101
+ export declare const TierDefinition: S.Schema<TierDefinition>;
1102
+ export type PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList = Array<TierDefinition>;
1103
+ export declare const PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList: S.Schema<PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList>;
1104
+ export interface PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest {
1105
+ classification_prompt?: string | null;
1106
+ context_window_size?: number;
1107
+ tier_definitions: PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequestTierDefinitionsList;
1108
+ }
1109
+ export declare const PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest: S.Schema<PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest>;
1110
+ /** When adaptive=True: 'all' scores every pool model with a tier-distance penalty (soft floors); 'classified_tier' Thompson-samples only inside the classified tier's pool */
1111
+ export type RequestComplexityRouterConfigAdaptiveEligible = "all" | "classified_tier";
1112
+ export declare const RequestComplexityRouterConfigAdaptiveEligible: any;
1113
+ export interface AdaptiveRouterWeights {
1114
+ cost?: number;
1115
+ quality?: number;
1116
+ }
1117
+ export declare const AdaptiveRouterWeights: S.Schema<AdaptiveRouterWeights>;
1118
+ /** What classifies the request when the LLM classifier errors, times out, or returns an unparseable response. 'heuristic' runs the local complexity scorer, which is right when the classifier grades complexity too. 'default_model' skips scoring and routes to default_model, which is what a classifier on some other taxonomy wants: a prompt that grades data sensitivity has no use for a complexity score, and scoring one produces a tier unrelated to what the operator configured. Requires default_model when set to 'default_model'. Only applies when classifier_type is 'llm', 'custom', or 'heuristic_first'. */
1119
+ export type RequestComplexityRouterConfigClassifierFallback = "heuristic" | "default_model";
1120
+ export declare const RequestComplexityRouterConfigClassifierFallback: any;
1121
+ /** Configuration for the LLM-based complexity classifier. */
1122
+ export interface ClassifierLLMConfig {
1123
+ /** Which calibration examples the built-in rubric carries. 'agentic' anchors routine installs, builds, multi-file edits, and standard debugging at MEDIUM, so ordinary engineering does not route to the most expensive tier; it suits agent, terminal, and coding-assistant traffic as well as mixed traffic. 'chat' omits those engineering anchors, for a deployment serving only conversational traffic. 'business' carries business/sales anchors and business-flavored tier criteria that keep routine drafting and summarizing off the expensive tiers and reserve the top tier for committing to decisions under tradeoffs; it suits sales, support, and go-to-market traffic. Every preset keeps the same four tiers, so this moves where the boundary sits without changing the taxonomy. Leave unset for 'legacy', the rubric as it shipped before calibration examples existed, so an existing router's tier decisions and spend do not move on upgrade. Mutually exclusive with system_prompt, which replaces the rubric this would select. Only applies when classifier_type is 'llm'. */
1124
+ classification_rubric?: ClassificationRubric | (string & {}) | null;
1125
+ /** Model name (from the router's model_list) to call for classification */
1126
+ model: string;
1127
+ /** Replaces the built-in complexity rubric as the classifier's entire system role. When set, neither the default rubric nor the context-window closing line is appended, so the prompt owns the whole taxonomy and the tier names SIMPLE/MEDIUM/COMPLEX/REASONING become whatever buckets it defines: a prompt that classifies data sensitivity routes on that instead of on difficulty. Two consequences of full replacement. The default rubric's closing paragraph is the classifier's prompt-injection defense, telling it that the caller's quoted system prompt and prior turns are material to judge and never instructions; a replacement that omits it lets a caller ask for a tier and get it. And the heuristic fallback still scores complexity, so a router on some other taxonomy wants classifier_fallback='default_model'. Leave unset for the built-in rubric. Only applies when classifier_type is 'llm'. */
1128
+ system_prompt?: string | null;
1129
+ /** Timeout budget for the classification call, in milliseconds */
1130
+ timeout_ms?: number;
1131
+ }
1132
+ export declare const ClassifierLLMConfig: S.Schema<ClassifierLLMConfig>;
1133
+ /** Classification strategy: local regex/keyword scoring, an LLM call, a custom classifier plugin, or 'heuristic_first', which scores locally and only pays for the LLM classifier when the local scorer does not confidently land a cheap tier */
1134
+ export type RequestComplexityRouterConfigClassifierType = "heuristic" | "llm" | "custom" | "heuristic_first";
1135
+ export declare const RequestComplexityRouterConfigClassifierType: any;
1136
+ export type RequestComplexityRouterConfigCodeKeywordsList = Array<string>;
1137
+ export declare const RequestComplexityRouterConfigCodeKeywordsList: S.Schema<RequestComplexityRouterConfigCodeKeywordsList>;
1138
+ export type RequestComplexityRouterConfigCustomTechnicalKeywordsList = Array<string>;
1139
+ export declare const RequestComplexityRouterConfigCustomTechnicalKeywordsList: S.Schema<RequestComplexityRouterConfigCustomTechnicalKeywordsList>;
1140
+ /** Weights for each scoring dimension */
1141
+ export type RequestComplexityRouterConfigDimensionWeightsMap = {
1142
+ [key: string]: number | undefined;
1143
+ };
1144
+ export declare const RequestComplexityRouterConfigDimensionWeightsMap: S.Schema<RequestComplexityRouterConfigDimensionWeightsMap>;
1145
+ export type RequestComplexityRouterConfigEscalationKeywordsList = Array<string>;
1146
+ export declare const RequestComplexityRouterConfigEscalationKeywordsList: S.Schema<RequestComplexityRouterConfigEscalationKeywordsList>;
1147
+ export type RequestComplexityRouterConfigHousekeepingPatternsList = Array<string>;
1148
+ export declare const RequestComplexityRouterConfigHousekeepingPatternsList: S.Schema<RequestComplexityRouterConfigHousekeepingPatternsList>;
1149
+ /** Keywords/phrases that trigger this rule (lexical or semantic match) */
1150
+ export type KeywordTierRuleKeywordsList = Array<string>;
1151
+ export declare const KeywordTierRuleKeywordsList: S.Schema<KeywordTierRuleKeywordsList>;
1152
+ /** A deterministic override: if any keyword matches, route to this tier. */
1153
+ export interface KeywordTierRule {
1154
+ /** Keywords/phrases that trigger this rule (lexical or semantic match) */
1155
+ keywords: KeywordTierRuleKeywordsList;
1156
+ /** Tier to route to when this rule matches: a built-in tier name, or with tier_definitions set, one of the defined tier names */
1157
+ tier: string;
1158
+ }
1159
+ export declare const KeywordTierRule: S.Schema<KeywordTierRule>;
1160
+ export type RequestComplexityRouterConfigKeywordTierRulesList = Array<KeywordTierRule>;
1161
+ export declare const RequestComplexityRouterConfigKeywordTierRulesList: S.Schema<RequestComplexityRouterConfigKeywordTierRulesList>;
1162
+ export type RequestComplexityRouterConfigPlanModePatternsList = Array<string>;
1163
+ export declare const RequestComplexityRouterConfigPlanModePatternsList: S.Schema<RequestComplexityRouterConfigPlanModePatternsList>;
1164
+ export type RequestComplexityRouterConfigReasoningKeywordsList = Array<string>;
1165
+ export declare const RequestComplexityRouterConfigReasoningKeywordsList: S.Schema<RequestComplexityRouterConfigReasoningKeywordsList>;
1166
+ /** One open/close delimiter pair a harness wraps injected context in. Normalizing here rather than at the scan is what makes matching case-insensitive: markers reach the scan already lowered, so it lowercases only the haystack and never the needles. Stripping keeps YAML indentation whitespace from becoming part of the delimiter. */
1167
+ export interface ReminderMarkerPair {
1168
+ /** Closing delimiter, e.g. '</system-reminder>' */
1169
+ close: string;
1170
+ /** Opening delimiter, e.g. '<system-reminder>' */
1171
+ open: string;
1172
+ }
1173
+ export declare const ReminderMarkerPair: S.Schema<ReminderMarkerPair>;
1174
+ export type RequestComplexityRouterConfigReminderMarkersList = Array<ReminderMarkerPair>;
1175
+ export declare const RequestComplexityRouterConfigReminderMarkersList: S.Schema<RequestComplexityRouterConfigReminderMarkersList>;
1176
+ export type RequestComplexityRouterConfigSimpleKeywordsList = Array<string>;
1177
+ export declare const RequestComplexityRouterConfigSimpleKeywordsList: S.Schema<RequestComplexityRouterConfigSimpleKeywordsList>;
1178
+ export type RequestComplexityRouterConfigTechnicalKeywordsList = Array<string>;
1179
+ export declare const RequestComplexityRouterConfigTechnicalKeywordsList: S.Schema<RequestComplexityRouterConfigTechnicalKeywordsList>;
1180
+ /** Score boundaries between tiers. These keys (simple_medium, medium_complex, complex_reasoning) name the gaps between the default tier names and are not renameable by tier_labels; they are scorer knobs persisted by name on every routing decision */
1181
+ export type RequestComplexityRouterConfigTierBoundariesMap = {
1182
+ [key: string]: number | undefined;
1183
+ };
1184
+ export declare const RequestComplexityRouterConfigTierBoundariesMap: S.Schema<RequestComplexityRouterConfigTierBoundariesMap>;
1185
+ export type RequestComplexityRouterConfigTierDefinitionsList = Array<TierDefinition>;
1186
+ export declare const RequestComplexityRouterConfigTierDefinitionsList: S.Schema<RequestComplexityRouterConfigTierDefinitionsList>;
1187
+ /** Display names for the complexity tiers, so a deployment can use its own vocabulary (e.g. Cheap/Standard/Premium/Deep) in the dashboard, spend logs, and the LLM classifier rubric. Purely operator-facing: config keys stay canonical (tiers, keyword_tier_rules[].tier, tier_boundaries), API callers never see these names, and the heuristic scorer never reads them. Unlisted tiers keep their canonical name. Partial maps are allowed. */
1188
+ export type RequestComplexityRouterConfigTierLabelsMap = {
1189
+ [key: string]: string | undefined;
1190
+ };
1191
+ export declare const RequestComplexityRouterConfigTierLabelsMap: S.Schema<RequestComplexityRouterConfigTierLabelsMap>;
1192
+ export type ComplexityTierModelLitellmParamsMap = {
1193
+ [key: string]: unknown | undefined;
1194
+ };
1195
+ export declare const ComplexityTierModelLitellmParamsMap: S.Schema<ComplexityTierModelLitellmParamsMap>;
1196
+ export interface ComplexityTierModel {
1197
+ litellm_params?: ComplexityTierModelLitellmParamsMap;
1198
+ model_name: string;
1199
+ }
1200
+ export declare const ComplexityTierModel: S.Schema<ComplexityTierModel>;
1201
+ export type RequestComplexityRouterConfigTierModelConfigsValueList = Array<ComplexityTierModel>;
1202
+ export declare const RequestComplexityRouterConfigTierModelConfigsValueList: S.Schema<RequestComplexityRouterConfigTierModelConfigsValueList>;
1203
+ export type RequestComplexityRouterConfigTierModelConfigsMap = {
1204
+ [key: string]: RequestComplexityRouterConfigTierModelConfigsValueList | undefined;
1205
+ };
1206
+ export declare const RequestComplexityRouterConfigTierModelConfigsMap: S.Schema<RequestComplexityRouterConfigTierModelConfigsMap>;
1207
+ export type RequestComplexityRouterConfigTiersValueCase1List = Array<string>;
1208
+ export declare const RequestComplexityRouterConfigTiersValueCase1List: S.Schema<RequestComplexityRouterConfigTiersValueCase1List>;
1209
+ export type RequestComplexityRouterConfigTiersValue = string | RequestComplexityRouterConfigTiersValueCase1List;
1210
+ export declare const RequestComplexityRouterConfigTiersValue: S.Schema<RequestComplexityRouterConfigTiersValue>;
1211
+ /** Mapping of complexity tiers to a model or model pool. A list is randomly picked from when adaptive=False, and used as a soft-floor home pool when adaptive=True */
1212
+ export type RequestComplexityRouterConfigTiersMap = {
1213
+ [key: string]: RequestComplexityRouterConfigTiersValue | undefined;
1214
+ };
1215
+ export declare const RequestComplexityRouterConfigTiersMap: S.Schema<RequestComplexityRouterConfigTiersMap>;
1216
+ /** Token count thresholds for simple/complex classification */
1217
+ export type RequestComplexityRouterConfigTokenThresholdsMap = {
1218
+ [key: string]: number | undefined;
1219
+ };
1220
+ export declare const RequestComplexityRouterConfigTokenThresholdsMap: S.Schema<RequestComplexityRouterConfigTokenThresholdsMap>;
1221
+ /** The part of a complexity-router config a request can carry. `plugins` holds live RoutingPlugin objects, which no JSON body can express and which have no OpenAPI schema, so it is closed off here rather than left as an arbitrary-type field. */
1222
+ export interface RequestComplexityRouterConfig {
1223
+ /** Enable adaptive bandit selection with soft complexity floors */
1224
+ adaptive?: boolean;
1225
+ /** When adaptive=True: 'all' scores every pool model with a tier-distance penalty (soft floors); 'classified_tier' Thompson-samples only inside the classified tier's pool */
1226
+ adaptive_eligible?: RequestComplexityRouterConfigAdaptiveEligible | (string & {});
1227
+ /** Quality vs cost weights for adaptive selection (used when adaptive=True) */
1228
+ adaptive_weights?: AdaptiveRouterWeights;
1229
+ /** Replaces the opening instructions of the LLM classifier rubric (the judging-criteria prose) for a custom tier set. The per-tier bullets and the trust-boundary paragraph telling the classifier to ignore tier requests embedded in quoted caller text are always appended after it and cannot be overridden. Requires tier_definitions; a built-in-tier router customizes its prompt via classifier_llm_config.system_prompt or classification_rubric instead. */
1230
+ classification_prompt?: string | null;
1231
+ /** Maximum characters of prior-turn text quoted to the LLM classifier, across the whole context window, per classification call. Turns are taken newest first and quoted whole while they fit, so a conversation small enough to quote entirely is never cut; once the budget runs out the older turns are dropped whole and only the turn straddling the boundary is truncated, into whatever space is left. The current ask and the caller's system prompt sit outside this budget and are always sent in full, as does the numbering each quoted turn carries. A budget under 120 leaves no room to quote a turn and suppresses the block; set classifier_context_window_size to 0 to turn context off deliberately. Only applies when classifier_type is 'llm'. */
1232
+ classifier_context_budget_chars?: number;
1233
+ /** Include assistant turns in the classifier context window, so difficulty stated by the model rather than by the user stays visible: a plan the assistant calls complex, which the user approves with 'yes', is classified on the work being approved instead of on the word 'yes'. When enabled, classifier_context_window_size counts the last N turns of the conversation across both roles rather than the last N user turns, and assistant text is sent to the classifier model, which may be a different deployment or provider than the routed completion model. Assistant replies spend classifier_context_budget_chars alongside user turns, so raise it if the oldest turns stop being quoted once replies join the window. Off by default because enabling it shifts tier decisions, and therefore spend, for an already-deployed router. Only applies when classifier_type is 'llm'. */
1234
+ classifier_context_include_assistant_turns?: boolean;
1235
+ /** Optional cap on each individual prior turn's text, applied before classifier_context_budget_chars bounds the block. Unset by default, so one long turn may spend the whole budget, which is usually what a follow-up needs; set it when no single turn should dominate the context the classifier sees. A capped turn keeps its opening and its ending with the middle elided. Only applies when classifier_type is 'llm'. */
1236
+ classifier_context_per_turn_chars?: number | null;
1237
+ /** Number of prior user turns (tool output and harness reminders excluded) to include as context in the LLM classifier prompt, so a follow-up like 'now do the same for the streaming path' is classified against what it refers to. Counts turns of both roles when classifier_context_include_assistant_turns is enabled. These turns are sent to the classifier model, which may be a different deployment or provider than the routed completion model; that call already carries the current user ask and the caller's system prompt in full. Set to 0 to send neither prior turns nor any conversation context beyond the current ask. Only applies when classifier_type is 'llm'. */
1238
+ classifier_context_window_size?: number;
1239
+ /** What classifies the request when the LLM classifier errors, times out, or returns an unparseable response. 'heuristic' runs the local complexity scorer, which is right when the classifier grades complexity too. 'default_model' skips scoring and routes to default_model, which is what a classifier on some other taxonomy wants: a prompt that grades data sensitivity has no use for a complexity score, and scoring one produces a tier unrelated to what the operator configured. Requires default_model when set to 'default_model'. Only applies when classifier_type is 'llm', 'custom', or 'heuristic_first'. */
1240
+ classifier_fallback?: RequestComplexityRouterConfigClassifierFallback | (string & {});
1241
+ /** Configuration for the LLM classifier; required when classifier_type is 'llm' or 'heuristic_first' */
1242
+ classifier_llm_config?: ClassifierLLMConfig | null;
1243
+ /** Not settable over HTTP; the classifier plugin is a runtime object */
1244
+ classifier_plugin?: unknown | null;
1245
+ /** Timeout budget for the classifier plugin call, in milliseconds. On expiry the fallback path decides the tier. Only applies when classifier_type is 'custom'. */
1246
+ classifier_plugin_timeout_ms?: number;
1247
+ /** Classification strategy: local regex/keyword scoring, an LLM call, a custom classifier plugin, or 'heuristic_first', which scores locally and only pays for the LLM classifier when the local scorer does not confidently land a cheap tier */
1248
+ classifier_type?: RequestComplexityRouterConfigClassifierType | (string & {});
1249
+ /** Keywords indicating code-related content */
1250
+ code_keywords?: RequestComplexityRouterConfigCodeKeywordsList | null;
1251
+ /** Domain-specific technical keywords appended to the effective base list (technical_keywords if set, otherwise DEFAULT_TECHNICAL_KEYWORDS). Order is preserved; duplicates are removed case-insensitively against the base list and within this list. */
1252
+ custom_technical_keywords?: RequestComplexityRouterConfigCustomTechnicalKeywordsList | null;
1253
+ /** Default model to use if tier cannot be determined */
1254
+ default_model?: string | null;
1255
+ /** When True and a session_id is resolvable on the request, pin the deployment chosen inside each routed model group and reuse it whenever the session returns to that group, without pinning which group the session routes to. Independent of session_affinity, which pins the model group instead (and always carries this deployment pin with it): with session_affinity off, every turn is still classified on its own merits while a session that escalates to a stronger tier and comes back still lands on the deployment it used before, which is what keeps a provider prompt cache warm. Pins are held per model group, so switching tiers does not disturb the pin left behind in the previous group. On by default because re-shuffling a conversation across deployments of the same model discards that cache for no benefit; set False to keep every turn load-balanced across the group, which is what a deployment set with tight per-deployment rate limits wants. Inert when no session_id is resolvable, since there is nothing to key a pin on, and suppressed when plugins are configured, for the same reason session_affinity is. */
1256
+ deployment_affinity?: boolean;
1257
+ /** Weights for each scoring dimension */
1258
+ dimension_weights?: RequestComplexityRouterConfigDimensionWeightsMap;
1259
+ /** Embedding model (LiteLLM model name) used when semantic_keyword_matching is enabled */
1260
+ embedding_model?: string | null;
1261
+ /** Case-sensitive phrases a user can include to force a bump to the next-higher complexity tier when they aren't satisfied with results (they can force a stronger model, but not choose which one). Defaults to ['LITELLM ESCALATE'] when unset; set to an empty list to disable. */
1262
+ escalation_keywords?: RequestComplexityRouterConfigEscalationKeywordsList | null;
1263
+ /** Tier routed to when the LLM classifier fails (timeout, provider error, or an unparseable reply). Required with tier_definitions and must name a defined tier; the heuristic scorer cannot produce custom tiers, so this replaces the heuristic fallback for custom tier sets. */
1264
+ fallback_tier?: string | null;
1265
+ /** The highest tier the local scorer may decide on its own; required when classifier_type is 'heuristic_first' and rejected otherwise. A request whose heuristic tier is at or below this one skips the LLM classifier and routes straight to that heuristic tier, so the classifier call is only paid for on traffic the scorer could not place cheaply. The scorer must also have produced at least one signal: a prompt where no dimension fired scores 0.0 and would otherwise land SIMPLE by default rather than by evidence, which is how a chained router would silently send unclassified traffic to the cheapest model. Names a built-in tier, and may not name the highest one, since that would make the LLM classifier unreachable. */
1266
+ heuristic_first_max_tier?: string | null;
1267
+ /** Additional case-sensitive literal sentinels that mark a request as client housekeeping, on top of the built-in conversation-title ones. For clients whose wording the built-ins don't cover, or after a client release changes its strings. */
1268
+ housekeeping_patterns?: RequestComplexityRouterConfigHousekeepingPatternsList | null;
1269
+ /** Rules that force a specific tier when their keywords match the prompt */
1270
+ keyword_tier_rules?: RequestComplexityRouterConfigKeywordTierRulesList | null;
1271
+ /** Minimum cosine similarity for a semantic keyword match */
1272
+ match_threshold?: number;
1273
+ /** When set, requests carrying a coding-agent plan-mode sentinel (Claude Code plan mode, VS Code Copilot Plan mode, Copilot CLI's exit_plan_mode tool) are routed to at least this tier: the classified tier still wins when it is higher, and the floor also overrides a session-affinity pin to a lower tier for exactly the turns carrying the sentinel, without rewriting the pin -- the first turn after plan mode exits routes as if plan mode had never happened. Names a built-in tier, or with tier_definitions set, one of the defined tier names (list order is ascending severity, same as keyword_tier_rules). Unset disables detection entirely. The sentinels ride in client-injected prompt text, so a caller who pastes one can spend up to this tier's models -- never down, and never outside the configured pools. */
1274
+ plan_mode_min_tier?: string | null;
1275
+ /** Additional case-sensitive literal sentinels that mark a request as plan mode, on top of the built-in Claude Code and Copilot ones. For clients whose plan-mode wording the built-ins don't cover, or after a client release changes its strings. */
1276
+ plan_mode_patterns?: RequestComplexityRouterConfigPlanModePatternsList | null;
1277
+ /** Not settable over HTTP; routing plugins are runtime objects */
1278
+ plugins?: unknown | null;
1279
+ /** Keywords indicating reasoning-required content */
1280
+ reasoning_keywords?: RequestComplexityRouterConfigReasoningKeywordsList | null;
1281
+ /** Minimum weighted score a request must reach before 2+ reasoning markers may promote it to the reasoning tier. Unset tracks tier_boundaries.simple_medium, so the override never rescues a request the scorer placed in the cheapest tier; 0 restores the unconditional override */
1282
+ reasoning_override_min_score?: number | null;
1283
+ /** Override the delimiter pairs used to recognize and strip harness-injected reminder blocks before classification. A harness that wraps injected context differently per agent type (main, subagent, cron) lists every pair it emits. Replaces, rather than adds to, the built-in default of ('<system-reminder>', '</system-reminder>'), so a harness that also emits that pair lists it too. Matching is case-insensitive. */
1284
+ reminder_markers?: RequestComplexityRouterConfigReminderMarkersList | null;
1285
+ /** Return the resolved raw model name in the response model field instead of the client-requested complexity-router alias */
1286
+ return_raw_model_name?: boolean;
1287
+ /** Route a coding agent's own housekeeping calls to the cheapest configured tier without classifying them. A client names the conversation by quoting the whole session and asking for a title, so the ask reads as the session's engineering work and lands on the most expensive tier, which is the reverse of what the call is worth. Detection is a literal match against client-owned sentinels on the newest ask only, so it cannot fire on an earlier turn, and it never lowers what anyone else asked for: a keyword_tier_rule or a session pin still decides instead, and an escalation keyword or the plan-mode floor still raises the tier from here. Only the classifier is displaced, and its call is skipped, so a matched request costs nothing to route. Set false to classify these calls like any other. */
1288
+ route_housekeeping_to_cheapest_tier?: boolean;
1289
+ /** Match keyword_tier_rules by embedding similarity instead of literal text */
1290
+ semantic_keyword_matching?: boolean;
1291
+ /** When True and a session_id is resolvable on the request, pin the model chosen on the session's first turn and reuse it for every later turn, skipping re-classification. Off by default so every turn is classified on its own merits and routed to the cheapest adequate tier. Set True to keep a multi-turn session on one model, which preserves provider prompt caches and avoids cross-model conversation-history errors. Always implies the deployment pin regardless of deployment_affinity: the session sticks to one deployment of the pinned model, since freezing the model while re-shuffling its deployments would still go cache-cold. */
1292
+ session_affinity?: boolean;
1293
+ /** TTL for the session affinity pin; refreshed on every cache hit. Bounds both the session_affinity model pin and the deployment_affinity deployment pin, so it measures idle time for the session's routing decisions rather than total session length */
1294
+ session_affinity_ttl_seconds?: number;
1295
+ /** Keywords indicating simple/basic queries */
1296
+ simple_keywords?: RequestComplexityRouterConfigSimpleKeywordsList | null;
1297
+ /** Keywords indicating technical content */
1298
+ technical_keywords?: RequestComplexityRouterConfigTechnicalKeywordsList | null;
1299
+ /** Score boundaries between tiers. These keys (simple_medium, medium_complex, complex_reasoning) name the gaps between the default tier names and are not renameable by tier_labels; they are scorer knobs persisted by name on every routing decision */
1300
+ tier_boundaries?: RequestComplexityRouterConfigTierBoundariesMap;
1301
+ /** Operator-defined tier set replacing the built-in SIMPLE/MEDIUM/COMPLEX/REASONING. Each entry's name becomes a value the LLM classifier can return and its description becomes that tier's rubric bullet; entries named after a built-in tier may omit the description and inherit the built-in criteria. List order is ascending severity and decides which tier wins when several keyword_tier_rules match. Requires classifier_type 'llm' or 'custom', a fallback_tier, and `tiers` keys matching the defined names exactly. Escalation, adaptive selection, session affinity, plugins, tier_labels, and the calibration-example rubric presets are unavailable with a custom tier set: the first four are built on the built-in tier ladder, and the last two rename or exemplify tiers the set replaces. */
1302
+ tier_definitions?: RequestComplexityRouterConfigTierDefinitionsList | null;
1303
+ /** Score penalty per tier-step away from the classified tier when adaptive=True */
1304
+ tier_distance_penalty?: number;
1305
+ /** Display names for the complexity tiers, so a deployment can use its own vocabulary (e.g. Cheap/Standard/Premium/Deep) in the dashboard, spend logs, and the LLM classifier rubric. Purely operator-facing: config keys stay canonical (tiers, keyword_tier_rules[].tier, tier_boundaries), API callers never see these names, and the heuristic scorer never reads them. Unlisted tiers keep their canonical name. Partial maps are allowed. */
1306
+ tier_labels?: RequestComplexityRouterConfigTierLabelsMap;
1307
+ tier_model_configs?: RequestComplexityRouterConfigTierModelConfigsMap;
1308
+ /** Mapping of complexity tiers to a model or model pool. A list is randomly picked from when adaptive=False, and used as a soft-floor home pool when adaptive=True */
1309
+ tiers?: RequestComplexityRouterConfigTiersMap;
1310
+ /** Token count thresholds for simple/complex classification */
1311
+ token_thresholds?: RequestComplexityRouterConfigTokenThresholdsMap;
1312
+ }
1313
+ export declare const RequestComplexityRouterConfig: S.Schema<RequestComplexityRouterConfig>;
1314
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap = {
1315
+ [key: string]: unknown | undefined;
1316
+ };
1317
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap>;
1318
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList = Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesItemMap>;
1319
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList>;
1320
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap = {
1321
+ [key: string]: unknown | undefined;
1322
+ };
1323
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap>;
1324
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List = Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1ItemMap>;
1325
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List>;
1326
+ /** The top-level system prompt an Anthropic /v1/messages body carries beside its messages */
1327
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem = string | PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystemCase1List;
1328
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem>;
1329
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap = {
1330
+ [key: string]: unknown | undefined;
1331
+ };
1332
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap>;
1333
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList = Array<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsItemMap>;
1334
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList>;
1335
+ export interface PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest {
1336
+ /** The complexity router config to route against, in the shape /model/new accepts */
1337
+ complexity_router_config: RequestComplexityRouterConfig;
1338
+ /** Model to route to when no tier resolves, i.e. complexity_router_default_model */
1339
+ default_model?: string | null;
1340
+ /** The full message list to route, exactly as the serving path would receive it. Mutually exclusive with prompt */
1341
+ messages?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestMessagesList | null;
1342
+ /** A single ask to route, as an end user would send it. Mutually exclusive with messages */
1343
+ prompt?: string | null;
1344
+ /** Name reported as the router in the routing decision. Display only */
1345
+ router_name?: string;
1346
+ /** The top-level system prompt an Anthropic /v1/messages body carries beside its messages */
1347
+ system?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestSystem | null;
1348
+ /** Team the router is being created for. Required for a team admin, who may only test their own team's routers */
1349
+ team_id?: string | null;
1350
+ /** The tool definitions the request advertises, which decide whether the plan-mode floor applies */
1351
+ tools?: PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequestToolsList | null;
1352
+ }
1353
+ export declare const PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest: S.Schema<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest>;
1354
+ export type StandardLoggingRoutingDecisionCause = "heuristic_scorer" | "reasoning_override" | "llm_classifier" | "heuristic_first_short_circuit" | "classifier_plugin" | "classifier_fallback" | "default_model_fallback" | "literal_keyword_match" | "semantic_keyword_match" | "plan_mode" | "housekeeping" | "session_affinity_pin" | "session_affinity_escalation" | "default_fallback" | "keyword" | "quality_tier" | "bandit";
1355
+ export declare const StandardLoggingRoutingDecisionCause: any;
1356
+ export type StandardLoggingRoutingDecisionRouterType = "complexity" | "adaptive" | "quality";
1357
+ export declare const StandardLoggingRoutingDecisionRouterType: any;
1358
+ export type StandardLoggingRoutingDecisionSignalsList = Array<string>;
1359
+ export declare const StandardLoggingRoutingDecisionSignalsList: S.Schema<StandardLoggingRoutingDecisionSignalsList>;
1360
+ /** Snapshot of the complexity scorer's tier boundaries at decision time, so a historical spend log row stays explainable after the router config changes. */
1361
+ export interface StandardLoggingRoutingDecisionTierBoundaries {
1362
+ complex_reasoning: number;
1363
+ medium_complex: number;
1364
+ simple_medium: number;
1365
+ }
1366
+ export declare const StandardLoggingRoutingDecisionTierBoundaries: S.Schema<StandardLoggingRoutingDecisionTierBoundaries>;
1367
+ export type StandardLoggingRoutingDecisionTierLitellmParamsMap = {
1368
+ [key: string]: unknown | undefined;
1369
+ };
1370
+ export declare const StandardLoggingRoutingDecisionTierLitellmParamsMap: S.Schema<StandardLoggingRoutingDecisionTierLitellmParamsMap>;
1371
+ /** Per-request provenance for a pre-routing strategy (auto-router) decision. */
1372
+ export interface StandardLoggingRoutingDecision {
1373
+ cause?: StandardLoggingRoutingDecisionCause;
1374
+ classifier_cost?: number;
1375
+ classifier_model?: string;
1376
+ conversation_continuing?: boolean;
1377
+ escalated?: boolean;
1378
+ escalation_keyword?: string;
1379
+ matched_keyword?: string;
1380
+ reasoning_override_min_score?: number;
1381
+ request_type?: string;
1382
+ routed_model?: string;
1383
+ router_model_name?: string;
1384
+ router_type?: StandardLoggingRoutingDecisionRouterType;
1385
+ savings_baseline_deployment_id?: string;
1386
+ savings_baseline_model?: string;
1387
+ score?: number;
1388
+ signals?: StandardLoggingRoutingDecisionSignalsList;
1389
+ tier?: string;
1390
+ tier_boundaries?: StandardLoggingRoutingDecisionTierBoundaries;
1391
+ tier_label?: string;
1392
+ tier_litellm_params?: StandardLoggingRoutingDecisionTierLitellmParamsMap;
1393
+ }
1394
+ export declare const StandardLoggingRoutingDecision: S.Schema<StandardLoggingRoutingDecision>;
1395
+ /** Where one prompt would have been routed, and why. */
1396
+ export interface AutoRouterRoutingTestResponse {
1397
+ /** The model group the router picked */
1398
+ routed_model: string;
1399
+ /** Whether routed_model is a model group available to the caller, scoped to team_id when given. Never confirms models the caller could not use */
1400
+ routed_model_configured: boolean;
1401
+ /** The decision record this request would have written to its log row */
1402
+ routing_decision: StandardLoggingRoutingDecision;
1403
+ }
1404
+ export declare const AutoRouterRoutingTestResponse: S.Schema<AutoRouterRoutingTestResponse>;
1405
+ export interface PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest {
1406
+ }
1407
+ export declare const PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest: S.Schema<PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest>;
1408
+ export interface PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse {
1409
+ body: unknown;
1410
+ }
1411
+ export declare const PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse: S.Schema<PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse>;
1412
+ export interface PostReloadModelCostMapReloadModelCostMapRequest {
1413
+ }
1414
+ export declare const PostReloadModelCostMapReloadModelCostMapRequest: S.Schema<PostReloadModelCostMapReloadModelCostMapRequest>;
1415
+ export interface PostReloadModelCostMapReloadModelCostMapResponse {
1416
+ body: unknown;
1417
+ }
1418
+ export declare const PostReloadModelCostMapReloadModelCostMapResponse: S.Schema<PostReloadModelCostMapReloadModelCostMapResponse>;
1419
+ export interface PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest {
1420
+ hours: number;
1421
+ }
1422
+ export declare const PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest: S.Schema<PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest>;
1423
+ export interface PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse {
1424
+ body: unknown;
1425
+ }
1426
+ export declare const PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse: S.Schema<PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse>;
1427
+ export interface PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest {
1428
+ hours: number;
1429
+ }
1430
+ export declare const PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest: S.Schema<PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest>;
1431
+ export interface PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse {
1432
+ body: unknown;
1433
+ }
1434
+ export declare const PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse: S.Schema<PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse>;
1435
+ export interface PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest {
1436
+ access_group: string;
1437
+ budget_duration?: string | null;
1438
+ budget_id?: string | null;
1439
+ max_budget?: number | null;
1440
+ soft_budget?: number | null;
1441
+ }
1442
+ export declare const PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest: S.Schema<PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest>;
1443
+ export interface UnblockModelModelUnblockPostRequest {
1444
+ model_id: string;
1445
+ }
1446
+ export declare const UnblockModelModelUnblockPostRequest: S.Schema<UnblockModelModelUnblockPostRequest>;
1447
+ export interface UnblockModelModelUnblockPostResponse {
1448
+ body: LiteLLMProxyModelTable | null;
1449
+ }
1450
+ export declare const UnblockModelModelUnblockPostResponse: S.Schema<UnblockModelModelUnblockPostResponse>;
1451
+ export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList = Array<string>;
1452
+ export declare const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList: S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList>;
1453
+ export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList = Array<string>;
1454
+ export declare const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList: S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList>;
1455
+ export interface UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest {
1456
+ access_group: string;
1457
+ model_ids?: UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelIdsList | null;
1458
+ model_names?: UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequestModelNamesList | null;
1459
+ }
1460
+ export declare const UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest: S.Schema<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest>;
1461
+ export interface UpdateModelModelUpdatePostRequest {
1462
+ blocked?: boolean | null;
1463
+ litellm_params?: UpdateLiteLLMParams | null;
1464
+ model_info?: LitellmTypesRouterModelInfo | null;
1465
+ model_name?: string | null;
1466
+ }
1467
+ export declare const UpdateModelModelUpdatePostRequest: S.Schema<UpdateModelModelUpdatePostRequest>;
1468
+ export interface UpdateModelModelUpdatePostResponse {
1469
+ body: unknown;
1470
+ }
1471
+ export declare const UpdateModelModelUpdatePostResponse: S.Schema<UpdateModelModelUpdatePostResponse>;
1472
+ /** List of model group names to make public */
1473
+ export type UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList = Array<string>;
1474
+ export declare const UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList: S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList>;
1475
+ export interface UpdatePublicModelGroupsModelGroupMakePublicPostRequest {
1476
+ /** List of model group names to make public */
1477
+ model_groups: UpdatePublicModelGroupsModelGroupMakePublicPostRequestModelGroupsList;
1478
+ }
1479
+ export declare const UpdatePublicModelGroupsModelGroupMakePublicPostRequest: S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostRequest>;
1480
+ export interface UpdatePublicModelGroupsModelGroupMakePublicPostResponse {
1481
+ body: unknown;
1482
+ }
1483
+ export declare const UpdatePublicModelGroupsModelGroupMakePublicPostResponse: S.Schema<UpdatePublicModelGroupsModelGroupMakePublicPostResponse>;
1484
+ export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map = {
1485
+ [key: string]: unknown | undefined;
1486
+ };
1487
+ export declare const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map: S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map>;
1488
+ export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue = string | UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValueCase1Map;
1489
+ export declare const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue: S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue>;
1490
+ export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap = {
1491
+ [key: string]: UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksValue | undefined;
1492
+ };
1493
+ export declare const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap: S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap>;
1494
+ export interface UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest {
1495
+ useful_links: UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequestUsefulLinksMap;
1496
+ }
1497
+ export declare const UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest: S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest>;
1498
+ export interface UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse {
1499
+ body: unknown;
1500
+ }
1501
+ export declare const UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse: S.Schema<UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse>;
1502
+ export type ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap = {
1503
+ [key: string]: unknown | undefined;
1504
+ };
1505
+ export declare const ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap: S.Schema<ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap>;
1506
+ export interface ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest {
1507
+ complexity_router_config: ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequestComplexityRouterConfigMap;
1508
+ /** Team the router is being created for. Required for a team admin, who may only validate their own team's routers */
1509
+ team_id?: string | null;
1510
+ }
1511
+ export declare const ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest: S.Schema<ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest>;
1512
+ export interface ComplexityRouterConfigValidationResponse {
1513
+ error?: string | null;
1514
+ valid: boolean;
1515
+ }
1516
+ export declare const ComplexityRouterConfigValidationResponse: S.Schema<ComplexityRouterConfigValidationResponse>;
1517
+ export type AddNewModelModelNewPostError = UnprocessableEntity | LitellmOpError;
1518
+ /** Add New Model Allows adding new models to the model list in the config.yaml */
1519
+ export declare const addNewModelModelNewPost: API.OperationMethod<AddNewModelModelNewPostRequest, AddNewModelModelNewPostResponse, AddNewModelModelNewPostError, LitellmOpContext>;
1520
+ export type CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteError = LitellmOpError;
1521
+ /** Cancel Anthropic Beta Headers Reload ADMIN ONLY / MASTER KEY Only Endpoint Cancel the scheduled periodic reload of the Anthropic beta headers configuration. */
1522
+ export declare const cancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDelete: API.OperationMethod<CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteRequest, CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteResponse, CancelAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadDeleteError, LitellmOpContext>;
1523
+ export type CancelModelCostMapReloadScheduleModelCostMapReloadDeleteError = LitellmOpError;
1524
+ /** Cancel Model Cost Map Reload ADMIN ONLY / MASTER KEY Only Endpoint Cancel the scheduled periodic reload of the model cost map. */
1525
+ export declare const cancelModelCostMapReloadScheduleModelCostMapReloadDelete: API.OperationMethod<CancelModelCostMapReloadScheduleModelCostMapReloadDeleteRequest, CancelModelCostMapReloadScheduleModelCostMapReloadDeleteResponse, CancelModelCostMapReloadScheduleModelCostMapReloadDeleteError, LitellmOpContext>;
1526
+ export type CreateModelGroupAccessGroupNewPostError = UnprocessableEntity | LitellmOpError;
1527
+ /** Create Model Group Create a new access group containing multiple model names. An access group is a named collection of model groups that can be referenced by teams/keys for simplified access control. Example: ```bash curl -X POST 'http://localhost:4000/access_group/new' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "access_group": "production-models", "model_names": ["gpt-4", "claude-3-opus", "gemini-pro"] }' ``` Parameters: - access_group: str - The access group name (e.g., "production-models") - model_names: List[str] - List of existing model groups to include Returns: - NewModelGroupResponse with the created access group details Raises: - HTTPException 400: If any model names don't exist - HTTPException 500: If database operations fail */
1528
+ export declare const createModelGroupAccessGroupNewPost: API.OperationMethod<CreateModelGroupAccessGroupNewPostRequest, NewModelGroupResponse, CreateModelGroupAccessGroupNewPostError, LitellmOpContext>;
1529
+ export type DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteError = UnprocessableEntity | LitellmOpError;
1530
+ /** Delete Access Group Delete an access group. Removes the access group from all deployments that have it. Example: ```bash curl -X DELETE 'http://localhost:4000/access_group/production-models/delete' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - DeleteModelGroupResponse with deletion details Raises: - HTTPException 404: If access group not found */
1531
+ export declare const deleteAccessGroupAccessGroupAccessGroupDeleteDelete: API.OperationMethod<DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteRequest, DeleteModelGroupResponse, DeleteAccessGroupAccessGroupAccessGroupDeleteDeleteError, LitellmOpContext>;
1532
+ export type DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteError = UnprocessableEntity | LitellmOpError;
1533
+ /** Delete Access Group Budget Clear the shared budget of an access group, leaving the group itself in place. Example: ```bash curl -X DELETE 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - DeleteAccessGroupBudgetResponse; budget_deleted is false when there was nothing to clear Raises: - HTTPException 404: If access group not found */
1534
+ export declare const deleteAccessGroupBudgetAccessGroupAccessGroupBudgetDelete: API.OperationMethod<DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteRequest, DeleteAccessGroupBudgetResponse, DeleteAccessGroupBudgetAccessGroupAccessGroupBudgetDeleteError, LitellmOpContext>;
1535
+ export type DeleteModelModelDeletePostError = BadRequest | UnprocessableEntity | LitellmOpError;
1536
+ /** Delete Model Allows deleting models in the model list in the config.yaml */
1537
+ export declare const deleteModelModelDeletePost: API.OperationMethod<DeleteModelModelDeletePostRequest, DeleteModelModelDeletePostResponse, DeleteModelModelDeletePostError, LitellmOpContext>;
1538
+ export type GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetError = UnprocessableEntity | LitellmOpError;
1539
+ /** Get Access Group Budget Get the shared budget of an access group, and the spend drawn against it. Example: ```bash curl -X GET 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - AccessGroupBudgetResponse; budget is null when the group has no budget set Raises: - HTTPException 404: If access group not found */
1540
+ export declare const getAccessGroupBudgetAccessGroupAccessGroupBudgetGet: API.OperationMethod<GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetRequest, AccessGroupBudgetResponse, GetAccessGroupBudgetAccessGroupAccessGroupBudgetGetError, LitellmOpContext>;
1541
+ export type GetAccessGroupInfoAccessGroupAccessGroupInfoGetError = UnprocessableEntity | LitellmOpError;
1542
+ /** Get Access Group Info Get information about a specific access group. Example: ```bash curl -X GET 'http://localhost:4000/access_group/production-models/info' \ -H 'Authorization: Bearer sk-1234' ``` Parameters: - access_group: str - The access group name (URL path parameter) Returns: - AccessGroupInfo with the access group details, its shared budget and its spend Raises: - HTTPException 404: If access group not found */
1543
+ export declare const getAccessGroupInfoAccessGroupAccessGroupInfoGet: API.OperationMethod<GetAccessGroupInfoAccessGroupAccessGroupInfoGetRequest, AccessGroupInfo, GetAccessGroupInfoAccessGroupAccessGroupInfoGetError, LitellmOpContext>;
1544
+ export type GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetError = LitellmOpError;
1545
+ /** Get Anthropic Beta Headers Reload Status ADMIN ONLY / MASTER KEY Only Endpoint Get the status of the scheduled Anthropic beta headers reload job. */
1546
+ export declare const getAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGet: API.OperationMethod<GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetRequest, GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetResponse, GetAnthropicBetaHeadersReloadStatusScheduleAnthropicBetaHeadersReloadStatusGetError, LitellmOpContext>;
1547
+ export type GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetError = UnprocessableEntity | LitellmOpError;
1548
+ /** Get Auto Router Classifier Default Prompt Get the built-in system prompt used by an auto-router's LLM classifier */
1549
+ export declare const getAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGet: API.OperationMethod<GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetRequest, AutoRouterClassifierDefaultPromptResponse, GetAutoRouterClassifierDefaultPromptAutoRouterClassifierDefaultPromptGetError, LitellmOpContext>;
1550
+ export type GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetError = LitellmOpError;
1551
+ /** Get Model Cost Map Reload Status ADMIN ONLY / MASTER KEY Only Endpoint Get the status of the scheduled model cost map reload job. */
1552
+ export declare const getModelCostMapReloadStatusScheduleModelCostMapReloadStatusGet: API.OperationMethod<GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetRequest, GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetResponse, GetModelCostMapReloadStatusScheduleModelCostMapReloadStatusGetError, LitellmOpContext>;
1553
+ export type GetModelCostMapSourceModelCostMapSourceGetError = LitellmOpError;
1554
+ /** Get Model Cost Map Source ADMIN ONLY / MASTER KEY Only Endpoint Returns information about where the current model cost/pricing data was loaded from. Response fields: - source: "local" (bundled backup) or "remote" (fetched from URL) - url: the remote URL that was attempted (null when env-forced local) - is_env_forced: true if LITELLM_LOCAL_MODEL_COST_MAP=True forced local usage - fallback_reason: human-readable reason why remote failed (null on success) - model_count: number of models in the currently loaded cost map */
1555
+ export declare const getModelCostMapSourceModelCostMapSourceGet: API.OperationMethod<GetModelCostMapSourceModelCostMapSourceGetRequest, GetModelCostMapSourceModelCostMapSourceGetResponse, GetModelCostMapSourceModelCostMapSourceGetError, LitellmOpContext>;
1556
+ export type GetModelDeprecationsModelDeprecationError = UnprocessableEntity | LitellmOpError;
1557
+ /** Model Deprecations List models with known deprecation/sunset dates, bucketed by urgency. Reads `deprecation_date` metadata from `model_prices_and_context_window.json` (and any per-deployment `model_info.deprecation_date` overrides) for the models configured on this proxy. Parameters: warn_within_days: Window (in days) used to bucket "imminent" models, 30 by default. Returns: A payload with three lists of `ModelDeprecationInfo` entries: - `deprecated`: deprecation date is in the past, so these requests may fail at any time. - `imminent`: deprecation date is within `warn_within_days` from today. - `upcoming`: deprecation date is further out. Example: ```shell curl -X GET 'http://localhost:4000/model/deprecations' \ -H 'Authorization: Bearer sk-1234' ``` */
1558
+ export declare const getModelDeprecationsModelDeprecation: API.OperationMethod<GetModelDeprecationsModelDeprecationRequest, ModelDeprecationResponse, GetModelDeprecationsModelDeprecationError, LitellmOpContext>;
1559
+ export type GetModelDeprecationsV1ModelDeprecationError = UnprocessableEntity | LitellmOpError;
1560
+ /** Model Deprecations List models with known deprecation/sunset dates, bucketed by urgency. Reads `deprecation_date` metadata from `model_prices_and_context_window.json` (and any per-deployment `model_info.deprecation_date` overrides) for the models configured on this proxy. Parameters: warn_within_days: Window (in days) used to bucket "imminent" models, 30 by default. Returns: A payload with three lists of `ModelDeprecationInfo` entries: - `deprecated`: deprecation date is in the past, so these requests may fail at any time. - `imminent`: deprecation date is within `warn_within_days` from today. - `upcoming`: deprecation date is further out. Example: ```shell curl -X GET 'http://localhost:4000/model/deprecations' \ -H 'Authorization: Bearer sk-1234' ``` */
1561
+ export declare const getModelDeprecationsV1ModelDeprecation: API.OperationMethod<GetModelDeprecationsV1ModelDeprecationRequest, ModelDeprecationResponse, GetModelDeprecationsV1ModelDeprecationError, LitellmOpContext>;
1562
+ export type GetModelGroupInfoModelGroupInfoError = UnprocessableEntity | LitellmOpError;
1563
+ /** Model Group Info Get information about all the deployments on litellm proxy, including config.yaml descriptions (except api key and api base) - /model_group/info returns all model groups. End users of proxy should use /model_group/info since those models will be used for /chat/completions, /embeddings, etc. - /model_group/info?model_group=rerank-english-v3.0 returns all model groups for a specific model group (`model_name` in config.yaml) Example Request (All Models): ```shell curl -X 'GET' 'http://localhost:4000/model_group/info' -H 'accept: application/json' -H 'x-api-key: sk-1234' ``` Example Request (Specific Model Group): ```shell curl -X 'GET' 'http://localhost:4000/model_group/info?model_group=rerank-english-v3.0' -H 'accept: application/json' -H 'Authorization: Bearer sk-1234' ``` Example Request (Specific Wildcard Model Group): (e.g. `model_name: openai/*` on config.yaml) ```shell curl -X 'GET' 'http://localhost:4000/model_group/info?model_group=openai/tts-1' -H 'accept: application/json' -H 'Authorization: Bearersk-1234' ``` Learn how to use and set wildcard models [here](https://docs.litellm.ai/docs/wildcard_routing) Example Response: ```json { "data": [ { "model_group": "rerank-english-v3.0", "providers": [ "cohere" ], "max_input_tokens": null, "max_output_tokens": null, "input_cost_per_token": 0.0, "output_cost_per_token": 0.0, "mode": null, "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": false, "supports_function_calling": false, "supported_openai_params": [ "stream", "temperature", "max_tokens", "logit_bias", "top_p", "frequency_penalty", "presence_penalty", "stop", "n", "extra_headers" ] }, { "model_group": "gpt-3.5-turbo", "providers": [ "openai" ], "max_input_tokens": 16385.0, "max_output_tokens": 4096.0, "input_cost_per_token": 1.5e-06, "output_cost_per_token": 2e-06, "mode": "chat", "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": false, "supports_function_calling": true, "supported_openai_params": [ "frequency_penalty", "logit_bias", "logprobs", "top_logprobs", "max_tokens", "max_completion_tokens", "n", "presence_penalty", "seed", "stop", "stream", "stream_options", "temperature", "top_p", "tools", "tool_choice", "function_call", "functions", "max_retries", "extra_headers", "parallel_tool_calls", "response_format" ] }, { "model_group": "llava-hf", "providers": [ "openai" ], "max_input_tokens": null, "max_output_tokens": null, "input_cost_per_token": 0.0, "output_cost_per_token": 0.0, "mode": null, "tpm": null, "rpm": null, "supports_parallel_function_calling": false, "supports_vision": true, "supports_function_calling": false, "supported_openai_params": [ "frequency_penalty", "logit_bias", "logprobs", "top_logprobs", "max_tokens", "max_completion_tokens", "n", "presence_penalty", "seed", "stop", "stream", "stream_options", "temperature", "top_p", "tools", "tool_choice", "function_call", "functions", "max_retries", "extra_headers", "parallel_tool_calls", "response_format" ] } ] } ``` */
1564
+ export declare const getModelGroupInfoModelGroupInfo: API.OperationMethod<GetModelGroupInfoModelGroupInfoRequest, GetModelGroupInfoModelGroupInfoResponse, GetModelGroupInfoModelGroupInfoError, LitellmOpContext>;
1565
+ export type GetModelInfoModelsModelIdError = UnprocessableEntity | LitellmOpError;
1566
+ /** Model Info Retrieve information about a specific model accessible to your API key. Returns model details only if the model is available to your API key/team. Returns 404 if the model doesn't exist or is not accessible. Follows OpenAI API specification for individual model retrieval. https://platform.openai.com/docs/api-reference/models/retrieve Query parameters mirror `/v1/models` so the same caller context (team scoping, health filtering, paused deployments) drives both endpoints; the listing's public id must resolve to the same internal deployment here. */
1567
+ export declare const getModelInfoModelsModelId: API.OperationMethod<GetModelInfoModelsModelIdRequest, GetModelInfoModelsModelIdResponse, GetModelInfoModelsModelIdError, LitellmOpContext>;
1568
+ export type GetModelInfoV1ModelInfoError = UnprocessableEntity | LitellmOpError;
1569
+ /** Model Info V1 Provides more info about each model in /models, including config.yaml descriptions (except api key and api base) Parameters: litellm_model_id: Optional[str] = None (this is the value of `x-litellm-model-id` returned in response headers) - When litellm_model_id is passed, it will return the info for that specific model - When litellm_model_id is not passed, it will return the info for all models - include_team_models: When true, filter to deployments the caller can use (same as /v2/model/info). - teamId: Filter to models accessible by the given team. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks, matching `/v1/models?healthy_only=true`. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true`, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Ignored when `litellm_model_id` is passed, since that is a direct lookup of one deployment rather than a listing. Hiding is presentation-only: a hidden model can still be called directly. Each model in the list response includes `model_info.access_via_team_ids` and `model_info.direct_access` when the proxy database is connected. Returns: Returns a dictionary containing information about each model. Example Response: ```json { "data": [ { "model_name": "fake-openai-endpoint", "litellm_params": { "api_base": "https://exampleopenaiendpoint-production.up.railway.app/", "model": "openai/fake" }, "model_info": { "id": "112f74fab24a7a5245d2ced3536dd8f5f9192c57ee6e332af0f0512e08bed5af", "db_model": false } } ] } ``` */
1570
+ export declare const getModelInfoV1ModelInfo: API.OperationMethod<GetModelInfoV1ModelInfoRequest, GetModelInfoV1ModelInfoResponse, GetModelInfoV1ModelInfoError, LitellmOpContext>;
1571
+ export type GetModelInfoV1ModelsModelIdError = UnprocessableEntity | LitellmOpError;
1572
+ /** Model Info Retrieve information about a specific model accessible to your API key. Returns model details only if the model is available to your API key/team. Returns 404 if the model doesn't exist or is not accessible. Follows OpenAI API specification for individual model retrieval. https://platform.openai.com/docs/api-reference/models/retrieve Query parameters mirror `/v1/models` so the same caller context (team scoping, health filtering, paused deployments) drives both endpoints; the listing's public id must resolve to the same internal deployment here. */
1573
+ export declare const getModelInfoV1ModelsModelId: API.OperationMethod<GetModelInfoV1ModelsModelIdRequest, GetModelInfoV1ModelsModelIdResponse, GetModelInfoV1ModelsModelIdError, LitellmOpContext>;
1574
+ export type GetModelInfoV1V1ModelInfoError = UnprocessableEntity | LitellmOpError;
1575
+ /** Model Info V1 Provides more info about each model in /models, including config.yaml descriptions (except api key and api base) Parameters: litellm_model_id: Optional[str] = None (this is the value of `x-litellm-model-id` returned in response headers) - When litellm_model_id is passed, it will return the info for that specific model - When litellm_model_id is not passed, it will return the info for all models - include_team_models: When true, filter to deployments the caller can use (same as /v2/model/info). - teamId: Filter to models accessible by the given team. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks, matching `/v1/models?healthy_only=true`. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true`, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Ignored when `litellm_model_id` is passed, since that is a direct lookup of one deployment rather than a listing. Hiding is presentation-only: a hidden model can still be called directly. Each model in the list response includes `model_info.access_via_team_ids` and `model_info.direct_access` when the proxy database is connected. Returns: Returns a dictionary containing information about each model. Example Response: ```json { "data": [ { "model_name": "fake-openai-endpoint", "litellm_params": { "api_base": "https://exampleopenaiendpoint-production.up.railway.app/", "model": "openai/fake" }, "model_info": { "id": "112f74fab24a7a5245d2ced3536dd8f5f9192c57ee6e332af0f0512e08bed5af", "db_model": false } } ] } ``` */
1576
+ export declare const getModelInfoV1V1ModelInfo: API.OperationMethod<GetModelInfoV1V1ModelInfoRequest, GetModelInfoV1V1ModelInfoResponse, GetModelInfoV1V1ModelInfoError, LitellmOpContext>;
1577
+ export type GetModelInfoV2V2ModelInfoError = UnprocessableEntity | LitellmOpError;
1578
+ /** Model Info V2 Paginated model metadata for proxy deployments (pricing, provider, team access). Returns configured router deployments with enriched `model_info` (costs, provider, context window, etc.). Sensitive fields such as API keys and api_base are omitted. Query parameters: model: Filter to a single public `model_name`. user_models_only: When true, only return models created by the calling user. include_team_models: When true, populate `access_via_team_ids` and `direct_access` on each model and filter to deployments the caller can use. page / size: Pagination controls (defaults: page=1, size=50). search: Case-insensitive partial match on model name or team public name. modelId: Return a single deployment by LiteLLM model id. teamId: Filter to models with direct access or team membership for this team id. sortBy / sortOrder: Sort by model_name, created_at, updated_at, costs, or status. Example request: ``` curl -X GET 'http://localhost:4000/v2/model/info?include_team_models=true&page=1&size=50' \ --header 'Authorization: Bearer sk-1234' ``` Example response: ```json { "data": [ { "model_name": "gpt-4", "litellm_params": {"model": "openai/gpt-4.1"}, "model_info": { "id": "abc123", "litellm_provider": "openai", "access_via_team_ids": ["team-1"], "direct_access": true } } ], "total_count": 1, "current_page": 1, "total_pages": 1, "size": 50 } ``` */
1579
+ export declare const getModelInfoV2V2ModelInfo: API.OperationMethod<GetModelInfoV2V2ModelInfoRequest, GetModelInfoV2V2ModelInfoResponse, GetModelInfoV2V2ModelInfoError, LitellmOpContext>;
1580
+ export type GetModelMetricsExceptionsModelMetricsExceptionError = UnprocessableEntity | LitellmOpError;
1581
+ /** Model Metrics Exceptions View number of failed requests per model on config.yaml */
1582
+ export declare const getModelMetricsExceptionsModelMetricsException: API.OperationMethod<GetModelMetricsExceptionsModelMetricsExceptionRequest, GetModelMetricsExceptionsModelMetricsExceptionResponse, GetModelMetricsExceptionsModelMetricsExceptionError, LitellmOpContext>;
1583
+ export type GetModelMetricsModelMetricsError = UnprocessableEntity | LitellmOpError;
1584
+ /** Model Metrics View number of requests & avg latency per model on config.yaml */
1585
+ export declare const getModelMetricsModelMetrics: API.OperationMethod<GetModelMetricsModelMetricsRequest, GetModelMetricsModelMetricsResponse, GetModelMetricsModelMetricsError, LitellmOpContext>;
1586
+ export type GetModelMetricsSlowResponsesModelMetricsSlowResponseError = UnprocessableEntity | LitellmOpError;
1587
+ /** Model Metrics Slow Responses View number of hanging requests per model_group */
1588
+ export declare const getModelMetricsSlowResponsesModelMetricsSlowResponse: API.OperationMethod<GetModelMetricsSlowResponsesModelMetricsSlowResponseRequest, GetModelMetricsSlowResponsesModelMetricsSlowResponseResponse, GetModelMetricsSlowResponsesModelMetricsSlowResponseError, LitellmOpContext>;
1589
+ export type GetModelSettingsModelSettingsError = LitellmOpError;
1590
+ /** Model Settings Returns provider name, description, and required parameters for each provider */
1591
+ export declare const getModelSettingsModelSettings: API.OperationMethod<GetModelSettingsModelSettingsRequest, GetModelSettingsModelSettingsResponse, GetModelSettingsModelSettingsError, LitellmOpContext>;
1592
+ export type GetModelStreamingMetricsModelStreamingMetricsError = UnprocessableEntity | LitellmOpError;
1593
+ /** Model Streaming Metrics View time to first token for models in spend logs */
1594
+ export declare const getModelStreamingMetricsModelStreamingMetrics: API.OperationMethod<GetModelStreamingMetricsModelStreamingMetricsRequest, GetModelStreamingMetricsModelStreamingMetricsResponse, GetModelStreamingMetricsModelStreamingMetricsError, LitellmOpContext>;
1595
+ export type ListAccessGroupsAccessGroupListGetError = LitellmOpError;
1596
+ /** List Access Groups List all access groups. Returns a list of all access groups with their model names, deployment counts, shared budget and the spend drawn against it. Example: ```bash curl -X GET 'http://localhost:4000/access_group/list' \ -H 'Authorization: Bearer sk-1234' ``` Returns: - ListAccessGroupsResponse with all access groups */
1597
+ export declare const listAccessGroupsAccessGroupListGet: API.OperationMethod<ListAccessGroupsAccessGroupListGetRequest, ListAccessGroupsResponse, ListAccessGroupsAccessGroupListGetError, LitellmOpContext>;
1598
+ export type ModelListModelsGetError = BadRequest | UnprocessableEntity | LitellmOpError;
1599
+ /** Model List Use `/model/info` - to get detailed model information, example - pricing, mode, etc. This is just for compatibility with openai projects like aider. Query Parameters: - include_metadata: Include additional metadata in the response with fallback information - fallback_type: Type of fallbacks to include ("general", "context_window", "content_policy") Defaults to "general" when include_metadata=true - scope: Optional scope parameter. Currently only accepts "expand". When scope=expand is passed, proxy admins, team admins, and org admins will receive all proxy models as if they are a proxy admin. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true` in general_settings, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Models expanded from wildcard routes (e.g. `openai/*`) are not filtered, and nothing is hidden when `allowed_fails_policy` is configured (cooldown remains the sole exclusion mechanism). Hiding is presentation-only: a hidden model can still be called directly. */
1600
+ export declare const modelListModelsGet: API.OperationMethod<ModelListModelsGetRequest, ModelListModelsGetResponse, ModelListModelsGetError, LitellmOpContext>;
1601
+ export type ModelListV1ModelsGetError = UnprocessableEntity | LitellmOpError;
1602
+ /** Model List Use `/model/info` - to get detailed model information, example - pricing, mode, etc. This is just for compatibility with openai projects like aider. Query Parameters: - include_metadata: Include additional metadata in the response with fallback information - fallback_type: Type of fallbacks to include ("general", "context_window", "content_policy") Defaults to "general" when include_metadata=true - scope: Optional scope parameter. Currently only accepts "expand". When scope=expand is passed, proxy admins, team admins, and org admins will receive all proxy models as if they are a proxy admin. - healthy_only: When true, hide models whose backing deployments are all marked unhealthy by background health checks. Set `general_settings.model_list_healthy_only: true` to apply this to every caller without the query parameter. Requires `background_health_checks: true` in general_settings, plus either `model_list_healthy_only` or `enable_health_check_routing` to keep deployment health state cached; without health state the listing is returned unfiltered (fail open). Models expanded from wildcard routes (e.g. `openai/*`) are not filtered, and nothing is hidden when `allowed_fails_policy` is configured (cooldown remains the sole exclusion mechanism). Hiding is presentation-only: a hidden model can still be called directly. */
1603
+ export declare const modelListV1ModelsGet: API.OperationMethod<ModelListV1ModelsGetRequest, ModelListV1ModelsGetResponse, ModelListV1ModelsGetError, LitellmOpContext>;
1604
+ export type PatchModelModelModelIdUpdatePatchError = UnprocessableEntity | LitellmOpError;
1605
+ /** Patch Model PATCH Endpoint for partial model updates. Only updates the fields specified in the request while preserving other existing values. Follows proper PATCH semantics by only modifying provided fields. Args: model_id: The ID of the model to update patch_data: The fields to update and their new values user_api_key_dict: User authentication information Returns: Updated model information Raises: ProxyException: For various error conditions including authentication and database errors */
1606
+ export declare const patchModelModelModelIdUpdatePatch: API.OperationMethod<PatchModelModelModelIdUpdatePatchRequest, PatchModelModelModelIdUpdatePatchResponse, PatchModelModelModelIdUpdatePatchError, LitellmOpContext>;
1607
+ export type PostBlockModelModelBlockError = UnprocessableEntity | LitellmOpError;
1608
+ /** Block Model Block a DB-stored model deployment from serving requests. Parameters: - model_id: str - The model deployment id to block. */
1609
+ export declare const postBlockModelModelBlock: API.OperationMethod<PostBlockModelModelBlockRequest, PostBlockModelModelBlockResponse, PostBlockModelModelBlockError, LitellmOpContext>;
1610
+ export type PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptError = UnprocessableEntity | LitellmOpError;
1611
+ /** Preview Auto Router Classifier Prompt Get the system prompt an auto-router's LLM classifier sends for an edited tier set */
1612
+ export declare const postPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPrompt: API.OperationMethod<PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptRequest, AutoRouterClassifierDefaultPromptResponse, PostPreviewAutoRouterClassifierPromptAutoRouterClassifierDefaultPromptError, LitellmOpContext>;
1613
+ export type PostPreviewAutoRouterRoutingAutoRouterTestRoutingError = UnprocessableEntity | LitellmOpError;
1614
+ /** Preview Auto Router Routing Route a single request through a complexity-router config and report where it landed. Answers "which model would this request get?" for a config that only exists in a form, so an auto router can be checked before it is created. The request is classified by the same pre-routing hook a live request runs, over the same messages, system prompt and tool definitions, then dropped: nothing is sent to the model it routed to, and no auto router is created. A heuristic config therefore spends nothing, while an `llm` classifier or semantic keyword matching bills its classifier/embedding call to the calling key, like Test Connection does. Send `messages` to classify a real turn, with `system` and `tools` beside it when the surface carries them top level, as Anthropic /v1/messages does. `prompt` is the single-ask shorthand and routes as one user turn with nothing around it. **Example Request:** ```json { "messages": [ {"role": "system", "content": "You are a database migration assistant"}, {"role": "user", "content": "the index is not unique"}, {"role": "assistant", "content": "Then two workers can both insert. Add a unique index"}, {"role": "user", "content": "ok do it"} ], "tools": [{"type": "function", "function": {"name": "Bash", "description": "Run a command"}}], "complexity_router_config": { "tiers": {"SIMPLE": ["gpt-4o-mini"], "REASONING": ["o3"]}, "classifier_type": "heuristic" } } ``` */
1615
+ export declare const postPreviewAutoRouterRoutingAutoRouterTestRouting: API.OperationMethod<PostPreviewAutoRouterRoutingAutoRouterTestRoutingRequest, AutoRouterRoutingTestResponse, PostPreviewAutoRouterRoutingAutoRouterTestRoutingError, LitellmOpContext>;
1616
+ export type PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderError = LitellmOpError;
1617
+ /** Reload Anthropic Beta Headers ADMIN ONLY / MASTER KEY Only Endpoint Manually reload the Anthropic beta headers configuration from the remote source. This will fetch fresh configuration from the anthropic_beta_headers_config.json file. */
1618
+ export declare const postReloadAnthropicBetaHeadersReloadAnthropicBetaHeader: API.OperationMethod<PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderRequest, PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderResponse, PostReloadAnthropicBetaHeadersReloadAnthropicBetaHeaderError, LitellmOpContext>;
1619
+ export type PostReloadModelCostMapReloadModelCostMapError = LitellmOpError;
1620
+ /** Reload Model Cost Map ADMIN ONLY / MASTER KEY Only Endpoint Manually reload the model cost map from the remote source. This will fetch fresh pricing data from the model_prices_and_context_window.json file. */
1621
+ export declare const postReloadModelCostMapReloadModelCostMap: API.OperationMethod<PostReloadModelCostMapReloadModelCostMapRequest, PostReloadModelCostMapReloadModelCostMapResponse, PostReloadModelCostMapReloadModelCostMapError, LitellmOpContext>;
1622
+ export type PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadError = UnprocessableEntity | LitellmOpError;
1623
+ /** Schedule Anthropic Beta Headers Reload ADMIN ONLY / MASTER KEY Only Endpoint Schedule periodic reload of the Anthropic beta headers configuration. This will create a background job that reloads the configuration every specified hours. */
1624
+ export declare const postScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReload: API.OperationMethod<PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadRequest, PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadResponse, PostScheduleAnthropicBetaHeadersReloadScheduleAnthropicBetaHeadersReloadError, LitellmOpContext>;
1625
+ export type PostScheduleModelCostMapReloadScheduleModelCostMapReloadError = UnprocessableEntity | LitellmOpError;
1626
+ /** Schedule Model Cost Map Reload ADMIN ONLY / MASTER KEY Only Endpoint Schedule periodic reload of the model cost map. This will create a background job that reloads the model cost map every specified hours. */
1627
+ export declare const postScheduleModelCostMapReloadScheduleModelCostMapReload: API.OperationMethod<PostScheduleModelCostMapReloadScheduleModelCostMapReloadRequest, PostScheduleModelCostMapReloadScheduleModelCostMapReloadResponse, PostScheduleModelCostMapReloadScheduleModelCostMapReloadError, LitellmOpContext>;
1628
+ export type PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetError = UnprocessableEntity | LitellmOpError;
1629
+ /** Set Access Group Budget Set or replace the shared budget of an access group. Idempotent. Every key that can reach a model in the group draws from this one budget. Example: ```bash curl -X PUT 'http://localhost:4000/access_group/production-models/budget' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "max_budget": 100.0, "budget_duration": "30d" }' ``` Parameters: - access_group: str - The access group name (URL path parameter) - max_budget: Optional[float] - Requests fail once the group's shared spend exceeds this - soft_budget: Optional[float] - Fires an alert when reached; requests still succeed - budget_duration: Optional[str] - Frequency of resetting the group's spend (e.g. '30d') - budget_id: Optional[str] - Link an existing budget instead of creating one Returns: - AccessGroupBudgetResponse with the stored budget and current spend Raises: - HTTPException 400: If no budget field is given, or budget_duration cannot be parsed - HTTPException 404: If access group not found */
1630
+ export declare const putSetAccessGroupBudgetAccessGroupAccessGroupBudget: API.OperationMethod<PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetRequest, AccessGroupBudgetResponse, PutSetAccessGroupBudgetAccessGroupAccessGroupBudgetError, LitellmOpContext>;
1631
+ export type UnblockModelModelUnblockPostError = UnprocessableEntity | LitellmOpError;
1632
+ /** Unblock Model Unblock a DB-stored model deployment so it can serve requests again. Parameters: - model_id: str - The model deployment id to unblock. */
1633
+ export declare const unblockModelModelUnblockPost: API.OperationMethod<UnblockModelModelUnblockPostRequest, UnblockModelModelUnblockPostResponse, UnblockModelModelUnblockPostError, LitellmOpContext>;
1634
+ export type UpdateAccessGroupAccessGroupAccessGroupUpdatePutError = UnprocessableEntity | LitellmOpError;
1635
+ /** Update Access Group Update an access group's model names. This will: 1. Remove the access group from all current deployments 2. Add the access group to all deployments for the new model_names list Example: ```bash curl -X PUT 'http://localhost:4000/access_group/production-models/update' \ -H 'Authorization: Bearer sk-1234' \ -H 'Content-Type: application/json' \ -d '{ "model_names": ["gpt-4", "claude-3-sonnet"] }' ``` Parameters: - access_group: str - The access group name (URL path parameter) - model_names: List[str] - New list of model groups to include Returns: - NewModelGroupResponse with the updated access group details Raises: - HTTPException 400: If any model names don't exist - HTTPException 404: If access group not found */
1636
+ export declare const updateAccessGroupAccessGroupAccessGroupUpdatePut: API.OperationMethod<UpdateAccessGroupAccessGroupAccessGroupUpdatePutRequest, NewModelGroupResponse, UpdateAccessGroupAccessGroupAccessGroupUpdatePutError, LitellmOpContext>;
1637
+ export type UpdateModelModelUpdatePostError = BadRequest | UnprocessableEntity | LitellmOpError;
1638
+ /** Update Model Edit existing model params */
1639
+ export declare const updateModelModelUpdatePost: API.OperationMethod<UpdateModelModelUpdatePostRequest, UpdateModelModelUpdatePostResponse, UpdateModelModelUpdatePostError, LitellmOpContext>;
1640
+ export type UpdatePublicModelGroupsModelGroupMakePublicPostError = UnprocessableEntity | LitellmOpError;
1641
+ /** Update Public Model Groups Update which model groups are public */
1642
+ export declare const updatePublicModelGroupsModelGroupMakePublicPost: API.OperationMethod<UpdatePublicModelGroupsModelGroupMakePublicPostRequest, UpdatePublicModelGroupsModelGroupMakePublicPostResponse, UpdatePublicModelGroupsModelGroupMakePublicPostError, LitellmOpContext>;
1643
+ export type UpdateUsefulLinksModelHubUpdateUsefulLinksPostError = UnprocessableEntity | LitellmOpError;
1644
+ /** Update Useful Links Update useful links */
1645
+ export declare const updateUsefulLinksModelHubUpdateUsefulLinksPost: API.OperationMethod<UpdateUsefulLinksModelHubUpdateUsefulLinksPostRequest, UpdateUsefulLinksModelHubUpdateUsefulLinksPostResponse, UpdateUsefulLinksModelHubUpdateUsefulLinksPostError, LitellmOpContext>;
1646
+ export type ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostError = UnprocessableEntity | LitellmOpError;
1647
+ /** Validate Complexity Router Config Validate a complexity-router config without saving it. Runs the same check every write path runs (the router's own pydantic model), so a form can show the backend's exact verdict while the operator is still editing rather than after a rejected save. Gated exactly like the save it rehearses: a proxy admin, or a team admin naming their own team. Nothing is created, routed, or billed. */
1648
+ export declare const validateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPost: API.OperationMethod<ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostRequest, ComplexityRouterConfigValidationResponse, ValidateComplexityRouterConfigAutoRouterValidateComplexityRouterConfigPostError, LitellmOpContext>;
1649
+ //# sourceMappingURL=model_management.d.ts.map