@agentjido/llmdb 2026.8.2 → 2026.8.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (180) hide show
  1. package/dist/full.js +177 -169
  2. package/dist/generated/manifest.d.ts +1 -1
  3. package/dist/generated/manifest.js +1 -1
  4. package/dist/generated/provider-loaders.js +8 -0
  5. package/dist/providers/abacus.d.ts +1 -1
  6. package/dist/providers/abacus.js +1 -1
  7. package/dist/providers/aiand.d.ts +1 -1
  8. package/dist/providers/aiand.js +1 -1
  9. package/dist/providers/aihubmix.d.ts +1 -1
  10. package/dist/providers/aihubmix.js +1 -1
  11. package/dist/providers/alibaba.d.ts +1 -1
  12. package/dist/providers/alibaba.js +1 -1
  13. package/dist/providers/alibaba_cn.js +1 -1
  14. package/dist/providers/alibaba_token_plan.d.ts +1 -1
  15. package/dist/providers/alibaba_token_plan.js +1 -1
  16. package/dist/providers/alibaba_token_plan_cn.d.ts +1 -1
  17. package/dist/providers/alibaba_token_plan_cn.js +1 -1
  18. package/dist/providers/amazon_bedrock.d.ts +1 -1
  19. package/dist/providers/amazon_bedrock.js +1 -1
  20. package/dist/providers/ambient.js +1 -1
  21. package/dist/providers/amd.d.ts +11 -0
  22. package/dist/providers/amd.js +5 -0
  23. package/dist/providers/anthropic.js +1 -1
  24. package/dist/providers/azure.d.ts +1 -1
  25. package/dist/providers/azure.js +1 -1
  26. package/dist/providers/azure_cognitive_services.d.ts +1 -1
  27. package/dist/providers/azure_cognitive_services.js +1 -1
  28. package/dist/providers/baseten.d.ts +1 -1
  29. package/dist/providers/baseten.js +1 -1
  30. package/dist/providers/berget.js +1 -1
  31. package/dist/providers/blueclaw.js +1 -1
  32. package/dist/providers/chutes.d.ts +1 -1
  33. package/dist/providers/chutes.js +1 -1
  34. package/dist/providers/clarifai.js +1 -1
  35. package/dist/providers/cloudferro_sherlock.js +1 -1
  36. package/dist/providers/cloudflare_ai_gateway.d.ts +1 -1
  37. package/dist/providers/cloudflare_ai_gateway.js +1 -1
  38. package/dist/providers/cloudflare_workers_ai.d.ts +1 -1
  39. package/dist/providers/cloudflare_workers_ai.js +1 -1
  40. package/dist/providers/coralbricks.d.ts +11 -0
  41. package/dist/providers/coralbricks.js +5 -0
  42. package/dist/providers/cortecs.d.ts +1 -1
  43. package/dist/providers/cortecs.js +1 -1
  44. package/dist/providers/crof.js +1 -1
  45. package/dist/providers/crossmodel.d.ts +1 -1
  46. package/dist/providers/crossmodel.js +1 -1
  47. package/dist/providers/crusoe.d.ts +11 -0
  48. package/dist/providers/crusoe.js +5 -0
  49. package/dist/providers/daoxe.js +1 -1
  50. package/dist/providers/deepinfra.d.ts +1 -1
  51. package/dist/providers/deepinfra.js +1 -1
  52. package/dist/providers/deepseek.js +1 -1
  53. package/dist/providers/digitalocean.d.ts +1 -1
  54. package/dist/providers/digitalocean.js +1 -1
  55. package/dist/providers/edenai.d.ts +11 -0
  56. package/dist/providers/edenai.js +5 -0
  57. package/dist/providers/empiriolabs.d.ts +1 -1
  58. package/dist/providers/empiriolabs.js +1 -1
  59. package/dist/providers/evroc.d.ts +1 -1
  60. package/dist/providers/evroc.js +1 -1
  61. package/dist/providers/fastrouter.js +1 -1
  62. package/dist/providers/fireworks_ai.d.ts +1 -1
  63. package/dist/providers/fireworks_ai.js +1 -1
  64. package/dist/providers/friendli.d.ts +1 -1
  65. package/dist/providers/friendli.js +1 -1
  66. package/dist/providers/github_copilot.d.ts +1 -1
  67. package/dist/providers/github_copilot.js +1 -1
  68. package/dist/providers/gmicloud.js +1 -1
  69. package/dist/providers/google.d.ts +1 -1
  70. package/dist/providers/google.js +1 -1
  71. package/dist/providers/google_vertex.d.ts +1 -1
  72. package/dist/providers/google_vertex.js +1 -1
  73. package/dist/providers/greenpt.js +1 -1
  74. package/dist/providers/groq.js +1 -1
  75. package/dist/providers/hetzner.d.ts +1 -1
  76. package/dist/providers/hetzner.js +1 -1
  77. package/dist/providers/huggingface.d.ts +1 -1
  78. package/dist/providers/huggingface.js +1 -1
  79. package/dist/providers/hyper.d.ts +1 -1
  80. package/dist/providers/hyper.js +1 -1
  81. package/dist/providers/impossibl.js +1 -1
  82. package/dist/providers/inceptron.d.ts +1 -1
  83. package/dist/providers/inceptron.js +1 -1
  84. package/dist/providers/inferx.d.ts +1 -1
  85. package/dist/providers/inferx.js +1 -1
  86. package/dist/providers/infomaniak.js +1 -1
  87. package/dist/providers/io_net.js +1 -1
  88. package/dist/providers/kenari.js +1 -1
  89. package/dist/providers/kilo.d.ts +1 -1
  90. package/dist/providers/kilo.js +1 -1
  91. package/dist/providers/llmgateway.d.ts +1 -1
  92. package/dist/providers/llmgateway.js +1 -1
  93. package/dist/providers/llmtr.d.ts +1 -1
  94. package/dist/providers/llmtr.js +1 -1
  95. package/dist/providers/lmstudio.js +1 -1
  96. package/dist/providers/meganova.js +1 -1
  97. package/dist/providers/merge_gateway.d.ts +1 -1
  98. package/dist/providers/merge_gateway.js +1 -1
  99. package/dist/providers/meta.d.ts +1 -1
  100. package/dist/providers/meta.js +1 -1
  101. package/dist/providers/modal.js +1 -1
  102. package/dist/providers/modelscope.js +1 -1
  103. package/dist/providers/nano_gpt.d.ts +1 -1
  104. package/dist/providers/nano_gpt.js +1 -1
  105. package/dist/providers/nearai.js +1 -1
  106. package/dist/providers/nebius.d.ts +1 -1
  107. package/dist/providers/nebius.js +1 -1
  108. package/dist/providers/neon.d.ts +1 -1
  109. package/dist/providers/neon.js +1 -1
  110. package/dist/providers/neuralwatt.d.ts +1 -1
  111. package/dist/providers/neuralwatt.js +1 -1
  112. package/dist/providers/novita_ai.js +1 -1
  113. package/dist/providers/nvidia.d.ts +1 -1
  114. package/dist/providers/nvidia.js +1 -1
  115. package/dist/providers/ofox.d.ts +1 -1
  116. package/dist/providers/ofox.js +1 -1
  117. package/dist/providers/opencode.d.ts +1 -1
  118. package/dist/providers/opencode.js +1 -1
  119. package/dist/providers/opencode_go.d.ts +1 -1
  120. package/dist/providers/opencode_go.js +1 -1
  121. package/dist/providers/openrouter.d.ts +1 -1
  122. package/dist/providers/openrouter.js +1 -1
  123. package/dist/providers/perplexity_agent.d.ts +1 -1
  124. package/dist/providers/perplexity_agent.js +1 -1
  125. package/dist/providers/pioneer.d.ts +1 -1
  126. package/dist/providers/pioneer.js +1 -1
  127. package/dist/providers/privatemode_ai.d.ts +1 -1
  128. package/dist/providers/privatemode_ai.js +1 -1
  129. package/dist/providers/regolo_ai.d.ts +1 -1
  130. package/dist/providers/regolo_ai.js +1 -1
  131. package/dist/providers/requesty.d.ts +1 -1
  132. package/dist/providers/requesty.js +1 -1
  133. package/dist/providers/runinfra.d.ts +11 -0
  134. package/dist/providers/runinfra.js +5 -0
  135. package/dist/providers/salad_cloud.d.ts +11 -0
  136. package/dist/providers/salad_cloud.js +5 -0
  137. package/dist/providers/scnet_token_plan.d.ts +11 -0
  138. package/dist/providers/scnet_token_plan.js +5 -0
  139. package/dist/providers/siliconflow.js +1 -1
  140. package/dist/providers/siliconflow_cn.js +1 -1
  141. package/dist/providers/snowflake_cortex.d.ts +1 -1
  142. package/dist/providers/snowflake_cortex.js +1 -1
  143. package/dist/providers/stackit.js +1 -1
  144. package/dist/providers/submodel.js +1 -1
  145. package/dist/providers/synthetic.js +1 -1
  146. package/dist/providers/tensorx.js +1 -1
  147. package/dist/providers/thinkingmachines.js +1 -1
  148. package/dist/providers/tinfoil.d.ts +1 -1
  149. package/dist/providers/tinfoil.js +1 -1
  150. package/dist/providers/togetherai.d.ts +1 -1
  151. package/dist/providers/togetherai.js +1 -1
  152. package/dist/providers/umans_ai.d.ts +1 -1
  153. package/dist/providers/umans_ai.js +1 -1
  154. package/dist/providers/umans_ai_coding_plan.d.ts +1 -1
  155. package/dist/providers/umans_ai_coding_plan.js +1 -1
  156. package/dist/providers/upstage.d.ts +1 -1
  157. package/dist/providers/upstage.js +1 -1
  158. package/dist/providers/venice.d.ts +1 -1
  159. package/dist/providers/venice.js +1 -1
  160. package/dist/providers/vercel.d.ts +1 -1
  161. package/dist/providers/vercel.js +1 -1
  162. package/dist/providers/vivgrid.d.ts +1 -1
  163. package/dist/providers/vivgrid.js +1 -1
  164. package/dist/providers/vultr.js +1 -1
  165. package/dist/providers/wandb.d.ts +1 -1
  166. package/dist/providers/wandb.js +1 -1
  167. package/dist/providers/watsonx.d.ts +11 -0
  168. package/dist/providers/watsonx.js +5 -0
  169. package/dist/providers/xai.d.ts +1 -1
  170. package/dist/providers/xai.js +1 -1
  171. package/dist/providers/zai_coding_plan.d.ts +1 -1
  172. package/dist/providers/zai_coding_plan.js +1 -1
  173. package/dist/providers/zeldoc.d.ts +1 -1
  174. package/dist/providers/zeldoc.js +1 -1
  175. package/dist/providers/zenmux.d.ts +1 -1
  176. package/dist/providers/zenmux.js +1 -1
  177. package/dist/providers/zhipuai_coding_plan.d.ts +1 -1
  178. package/dist/providers/zhipuai_coding_plan.js +1 -1
  179. package/dist/snapshot.js +352 -336
  180. package/package.json +2 -2
@@ -1,5 +1,5 @@
1
1
  import { createProviderCatalog } from "../provider.js";
2
2
 
3
- export const data = {"alias_of":null,"base_url":"https://api.cortecs.ai/v1","catalog_only":true,"config_schema":null,"doc":"https://api.cortecs.ai/v1/models","env":["CORTECS_API_KEY"],"exclude_models":null,"extra":null,"id":"cortecs","models":{"claude-4-5-sonnet":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":3.259,"output":16.296},"deprecated":false,"extra":{"attachment":true,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"temperature":true},"family":null,"id":"claude-4-5-sonnet","knowledge":"2025-07-31","last_updated":"2025-09-29","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude 4.5 Sonnet","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-09-29","retired":false,"tags":null},"claude-4-6-sonnet":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":3.59,"output":17.92},"deprecated":false,"extra":{"attachment":true,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"temperature":true},"family":null,"id":"claude-4-6-sonnet","knowledge":"2025-08-31","last_updated":"2026-03-13","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Sonnet 4.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-17","retired":false,"tags":null},"claude-haiku-4-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.09,"output":5.43},"deprecated":false,"extra":{"attachment":true,"description":"Fast Claude model for responsive assistance, classification, and lightweight agents","family":"claude-haiku","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"temperature":true},"family":null,"id":"claude-haiku-4-5","knowledge":"2025-02-28","last_updated":"2025-10-15","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Haiku 4.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-15","retired":false,"tags":null},"claude-opus4-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":5.98,"output":29.89},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"temperature":true},"family":null,"id":"claude-opus4-5","knowledge":"2025-03-31","last_updated":"2025-11-24","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Opus 4.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-11-24","retired":false,"tags":null},"claude-opus4-6":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":5.98,"output":29.89},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"temperature":true},"family":null,"id":"claude-opus4-6","knowledge":"2025-05-31","last_updated":"2026-03-13","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Opus 4.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-05","retired":false,"tags":null},"claude-opus4-7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.56,"cache_write":6.99,"input":5.6,"output":27.99},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":false},"family":null,"id":"claude-opus4-7","knowledge":"2026-01-31","last_updated":"2026-04-16","lifecycle":null,"limits":{"context":1000000,"output":128000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Opus 4.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-16","retired":false,"tags":null},"claude-opus4-8":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.563,"cache_write":7.049,"input":5.64,"output":28.198},"deprecated":false,"extra":{"attachment":true,"description":"Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":false},"family":null,"id":"claude-opus4-8","knowledge":"2026-01","last_updated":"2026-05-28","lifecycle":null,"limits":{"context":1000000,"output":128000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Opus 4.8","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-05-28","retired":false,"tags":null},"claude-sonnet-4":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":3.307,"output":16.536},"deprecated":false,"extra":{"attachment":false,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"temperature":true},"family":null,"id":"claude-sonnet-4","knowledge":"2025-03","last_updated":"2025-05-22","lifecycle":null,"limits":{"context":200000,"output":64000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"Claude Sonnet 4","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-22","retired":false,"tags":null},"codestral-2508":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.03,"input":0.3,"output":0.9},"deprecated":false,"extra":{"attachment":false,"description":"Mistral coding model for code completion, generation, and developer workflows","family":"mistral","open_weights":true,"temperature":true},"family":null,"id":"codestral-2508","knowledge":"2025-03","last_updated":"2025-07-30","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Codestral 2508","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-30","retired":false,"tags":null},"deepseek-r1-0528":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.585,"output":2.307},"deprecated":false,"extra":{"attachment":false,"description":"DeepSeek reasoning model for multi-step analysis, math, coding, and tools","family":"deepseek-thinking","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"deepseek-r1-0528","knowledge":"2024-07","last_updated":"2025-05-28","lifecycle":null,"limits":{"context":164000,"output":164000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek R1 0528","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-28","retired":false,"tags":null},"deepseek-v3-0324":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.551,"output":1.654},"deprecated":false,"extra":{"attachment":false,"description":"DeepSeek chat model for instruction following, coding, and analysis","family":"deepseek","open_weights":true,"temperature":true},"family":null,"id":"deepseek-v3-0324","knowledge":"2024-07","last_updated":"2025-03-24","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V3 0324","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-03-24","retired":false,"tags":null},"deepseek-v3.2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.266,"output":0.444},"deprecated":false,"extra":{"attachment":false,"description":"DeepSeek chat model for instruction following, coding, and analysis","family":"deepseek","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"deepseek-v3.2","knowledge":"2024-07","last_updated":"2025-12-01","lifecycle":null,"limits":{"context":163840,"output":163840},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V3.2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-01","retired":false,"tags":null},"deepseek-v4-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.0028,"input":0.133,"output":0.266},"deprecated":false,"extra":{"attachment":false,"description":"Fast DeepSeek V4 lane for economical reasoning, coding, and long-context work","family":"deepseek-flash","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"deepseek-v4-flash","knowledge":"2025-05","last_updated":"2026-04-24","lifecycle":null,"limits":{"context":1048576,"output":384000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V4 Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-24","retired":false,"tags":null},"deepseek-v4-flash-0731":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.0645,"input":0.258,"output":0.31},"deprecated":false,"extra":{"attachment":false,"description":"Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding","family":"deepseek-flash","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"deepseek-v4-flash-0731","knowledge":"2025-05","last_updated":"2026-07-31","lifecycle":null,"limits":{"context":1048576,"output":384000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V4 Flash 0731","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-31","retired":false,"tags":null},"deepseek-v4-pro":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.003625,"input":1.553,"output":3.106},"deprecated":false,"extra":{"attachment":false,"description":"Open MoE flagship with million-token context for coding and long agent runs","family":"deepseek-thinking","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"deepseek-v4-pro","knowledge":"2025-05","last_updated":"2026-04-24","lifecycle":null,"limits":{"context":1048576,"output":384000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V4 Pro","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-24","retired":false,"tags":null},"devstral-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0,"output":0},"deprecated":false,"extra":{"attachment":false,"description":"Mistral coding agent model for repository tasks and software engineering workflows","open_weights":true,"temperature":true},"family":null,"id":"devstral-2512","knowledge":"2025-12","last_updated":"2025-12-09","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Devstral 2 2512","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-09","retired":false,"tags":null},"gemini-2.5-pro":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.654,"output":11.024},"deprecated":false,"extra":{"attachment":false,"description":"Advanced Gemini model for complex reasoning, coding, and multimodal analysis","family":"gemini-pro","open_weights":false,"temperature":true},"family":null,"id":"gemini-2.5-pro","knowledge":"2025-01","last_updated":"2025-06-17","lifecycle":null,"limits":{"context":1048576,"output":65535},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Gemini 2.5 Pro","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-03-20","retired":false,"tags":null},"glm-4.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.67,"output":2.46},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-4.5","knowledge":"2025-04","last_updated":"2025-07-29","lifecycle":null,"limits":{"context":131072,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM 4.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-29","retired":false,"tags":null},"glm-4.5-air":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.22,"output":1.34},"deprecated":false,"extra":{"attachment":false,"description":"Efficient GLM model for fast reasoning, coding, and agent workflows","family":"glm-air","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-4.5-air","knowledge":"2025-04","last_updated":"2025-08-01","lifecycle":null,"limits":{"context":131072,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM 4.5 Air","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-01","retired":false,"tags":null},"glm-4.7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.45,"output":2.23},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-4.7","knowledge":"2025-04","last_updated":"2025-12-22","lifecycle":null,"limits":{"context":198000,"output":198000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM 4.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-22","retired":false,"tags":null},"glm-4.7-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.09,"output":0.53},"deprecated":false,"extra":{"attachment":false,"description":"Efficient GLM model for fast reasoning, coding, and agent workflows","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-4.7-flash","knowledge":"2025-04","last_updated":"2025-08-08","lifecycle":null,"limits":{"context":203000,"output":203000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-4.7-Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-08","retired":false,"tags":null},"glm-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.08,"output":3.44},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-5","knowledge":null,"last_updated":"2026-02-11","lifecycle":null,"limits":{"context":202752,"output":202752},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM 5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-11","retired":false,"tags":null},"glm-5-turbo":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.308,"cache_write":1.544,"input":1.235,"output":4.118},"deprecated":false,"extra":{"attachment":false,"description":"Faster GLM-5 lane for coding agents that need lower latency","family":"glm","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5-turbo","knowledge":null,"last_updated":"2026-03-16","lifecycle":null,"limits":{"context":200000,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5-Turbo","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-16","retired":false,"tags":null},"glm-5.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.24,"input":1.31,"output":4.1},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5.1","knowledge":null,"last_updated":"2026-04-14","lifecycle":null,"limits":{"context":204800,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-14","retired":false,"tags":null},"glm-5.2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.39,"input":1.44,"output":4.53},"deprecated":false,"extra":{"attachment":false,"description":"Open flagship GLM for long-horizon coding agents and million-token context work","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["high","max"]}],"structured_output":true,"temperature":true},"family":null,"id":"glm-5.2","knowledge":null,"last_updated":"2026-06-13","lifecycle":null,"limits":{"context":1000000,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5.2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-13","retired":false,"tags":null},"glm-5v-turbo":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.308,"cache_write":1.544,"input":1.235,"output":4.118},"deprecated":false,"extra":{"attachment":true,"description":"Fast GLM vision model for screenshots, documents, and multimodal agent tasks","family":"glm","open_weights":false,"reasoning_options":[],"temperature":true},"family":null,"id":"glm-5v-turbo","knowledge":null,"last_updated":"2026-04-01","lifecycle":null,"limits":{"context":200000,"output":131072},"modalities":{"input":["text","image","video","pdf"],"output":["text"]},"model":null,"name":"GLM-5V-Turbo","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-01","retired":false,"tags":null},"gpt-4.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":2.354,"output":9.417},"deprecated":false,"extra":{"attachment":false,"description":"GPT model for general reasoning, writing, coding, and tool-assisted tasks","family":"gpt","open_weights":false,"temperature":true},"family":null,"id":"gpt-4.1","knowledge":"2024-06","last_updated":"2025-04-14","lifecycle":null,"limits":{"context":1047576,"output":32768},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT 4.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"gpt-5.4":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.25,"input":3,"output":16.13},"deprecated":false,"extra":{"attachment":true,"description":"Agent-ready GPT for coding and computer-use workflows at a lower cost","family":"gpt","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":false},"family":null,"id":"gpt-5.4","knowledge":"2025-08-31","last_updated":"2026-03-05","lifecycle":null,"limits":{"context":1050000,"output":128000},"modalities":{"input":["text","image","pdf"],"output":["text"]},"model":null,"name":"GPT-5.4","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-05","retired":false,"tags":null},"gpt-oss-120b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0,"output":0},"deprecated":false,"extra":{"attachment":false,"description":"Open-weight GPT model for self-hosted reasoning and instruction-following workloads","family":"gpt-oss","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"gpt-oss-120b","knowledge":"2024-01","last_updated":"2025-08-05","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GPT Oss 120b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-05","retired":false,"tags":null},"hermes-4-70b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.116,"output":0.358},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"hermes-4-70b","knowledge":"2023-12","last_updated":"2025-08-26","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Hermes 4 70B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-26","retired":false,"tags":null},"hy3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.103,"input":0.41,"output":1.025},"deprecated":false,"extra":{"attachment":false,"description":"Tencent Hy reasoning model for coding, instruction following, and agent tasks","family":"Hy","open_weights":true,"reasoning_options":[{"type":"effort","values":["none","low","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"hy3","knowledge":null,"last_updated":"2026-07-06","lifecycle":null,"limits":{"context":262144,"output":64000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Hy3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-06","retired":false,"tags":null},"intellect-3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.219,"output":1.202},"deprecated":false,"extra":{"attachment":true,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"intellect-3","knowledge":"2025-11","last_updated":"2025-11-26","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"INTELLECT 3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-11-26","retired":false,"tags":null},"kimi-k2-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.551,"output":2.646},"deprecated":false,"extra":{"attachment":false,"description":"Kimi model for long-context chat, coding, and agentic reasoning","family":"kimi-k2","open_weights":true,"temperature":true},"family":null,"id":"kimi-k2-instruct","knowledge":"2024-07","last_updated":"2025-09-05","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Kimi K2 Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-11","retired":false,"tags":null},"kimi-k2-thinking":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.656,"output":2.731},"deprecated":false,"extra":{"attachment":true,"description":"Kimi reasoning model for long-horizon research, planning, and tool use","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"kimi-k2-thinking","knowledge":"2025-12","last_updated":"2025-12-08","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Kimi K2 Thinking","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-08","retired":false,"tags":null},"kimi-k2.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.55,"output":2.76},"deprecated":false,"extra":{"attachment":false,"description":"Kimi reasoning model for long-horizon research, planning, and tool use","family":"kimi-thinking","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"toggle"}],"temperature":true},"family":null,"id":"kimi-k2.5","knowledge":"2025-01","last_updated":"2026-01-27","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image","video"],"output":["text"]},"model":null,"name":"Kimi K2.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-01-27","retired":false,"tags":null},"kimi-k2.6":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.2,"input":0.81,"output":3.54},"deprecated":false,"extra":{"attachment":false,"description":"Kimi reasoning model for long-horizon research, planning, and tool use","family":"kimi-thinking","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"toggle"}],"temperature":true},"family":null,"id":"kimi-k2.6","knowledge":null,"last_updated":"2026-04-17","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K2.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-17","retired":false,"tags":null},"kimi-k2.7-code":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.32,"input":1.28,"output":4.63},"deprecated":false,"extra":{"attachment":true,"description":"Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking","family":"kimi-k2","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"kimi-k2.7-code","knowledge":"2025-01","last_updated":"2026-06-12","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text","image","video"],"output":["text"]},"model":null,"name":"Kimi K2.7 Code","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-12","retired":false,"tags":null},"kimi-k3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":3,"output":15},"deprecated":false,"extra":{"attachment":true,"description":"Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work","family":"kimi-k3","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"kimi-k3","knowledge":null,"last_updated":"2026-07-16","lifecycle":null,"limits":{"context":1048576,"output":131072},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-16","retired":false,"tags":null},"llama-3.1-405b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0,"output":0},"deprecated":false,"extra":{"attachment":false,"description":"Open Llama instruction model for multilingual chat, reasoning, and coding","family":"llama","open_weights":true,"temperature":true},"family":null,"id":"llama-3.1-405b-instruct","knowledge":"2023-12","last_updated":"2024-07-23","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Llama 3.1 405B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-07-23","retired":false,"tags":null},"llama-3.3-70b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.089,"output":0.275},"deprecated":false,"extra":{"attachment":false,"description":"Popular open Llama workhorse for multilingual chat, coding, and self-hosting","family":"llama","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"llama-3.3-70b-instruct","knowledge":"2023-12","last_updated":"2024-12-06","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Llama 3.3 70B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-12-06","retired":false,"tags":null},"llama-4-maverick":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.03,"cache_write":0.151,"input":0.124,"output":0.603},"deprecated":false,"extra":{"attachment":true,"description":"Open multimodal Llama for strong reasoning with efficient everyday serving","family":"llama","open_weights":true,"temperature":true},"family":null,"id":"llama-4-maverick","knowledge":"2024-08","last_updated":"2025-04-05","lifecycle":null,"limits":{"context":1000000,"output":16384},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Llama 4 Maverick 17B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-05","retired":false,"tags":null},"minimax-m2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.39,"output":1.57},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"minimax-m2","knowledge":"2024-11","last_updated":"2025-10-27","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-27","retired":false,"tags":null},"minimax-m2.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.34,"output":1.34},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"minimax-m2.1","knowledge":null,"last_updated":"2025-12-23","lifecycle":null,"limits":{"context":196000,"output":196000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-23","retired":false,"tags":null},"minimax-m2.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.32,"output":1.18},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"minimax-m2.5","knowledge":null,"last_updated":"2026-02-12","lifecycle":null,"limits":{"context":196608,"output":196608},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-12","retired":false,"tags":null},"minimax-m2.7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.47,"output":1.4},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m2.7","knowledge":null,"last_updated":"2026-03-18","lifecycle":null,"limits":{"context":202752,"output":196072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-m2.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-18","retired":false,"tags":null},"minimax-m3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.089,"input":0.355,"output":1.775},"deprecated":false,"extra":{"attachment":true,"description":"MiniMax multimodal model for long-context coding, perception, and agent planning","family":"minimax","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"minimax-m3","knowledge":null,"last_updated":"2026-06-01","lifecycle":null,"limits":{"context":512000,"output":128000},"modalities":{"input":["text","image","video"],"output":["text"]},"model":null,"name":"MiniMax-M3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-01","retired":false,"tags":null},"mistral-large-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.05,"input":0.5,"output":1.5},"deprecated":false,"extra":{"attachment":true,"description":"Mistral's largest general model for enterprise agents, coding, and multilingual reasoning","family":"mistral-large","open_weights":true,"temperature":true},"family":null,"id":"mistral-large-2512","knowledge":"2025-12","last_updated":"2025-12-01","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Mistral Large 3 2512","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-01","retired":false,"tags":null},"mixtral-8x7B-instruct-v0.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"text":true,"tool_calls":false},"tools":{"enabled":false,"parallel":false,"streaming":false,"strict":false}},"catalog_only":true,"cost":{"input":0.438,"output":0.68},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"mixtral-8x7B-instruct-v0.1","knowledge":"2023-09","last_updated":"2023-12-11","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Mixtral 8x7B Instruct v0.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2023-12-11","retired":false,"tags":null},"nemotron-3-super-120b-a12b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.266,"output":0.799},"deprecated":false,"extra":{"attachment":false,"description":"Nemotron middle tier for collaborative agents and high-volume reasoning workloads","family":"nemotron","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"nemotron-3-super-120b-a12b","knowledge":"2025-12","last_updated":"2026-03-11","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Nemotron 3 Super 120B A12B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-11","retired":false,"tags":null},"nova-pro-v1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.016,"output":4.061},"deprecated":false,"extra":{"attachment":false,"description":"Flagship model for demanding analysis, coding, and production agent workflows","family":"nova-pro","open_weights":false,"temperature":true},"family":null,"id":"nova-pro-v1","knowledge":"2024-04","last_updated":"2024-12-03","lifecycle":null,"limits":{"context":300000,"output":5000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Nova Pro 1.0","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-12-03","retired":false,"tags":null},"qwen-2.5-72b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.062,"output":0.231},"deprecated":false,"extra":{"attachment":false,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"temperature":true},"family":null,"id":"qwen-2.5-72b-instruct","knowledge":"2024-06","last_updated":"2024-09-19","lifecycle":null,"limits":{"context":33000,"output":33000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen2.5 72B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-09-19","retired":false,"tags":null},"qwen3-235b-a22b-instruct-2507":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.062,"output":0.408},"deprecated":false,"extra":{"attachment":false,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3-235b-a22b-instruct-2507","knowledge":"2025-04","last_updated":"2025-07-23","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 235B A22B Instruct 2507","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-23","retired":false,"tags":null},"qwen3-32b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.099,"output":0.33},"deprecated":false,"extra":{"attachment":false,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"temperature":true},"family":null,"id":"qwen3-32b","knowledge":"2024-12","last_updated":"2025-04-29","lifecycle":null,"limits":{"context":16384,"output":16384},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 32B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-29","retired":false,"tags":null},"qwen3-coder-30b-a3b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.053,"output":0.222},"deprecated":false,"extra":{"attachment":false,"description":"Smaller Qwen coder for efficient local agents and repo-level fixes","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3-coder-30b-a3b-instruct","knowledge":"2025-04","last_updated":"2025-07-31","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 Coder 30B A3B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-31","retired":false,"tags":null},"qwen3-coder-480b-a35b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.441,"output":1.984},"deprecated":false,"extra":{"attachment":false,"description":"Qwen coding model for software agents, repository edits, and code reasoning","family":"qwen","open_weights":true,"temperature":true},"family":null,"id":"qwen3-coder-480b-a35b-instruct","knowledge":"2025-01","last_updated":"2025-07-25","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 Coder 480B A35B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-25","retired":false,"tags":null},"qwen3-coder-next":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.158,"output":0.84},"deprecated":false,"extra":{"attachment":false,"description":"Qwen coding model for software agents, repository edits, and code reasoning","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3-coder-next","knowledge":"2025-04","last_updated":"2026-02-04","lifecycle":null,"limits":{"context":256000,"output":65536},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 Coder Next 80B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-04","retired":false,"tags":null},"qwen3-next-80b-a3b-thinking":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.164,"output":1.311},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3-next-80b-a3b-thinking","knowledge":"2025-04","last_updated":"2025-09-11","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 Next 80B A3B Thinking","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-09-11","retired":false,"tags":null},"qwen3.5-122b-a10b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.444,"output":3.106},"deprecated":false,"extra":{"attachment":false,"description":"Qwen vision-language model for visual reasoning, documents, and agent tasks","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3.5-122b-a10b","knowledge":"2026-01","last_updated":"2026-02-24","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3.5 122B A10B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-24","retired":false,"tags":null},"qwen3.5-397b-a17b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.6,"output":3.6},"deprecated":false,"extra":{"attachment":false,"description":"Large open Qwen multimodal MoE for visual agents and long technical tasks","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3.5-397b-a17b","knowledge":"2026-01","last_updated":"2026-02-16","lifecycle":null,"limits":{"context":250000,"output":250000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3.5 397B A17B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-16","retired":false,"tags":null}},"name":"Cortecs","pricing_defaults":null};
3
+ export const data = {"alias_of":null,"base_url":"https://api.cortecs.ai/v1","catalog_only":true,"config_schema":null,"doc":"https://api.cortecs.ai/v1/models","env":["CORTECS_API_KEY"],"exclude_models":null,"extra":null,"id":"cortecs","models":{"apertus-70b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.393,"output":2.228},"deprecated":false,"extra":{"attachment":false,"description":"Apertus 70B is an open, multilingual language model designed for research, long-context reasoning, and sovereignty-focused AI systems.","open_weights":true,"reasoning_options":[],"structured_output":false,"temperature":false},"family":null,"id":"apertus-70b","knowledge":"2025-09","last_updated":"2025-09-02","lifecycle":null,"limits":{"context":65536,"output":65536},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Apertus 70B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-09-02","retired":false,"tags":null},"claude-4-5-sonnet":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.326,"cache_write":4.078,"input":2.989,"output":14.945},"deprecated":false,"extra":{"attachment":true,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-4-5-sonnet","knowledge":"2025-07-31","last_updated":"2025-09-29","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Sonnet 4.5 (latest)","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-09-29","retired":false,"tags":null},"claude-4-6-sonnet":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.32,"cache_write":3.999,"input":3.196,"output":15.94},"deprecated":false,"extra":{"attachment":true,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-4-6-sonnet","knowledge":"2025-08-31","last_updated":"2026-03-13","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Sonnet 4.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-17","retired":false,"tags":null},"claude-haiku-4-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.099,"cache_write":1.186,"input":0.996,"output":4.982},"deprecated":false,"extra":{"attachment":true,"description":"Fast Claude model for responsive assistance, classification, and lightweight agents","family":"claude-haiku","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-haiku-4-5","knowledge":"2025-02-28","last_updated":"2025-10-15","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Haiku 4.5 (latest)","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-15","retired":false,"tags":null},"claude-opus-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.55,"cache_write":6.874,"input":5.5,"output":27.498},"deprecated":false,"extra":{"attachment":true,"description":"Strongest Claude Opus model for coding, agents, and professional work","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":false},"family":null,"id":"claude-opus-5","knowledge":"2026-05","last_updated":"2026-07-24","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Opus 5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-24","retired":false,"tags":null},"claude-opus4-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.531,"cache_write":6.645,"input":5.313,"output":26.568},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-opus4-5","knowledge":"2025-05","last_updated":"2025-11-24","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Opus 4.5 (latest)","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-11-24","retired":false,"tags":null},"claude-opus4-6":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.531,"cache_write":6.645,"input":5.313,"output":26.561},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-opus4-6","knowledge":"2025-05-31","last_updated":"2026-03-13","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Opus 4.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-05","retired":false,"tags":null},"claude-opus4-7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.544,"cache_write":6.797,"input":5.437,"output":27.186},"deprecated":false,"extra":{"attachment":true,"description":"Flagship Claude model for deep reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":false},"family":null,"id":"claude-opus4-7","knowledge":"2026-01-31","last_updated":"2026-04-16","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Opus 4.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-16","retired":false,"tags":null},"claude-opus4-8":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.544,"cache_write":6.797,"input":5.437,"output":27.186},"deprecated":false,"extra":{"attachment":true,"description":"Top Claude Opus tier for the hardest reasoning, coding, and long-horizon agents","family":"claude-opus","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":false},"family":null,"id":"claude-opus4-8","knowledge":"2026-01","last_updated":"2026-05-28","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Opus 4.8","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-05-28","retired":false,"tags":null},"claude-sonnet-4":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.29,"cache_write":3.624,"input":2.898,"output":14.493},"deprecated":false,"extra":{"attachment":true,"description":"Balanced Claude model for coding, analysis, agent workflows, and cost control","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":true},"family":null,"id":"claude-sonnet-4","knowledge":"2025-03-31","last_updated":"2025-05-22","lifecycle":null,"limits":{"context":200000,"output":200000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Sonnet 4 (latest)","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-22","retired":false,"tags":null},"claude-sonnet-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.219,"cache_write":2.749,"input":2.2,"output":11},"deprecated":false,"extra":{"attachment":true,"description":"Everyday Claude agent model for coding, planning, browsing, and general work","family":"claude-sonnet","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]},{"min":1024,"type":"budget_tokens"}],"structured_output":true,"temperature":false},"family":null,"id":"claude-sonnet-5","knowledge":"2026-01-31","last_updated":"2026-06-30","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Claude Sonnet 5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-30","retired":false,"tags":null},"codestral-2508":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.033,"input":0.334,"output":1.003},"deprecated":false,"extra":{"attachment":false,"description":"Mistral coding model for code completion, generation, and developer workflows","family":"mistral","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"codestral-2508","knowledge":"2025-03","last_updated":"2025-07-30","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Codestral 2508","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-30","retired":false,"tags":null},"cosmos3-super-reasoner":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.099,"output":0.296},"deprecated":false,"extra":{"attachment":false,"description":"Cosmos3 Super Reasoner is a high-capacity reasoning model designed for complex multi-agent tasks and advanced physical AI understanding.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"cosmos3-super-reasoner","knowledge":null,"last_updated":"2026-06-02","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"cosmos3-super-reasoner","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-02","retired":false,"tags":null},"deepseek-r1-0528":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.163,"input":0.652,"output":2.57},"deprecated":false,"extra":{"attachment":false,"description":"DeepSeek reasoning model for multi-step analysis, math, coding, and tools","family":"deepseek-thinking","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"deepseek-r1-0528","knowledge":"2024-07","last_updated":"2025-05-28","lifecycle":null,"limits":{"context":164000,"output":164000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek R1 0528","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-28","retired":false,"tags":null},"deepseek-v3.2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.075,"input":0.296,"output":0.495},"deprecated":false,"extra":{"attachment":false,"description":"DeepSeek chat model for instruction following, coding, and analysis","family":"deepseek","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"deepseek-v3.2","knowledge":"2024-07","last_updated":"2025-12-01","lifecycle":null,"limits":{"context":163840,"output":163840},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V3.2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-01","retired":false,"tags":null},"deepseek-v4-flash-0731":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.03,"input":0.13,"output":0.28},"deprecated":false,"extra":{"attachment":false,"description":"Official DeepSeek V4 Flash release with enhanced agentic capabilities and integrated DSpark speculative decoding","family":"deepseek-flash","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"deepseek-v4-flash-0731","knowledge":"2025-05","last_updated":"2026-07-31","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V4 Flash 0731","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-31","retired":false,"tags":null},"deepseek-v4-pro":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.432,"input":1.73,"output":3.46},"deprecated":false,"extra":{"attachment":false,"description":"Open MoE flagship with million-token context for coding and long agent runs","family":"deepseek-thinking","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"temperature":true},"family":null,"id":"deepseek-v4-pro","knowledge":"2025-05","last_updated":"2026-04-24","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"DeepSeek V4 Pro","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-24","retired":false,"tags":null},"devstral-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.045,"input":0.446,"output":2.228},"deprecated":false,"extra":{"attachment":false,"description":"Mistral coding agent model for repository tasks and software engineering workflows","family":"devstral","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"devstral-2512","knowledge":"2025-12","last_updated":"2025-12-09","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Devstral 2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-09","retired":false,"tags":null},"gemini-2.5-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.029,"cache_write":0.097,"input":0.299,"output":2.491},"deprecated":false,"extra":{"attachment":true,"description":"Fast Gemini workhorse for multimodal apps where latency and price matter","family":"gemini-flash","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"gemini-2.5-flash","knowledge":"2025-01","last_updated":"2025-06-17","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 2.5 Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-06-17","retired":false,"tags":null},"gemini-2.5-pro":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.242,"cache_write":0.434,"input":1.495,"output":9.964},"deprecated":false,"extra":{"attachment":true,"description":"Advanced Gemini model for complex reasoning, coding, and multimodal analysis","family":"gemini-pro","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"gemini-2.5-pro","knowledge":"2025-01","last_updated":"2025-06-17","lifecycle":null,"limits":{"context":1048576,"output":65535},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 2.5 Pro","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-06-17","retired":false,"tags":null},"gemini-3.1-flash-lite":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.025,"cache_write":0.082,"input":0.272,"output":1.631},"deprecated":false,"extra":{"attachment":true,"description":"Low-latency Gemini model for high-volume multimodal and agent workloads","family":"gemini-flash-lite","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"gemini-3.1-flash-lite","knowledge":"2025-01","last_updated":"2026-05-07","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 3.1 Flash Lite","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-05-07","retired":false,"tags":null},"gemini-3.5-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.165,"cache_write":1,"input":1.649,"output":9.899},"deprecated":false,"extra":{"attachment":true,"description":"Fast Gemini model balancing multimodal reasoning, tool use, and cost","family":"gemini-flash","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gemini-3.5-flash","knowledge":"2025-01","last_updated":"2026-05-19","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 3.5 Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-05-19","retired":false,"tags":null},"gemini-3.5-flash-lite":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.033,"input":0.33,"output":2.749},"deprecated":false,"extra":{"attachment":true,"description":"Fast Gemini model balancing multimodal reasoning, tool use, and cost","family":"gemini-flash-lite","open_weights":false,"reasoning_options":[{"type":"effort","values":["minimal","low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gemini-3.5-flash-lite","knowledge":"2026-03","last_updated":"2026-07-21","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 3.5 Flash Lite","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-21","retired":false,"tags":null},"gemini-3.6-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.075,"cache_write":0.038,"input":0.75,"output":3.75},"deprecated":false,"extra":{"attachment":true,"description":"Fast Gemini model balancing multimodal reasoning, tool use, and cost","family":"gemini-flash","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gemini-3.6-flash","knowledge":"2026-03","last_updated":"2026-07-21","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 3.6 Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-21","retired":false,"tags":null},"gemini-3.7-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.075,"cache_write":0.038,"input":0.75,"output":3.75},"deprecated":false,"extra":{"attachment":true,"description":"High-efficiency Gemini model for agentic workflows, coding, and multimodal reasoning","family":"gemini-flash","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gemini-3.7-flash","knowledge":"2026-03","last_updated":"2026-08-13","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image","audio"],"output":["text"]},"model":null,"name":"Gemini 3.7 Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-08-13","retired":false,"tags":null},"gemma-3-27b-it":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"text":true,"tool_calls":false},"tools":{"enabled":false,"parallel":false,"streaming":false,"strict":false}},"catalog_only":true,"cost":{"input":0.099,"output":0.299},"deprecated":false,"extra":{"attachment":true,"description":"Gemma 3 is a family of lightweight, multimodal models from Google, supporting text and image inputs, multilingual capabilities, and a 131K context window.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"gemma-3-27b-it","knowledge":null,"last_updated":"2025-03-12","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"gemma-3-27b-it","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-03-12","retired":false,"tags":null},"gemma-4-26b-a4b-it":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.111,"output":0.557},"deprecated":false,"extra":{"attachment":true,"description":"Open Gemma instruction model for efficient chat and self-hosted deployments","family":"gemma","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"gemma-4-26b-a4b-it","knowledge":null,"last_updated":"2026-04-02","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Gemma 4 26B A4B IT","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-02","retired":false,"tags":null},"gemma-4-31b-it":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.223,"output":0.39},"deprecated":false,"extra":{"attachment":true,"description":"Largest Gemma 4 instruction model for open, self-hosted chat and reasoning","family":"gemma","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"gemma-4-31b-it","knowledge":null,"last_updated":"2026-04-02","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Gemma 4 31B IT","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-02","retired":false,"tags":null},"glm-4.7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.78,"output":2.785},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-4.7","knowledge":"2025-04","last_updated":"2025-12-22","lifecycle":null,"limits":{"context":202752,"output":198000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-4.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-22","retired":false,"tags":null},"glm-4.7-flash":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.08,"output":0.478},"deprecated":false,"extra":{"attachment":false,"description":"Efficient GLM model for fast reasoning, coding, and agent workflows","family":"glm-flash","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-4.7-flash","knowledge":"2025-04","last_updated":"2026-01-19","lifecycle":null,"limits":{"context":203000,"output":203000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-4.7-Flash","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-01-19","retired":false,"tags":null},"glm-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.247,"input":0.988,"output":3.164},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5","knowledge":null,"last_updated":"2026-02-12","lifecycle":null,"limits":{"context":202752,"output":202752},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-12","retired":false,"tags":null},"glm-5-turbo":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.296,"cache_write":1.544,"input":1.186,"output":3.955},"deprecated":false,"extra":{"attachment":false,"description":"Faster GLM-5 lane for coding agents that need lower latency","family":"glm","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5-turbo","knowledge":null,"last_updated":"2026-03-16","lifecycle":null,"limits":{"context":202752,"output":202752},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5-Turbo","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-16","retired":false,"tags":null},"glm-5.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.346,"input":1.384,"output":4.348},"deprecated":false,"extra":{"attachment":false,"description":"Flagship GLM model for hybrid reasoning, coding, and agentic engineering","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5.1","knowledge":null,"last_updated":"2026-04-07","lifecycle":null,"limits":{"context":202752,"output":202752},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-07","retired":false,"tags":null},"glm-5.2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.26,"input":1.2,"output":4.2},"deprecated":false,"extra":{"attachment":false,"description":"Open flagship GLM for long-horizon coding agents and million-token context work","family":"glm","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"effort","values":["high","max"]}],"structured_output":true,"temperature":true},"family":null,"id":"glm-5.2","knowledge":null,"last_updated":"2026-06-13","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GLM-5.2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-13","retired":false,"tags":null},"glm-5v-turbo":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.296,"cache_write":1.544,"input":1.186,"output":3.955},"deprecated":false,"extra":{"attachment":true,"description":"Fast GLM vision model for screenshots, documents, and multimodal agent tasks","family":"glm","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"glm-5v-turbo","knowledge":null,"last_updated":"2026-04-01","lifecycle":null,"limits":{"context":202752,"output":202752},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GLM-5V-Turbo","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-01","retired":false,"tags":null},"gpt-4.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.546,"input":2.192,"output":8.769},"deprecated":false,"extra":{"attachment":true,"description":"GPT model for general reasoning, writing, coding, and tool-assisted tasks","family":"gpt","open_weights":false,"structured_output":false,"temperature":true},"family":null,"id":"gpt-4.1","knowledge":"2024-04","last_updated":"2025-04-14","lifecycle":null,"limits":{"context":1047576,"output":1047576},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-4.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"gpt-4.1-mini":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.134,"input":0.434,"output":1.704},"deprecated":false,"extra":{"attachment":true,"description":"Affordable GPT-4.1 lane for fast coding help and structured extraction","family":"gpt-mini","open_weights":false,"structured_output":false,"temperature":true},"family":null,"id":"gpt-4.1-mini","knowledge":"2024-04","last_updated":"2025-04-14","lifecycle":null,"limits":{"context":1047576,"output":1047576},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-4.1 mini","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"gpt-4.1-nano":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.056,"input":0.111,"output":0.434},"deprecated":false,"extra":{"attachment":true,"description":"Tiny GPT-4.1 option for classification, routing, and very high-volume tasks","family":"gpt-nano","open_weights":false,"structured_output":false,"temperature":true},"family":null,"id":"gpt-4.1-nano","knowledge":"2024-04","last_updated":"2025-04-14","lifecycle":null,"limits":{"context":1047576,"output":1047576},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-4.1 nano","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"gpt-4o":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":1.33,"input":2.659,"output":10.635},"deprecated":false,"extra":{"attachment":true,"description":"Omni-era GPT for multimodal chat, practical coding, and general assistants","family":"gpt","open_weights":false,"structured_output":false,"temperature":true},"family":null,"id":"gpt-4o","knowledge":"2023-09","last_updated":"2024-08-06","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-4o","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-05-13","retired":false,"tags":null},"gpt-4o-mini":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.081,"input":0.159,"output":0.638},"deprecated":false,"extra":{"attachment":true,"description":"Small omni GPT for cheap multimodal assistance and production-scale traffic","family":"gpt-mini","open_weights":false,"structured_output":false,"temperature":true},"family":null,"id":"gpt-4o-mini","knowledge":"2023-09","last_updated":"2024-07-18","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-4o mini","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-07-18","retired":false,"tags":null},"gpt-5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.156,"input":1.375,"output":10.96},"deprecated":false,"extra":{"attachment":true,"description":"Original GPT-5 workhorse for reasoning, coding, writing, and tool workflows","family":"gpt","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":false,"temperature":false},"family":null,"id":"gpt-5","knowledge":"2024-09-30","last_updated":"2025-08-07","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-07","retired":false,"tags":null},"gpt-5-mini":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.056,"input":0.279,"output":2.192},"deprecated":false,"extra":{"attachment":true,"description":"Small GPT-5 for responsive agents, coding help, and everyday automation","family":"gpt-mini","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":false,"temperature":false},"family":null,"id":"gpt-5-mini","knowledge":"2024-05-30","last_updated":"2025-08-07","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5 Mini","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-07","retired":false,"tags":null},"gpt-5-nano":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.019,"input":0.06,"output":0.439},"deprecated":false,"extra":{"attachment":true,"description":"Tiny GPT-5 lane for routing, extraction, classification, and bulk jobs","family":"gpt-nano","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":false,"temperature":false},"family":null,"id":"gpt-5-nano","knowledge":"2024-05-30","last_updated":"2025-08-07","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5 Nano","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-07","retired":false,"tags":null},"gpt-5.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.156,"input":1.375,"output":10.96},"deprecated":false,"extra":{"attachment":true,"description":"Sharper GPT-5 generation for coding, product work, and tool-assisted tasks","family":"gpt","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":false,"temperature":false},"family":null,"id":"gpt-5.1","knowledge":"2024-09-30","last_updated":"2025-11-13","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-11-13","retired":false,"tags":null},"gpt-5.4":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.242,"input":2.898,"output":15.453},"deprecated":false,"extra":{"attachment":true,"description":"Agent-ready GPT for coding and computer-use workflows at a lower cost","family":"gpt","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":false},"family":null,"id":"gpt-5.4","knowledge":"2025-08-31","last_updated":"2026-03-05","lifecycle":null,"limits":{"context":1050000,"output":1050000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5.4","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-05","retired":false,"tags":null},"gpt-5.6-luna":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.11,"cache_write":1.38,"input":1.1,"output":6.599},"deprecated":false,"extra":{"attachment":true,"description":"Cost-efficient GPT-5.6 model for fast, high-volume workloads","family":"gpt-luna","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":false},"family":null,"id":"gpt-5.6-luna","knowledge":"2026-02-16","last_updated":"2026-07-09","lifecycle":null,"limits":{"context":1050000,"output":1050000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5.6 Luna","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-09","retired":false,"tags":null},"gpt-5.6-sol":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.55,"cache_write":6.879,"input":5.5,"output":32.998},"deprecated":false,"extra":{"attachment":true,"description":"Frontier GPT-5.6 model for complex professional work, coding, and agentic workflows","family":"gpt-sol","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":false},"family":null,"id":"gpt-5.6-sol","knowledge":"2026-02-16","last_updated":"2026-07-09","lifecycle":null,"limits":{"context":1050000,"output":1050000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5.6 Sol","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-09","retired":false,"tags":null},"gpt-5.6-terra":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.275,"cache_write":3.437,"input":2.749,"output":16.498},"deprecated":false,"extra":{"attachment":true,"description":"Balanced GPT-5.6 model for capable, cost-efficient everyday work","family":"gpt-terra","open_weights":false,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":false},"family":null,"id":"gpt-5.6-terra","knowledge":"2026-02-16","last_updated":"2026-07-09","lifecycle":null,"limits":{"context":1050000,"output":1050000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"GPT-5.6 Terra","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-09","retired":false,"tags":null},"gpt-oss-120b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.01,"input":0.089,"output":0.446},"deprecated":false,"extra":{"attachment":false,"description":"Open-weight GPT model for self-hosted reasoning and instruction-following workloads","family":"gpt-oss","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gpt-oss-120b","knowledge":null,"last_updated":"2025-08-05","lifecycle":null,"limits":{"context":131000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GPT OSS 120B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-05","retired":false,"tags":null},"gpt-oss-20b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.045,"output":0.167},"deprecated":false,"extra":{"attachment":false,"description":"Open GPT reasoning model for self-hosted agents and controllable deployments","family":"gpt-oss","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gpt-oss-20b","knowledge":null,"last_updated":"2025-08-05","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GPT OSS 20B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-05","retired":false,"tags":null},"gpt-oss-safeguard-120b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.179,"output":0.697},"deprecated":false,"extra":{"attachment":false,"description":"Safety model for policy screening, moderation, and risk-aware routing workflows","family":"gpt-oss","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","high"]}],"structured_output":true,"temperature":true},"family":null,"id":"gpt-oss-safeguard-120b","knowledge":null,"last_updated":"2025-10-29","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"GPT OSS Safeguard 120B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-29","retired":false,"tags":null},"hermes-4-405b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.996,"output":2.989},"deprecated":false,"extra":{"attachment":false,"description":"Hermes 4 405B is a frontier hybrid-mode reasoning model built on Llama 3.1, optimized for advanced logic, math, coding, and structured output generation.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"hermes-4-405b","knowledge":null,"last_updated":"2024-08-13","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"hermes-4-405b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-08-13","retired":false,"tags":null},"hermes-4-70b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.129,"output":0.399},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"hermes-4-70b","knowledge":"2023-12","last_updated":"2025-08-26","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Hermes 4 70B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-08-26","retired":false,"tags":null},"kimi-k2.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.124,"input":0.495,"output":2.768},"deprecated":false,"extra":{"attachment":true,"description":"Kimi reasoning model for long-horizon research, planning, and tool use","family":"kimi-k2","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"toggle"}],"structured_output":true,"temperature":true},"family":null,"id":"kimi-k2.5","knowledge":"2025-01","last_updated":"2026-01","lifecycle":null,"limits":{"context":262144,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K2.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-01","retired":false,"tags":null},"kimi-k2.6":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.193,"input":0.773,"output":3.38},"deprecated":false,"extra":{"attachment":true,"description":"Kimi reasoning model for long-horizon research, planning, and tool use","family":"kimi-k2","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[{"type":"toggle"}],"structured_output":true,"temperature":true},"family":null,"id":"kimi-k2.6","knowledge":"2025-01","last_updated":"2026-04-21","lifecycle":null,"limits":{"context":262144,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K2.6","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-21","retired":false,"tags":null},"kimi-k2.7-code":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.201,"input":0.75,"output":3.5},"deprecated":false,"extra":{"attachment":true,"description":"Coding-focused Kimi model, stronger on long-horizon repo work with less overthinking","family":"kimi-k2","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"kimi-k2.7-code","knowledge":"2025-01","last_updated":"2026-06-12","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K2.7 Code","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-12","retired":false,"tags":null},"kimi-k3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":3,"output":14.999},"deprecated":false,"extra":{"attachment":true,"description":"Multimodal Kimi model with 1M context and toggleable max-effort thinking for long-horizon agent work","family":"kimi-k3","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"kimi-k3","knowledge":null,"last_updated":"2026-07-16","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Kimi K3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-07-16","retired":false,"tags":null},"llama-3.1-405b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.95,"output":1.95},"deprecated":false,"extra":{"attachment":false,"description":"Open Llama instruction model for multilingual chat, reasoning, and coding","family":"llama","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"llama-3.1-405b-instruct","knowledge":"2023-12","last_updated":"2024-07-23","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Llama 3.1 405B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-07-23","retired":false,"tags":null},"llama-3.1-8b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.167,"output":0.167},"deprecated":false,"extra":{"attachment":false,"description":"Optimized for dialogue, this LLM by Meta outperforms other open-source chat models in benchmarks while prioritizing helpfulness and safety.","family":"llama","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"llama-3.1-8b-instruct","knowledge":"2023-12","last_updated":"2024-07-23","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Llama-3.1-8B-Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-07-23","retired":false,"tags":null},"llama-3.1-nemotron-ultra-253b-v1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.598,"output":1.794},"deprecated":false,"extra":{"attachment":false,"description":"A reasoning-optimized LLM based on Llama 3.1, Nemotron Ultra 253B delivers strong performance in tasks like RAG and tool use, with high efficiency and reduced latency.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"llama-3.1-nemotron-ultra-253b-v1","knowledge":null,"last_updated":"2025-04-07","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"llama-3.1-nemotron-ultra-253b-v1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-07","retired":false,"tags":null},"llama-3.3-70b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.129,"output":0.399},"deprecated":false,"extra":{"attachment":false,"description":"Popular open Llama workhorse for multilingual chat, coding, and self-hosting","family":"llama","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"llama-3.3-70b-instruct","knowledge":"2023-12","last_updated":"2024-12-06","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Llama-3.3-70B-Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-12-06","retired":false,"tags":null},"minicpm-v-4.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.651,"output":1.097},"deprecated":false,"extra":{"attachment":false,"description":"MiniCPM-V 4.5 is a compact, high-performance vision-language model excelling in video understanding, OCR, and multimodal reasoning with efficient deployment.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"minicpm-v-4.5","knowledge":null,"last_updated":"2026-06-02","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"minicpm-v-4.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-02","retired":false,"tags":null},"minimax-m2":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.349,"output":1.405},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m2","knowledge":null,"last_updated":"2025-10-27","lifecycle":null,"limits":{"context":400000,"output":400000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-27","retired":false,"tags":null},"minimax-m2.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.359,"output":1.435},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m2.1","knowledge":null,"last_updated":"2025-12-23","lifecycle":null,"limits":{"context":196000,"output":196000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-23","retired":false,"tags":null},"minimax-m2.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.03,"input":0.296,"output":1.087},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","interleaved":{"field":"reasoning_content"},"open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m2.5","knowledge":null,"last_updated":"2026-02-12","lifecycle":null,"limits":{"context":196680,"output":196608},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-12","retired":false,"tags":null},"minimax-m2.7":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.668,"output":2.674},"deprecated":false,"extra":{"attachment":false,"description":"MiniMax model for chat, coding, office work, and agentic tasks","family":"minimax","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m2.7","knowledge":null,"last_updated":"2026-03-18","lifecycle":null,"limits":{"context":196608,"output":196072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"MiniMax-M2.7","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-18","retired":false,"tags":null},"minimax-m3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.099,"input":0.395,"output":1.977},"deprecated":false,"extra":{"attachment":true,"description":"MiniMax multimodal model for long-context coding, perception, and agent planning","family":"minimax","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"minimax-m3","knowledge":null,"last_updated":"2026-06-01","lifecycle":null,"limits":{"context":1048576,"output":1048576},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"MiniMax-M3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-06-01","retired":false,"tags":null},"ministral-14b-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.022,"input":0.223,"output":0.223},"deprecated":false,"extra":{"attachment":true,"description":"Ministral 3 14B is a frontier-level 14B multimodal model optimized for local deployment, delivering state-of-the-art text and vision reasoning with a 256K context window and strong agentic capabilities.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"ministral-14b-2512","knowledge":null,"last_updated":"2025-12-03","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"ministral-14b-2512","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-03","retired":false,"tags":null},"ministral-3b-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.011,"input":0.111,"output":0.111},"deprecated":false,"extra":{"attachment":true,"description":"Ministral 3 3B is a compact, efficient multimodal model with strong language, vision capabilities, and ideal for custom fine-tuning.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"ministral-3b-2512","knowledge":null,"last_updated":"2025-12-03","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"ministral-3b-2512","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-03","retired":false,"tags":null},"ministral-8b-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.017,"input":0.167,"output":0.167},"deprecated":false,"extra":{"attachment":true,"description":"Ministral 3 8B is a balanced, efficient multimodal model offering strong text and vision capabilities, optimized for edge and local deployment.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"ministral-8b-2512","knowledge":null,"last_updated":"2025-12-03","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"ministral-8b-2512","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-03","retired":false,"tags":null},"mistral-7b-instruct-v0.2":{"aliases":[],"base_url":null,"capabilities":null,"catalog_only":true,"cost":{"input":0.159,"output":0.219},"deprecated":false,"extra":{"attachment":false,"description":"Mistral 7B Instruct is a compact, 7B parameter model optimized for fast and efficient text and code generation with a 32K token context window.","open_weights":false,"structured_output":false,"temperature":false},"family":null,"id":"mistral-7b-instruct-v0.2","knowledge":null,"last_updated":"2025-05-26","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"mistral-7b-instruct-v0.2","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-26","retired":false,"tags":null},"mistral-7b-instruct-v0.3":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.111,"output":0.111},"deprecated":false,"extra":{"attachment":false,"description":"Mistral-7B-Instruct-v0.3 model is a fine-tuned version of the Mistral 7B base model, optimized for instruction-following tasks. Released in 2023, it is intended for demonstration purposes and does not include built-in guardrails or moderation features.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"mistral-7b-instruct-v0.3","knowledge":null,"last_updated":"2025-05-26","lifecycle":null,"limits":{"context":127000,"output":127000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"mistral-7b-instruct-v0.3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-26","retired":false,"tags":null},"mistral-large-2402":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":4.284,"output":12.952},"deprecated":false,"extra":{"attachment":false,"description":"Mistral Large (24.02) is Mistral AI’s most advanced language model, built for complex multilingual reasoning, code generation, and deep text understanding.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"mistral-large-2402","knowledge":null,"last_updated":"2025-05-26","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"mistral-large-2402","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-26","retired":false,"tags":null},"mistral-large-2512":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.056,"input":0.557,"output":1.671},"deprecated":false,"extra":{"attachment":true,"description":"Mistral's largest general model for enterprise agents, coding, and multilingual reasoning","family":"mistral-large","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"mistral-large-2512","knowledge":"2024-11","last_updated":"2025-12-02","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Mistral Large 3","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-11-01","retired":false,"tags":null},"mistral-medium-2508":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.045,"input":0.446,"output":2.228},"deprecated":false,"extra":{"attachment":true,"description":"Mistral Medium 2508 is a frontier-class multimodal LLM with a 128,000 token context window, optimized for reasoning, coding, and multimodal tasks.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"mistral-medium-2508","knowledge":null,"last_updated":"2024-08-07","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"mistral-medium-2508","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-08-07","retired":false,"tags":null},"mistral-medium-3.5":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.671,"output":5.57},"deprecated":false,"extra":{"attachment":true,"description":"Mistral Medium 3.5 is a frontier multimodal 128B model combining reasoning, coding, and instruction-following with strong agentic performance and efficient deployment.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"mistral-medium-3.5","knowledge":null,"last_updated":"2026-04-30","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"mistral-medium-3.5","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-30","retired":false,"tags":null},"mistral-nemo-instruct-2407":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.014,"input":0.145,"output":0.145},"deprecated":false,"extra":{"attachment":false,"description":"A 12B parameter, instruct-tuned language model by Mistral AI and NVIDIA, designed for advanced instruction following, multi-turn conversations, and generating text and code across multiple languages.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"mistral-nemo-instruct-2407","knowledge":null,"last_updated":"2024-08-07","lifecycle":null,"limits":{"context":128000,"output":131072},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"mistral-nemo-instruct-2407","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-08-07","retired":false,"tags":null},"mistral-small-2503":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.111,"output":0.334},"deprecated":false,"extra":{"attachment":true,"description":"Combines advanced text and vision capabilities with 24 billion parameters, supporting multilingual tasks and long contexts up to 131k tokens, making it versatile for various applications without sacrificing performance.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"mistral-small-2503","knowledge":null,"last_updated":"2025-03-20","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"mistral-small-2503","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-03-20","retired":false,"tags":null},"mistral-small-2603":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.014,"input":0.143,"output":0.568},"deprecated":false,"extra":{"attachment":true,"description":"Fast Mistral production model for chat, extraction, and cost-sensitive agents","family":"mistral-small","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"mistral-small-2603","knowledge":"2025-06","last_updated":"2026-03-16","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Mistral Small 4","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-03-16","retired":false,"tags":null},"mistral-small-3.2-24b-instruct-2506":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.1,"output":0.312},"deprecated":false,"extra":{"attachment":true,"description":"Mistral-Small-3.2-24B-Instruct-2506 is a 24B parameter instruction-tuned model with enhanced long-context support (128k) and state-of-the-art vision understanding.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"mistral-small-3.2-24b-instruct-2506","knowledge":null,"last_updated":"2025-05-26","lifecycle":null,"limits":{"context":131000,"output":131000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"mistral-small-3.2-24b-instruct-2506","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-26","retired":false,"tags":null},"mixtral-8x7B-instruct-v0.1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"text":true,"tool_calls":false},"tools":{"enabled":false,"parallel":false,"streaming":false,"strict":false}},"catalog_only":true,"cost":{"input":0.488,"output":0.758},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"mixtral-8x7B-instruct-v0.1","knowledge":"2023-09","last_updated":"2023-12-11","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Mixtral 8x7B Instruct v0.1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2023-12-11","retired":false,"tags":null},"nemotron-nano-v2-12b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.24,"output":0.707},"deprecated":false,"extra":{"attachment":true,"description":"NVIDIA Nemotron Nano v2 12B is a 12-billion-parameter multimodal reasoning model designed for advanced video understanding, document intelligence, and visual reasoning, built with a hybrid Transformer-Mamba architecture for high efficiency and low latency.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nemotron-nano-v2-12b","knowledge":null,"last_updated":"2025-10-31","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"nemotron-nano-v2-12b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-10-31","retired":false,"tags":null},"nova-2-lite":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.373,"output":3.144},"deprecated":false,"extra":{"attachment":true,"description":"Nova 2 Lite is an advanced multimodal reasoning model that combines efficiency and performance, delivering reliable AI for agentic workflows and enterprise applications.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nova-2-lite","knowledge":null,"last_updated":"2025-12-04","lifecycle":null,"limits":{"context":1000000,"output":1000000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"nova-2-lite","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-12-04","retired":false,"tags":null},"nova-lite-v1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.069,"output":0.275},"deprecated":false,"extra":{"attachment":true,"description":"Nova Lite is a fast, low-cost multimodal foundation model capable of reasoning over text, images, and video in 200+ languages.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nova-lite-v1","knowledge":null,"last_updated":"2025-04-14","lifecycle":null,"limits":{"context":300000,"output":300000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"nova-lite-v1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"nova-micro-v1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.04,"output":0.159},"deprecated":false,"extra":{"attachment":true,"description":"Nova Micro is a multilingual text-to-text foundation model with strong reasoning capabilities and broad language coverage across 200+ languages.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nova-micro-v1","knowledge":null,"last_updated":"2025-04-14","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"nova-micro-v1","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04-14","retired":false,"tags":null},"nova-pro-v1":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.918,"output":3.671},"deprecated":false,"extra":{"attachment":true,"description":"Flagship model for demanding analysis, coding, and production agent workflows","family":"nova-pro","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"nova-pro-v1","knowledge":"2024-04","last_updated":"2024-12-03","lifecycle":null,"limits":{"context":300000,"output":5000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Nova Pro 1.0","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-12-03","retired":false,"tags":null},"nvidia-nemotron-3-nano-30b-a3b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.06,"output":0.24},"deprecated":false,"extra":{"attachment":false,"description":"Nemotron-Nano-3-30B-A3B is a compact Mixture-of-Experts model optimized for efficient reasoning, chat, and coding, with strong multilingual support and long-context RAG and agent workflows.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nvidia-nemotron-3-nano-30b-a3b","knowledge":null,"last_updated":"2026-01-12","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"nvidia-nemotron-3-nano-30b-a3b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-01-12","retired":false,"tags":null},"nvidia-nemotron-3-nano-omni":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.059,"output":0.237},"deprecated":false,"extra":{"attachment":false,"description":"Nemotron-3-Nano-Omni is an open, efficient omni-modal reasoning model that unifies text, image, audio, and video for agentic AI workflows.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"nvidia-nemotron-3-nano-omni","knowledge":null,"last_updated":"2026-04-29","lifecycle":null,"limits":{"context":300000,"output":300000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"nvidia-nemotron-3-nano-omni","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-29","retired":false,"tags":null},"pixtral-12b-2409":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.223,"output":0.223},"deprecated":false,"extra":{"attachment":true,"description":"Pixtral 2409 12B is a state-of-the-art multimodal model with 12B parameters and a 400M vision encoder, natively trained on interleaved text and image data. It excels in tasks spanning vision-language reasoning, instruction following, and pure text understanding, making it highly effective for real-world multimodal applications.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"pixtral-12b-2409","knowledge":null,"last_updated":"2024-11-09","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"pixtral-12b-2409","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2024-11-09","retired":false,"tags":null},"pixtral-large-2502":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":1.993,"output":5.978},"deprecated":false,"extra":{"attachment":true,"description":"Pixtral Large (25.02) is a 124B open-weight multimodal model built on Mistral Large 2, offering advanced image understanding and strong performance across text and code tasks.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"pixtral-large-2502","knowledge":null,"last_updated":"2025-05-26","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"pixtral-large-2502","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-05-26","retired":false,"tags":null},"qwen2.5-vl-72b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"text":true,"tool_calls":false},"tools":{"enabled":false,"parallel":false,"streaming":false,"strict":false}},"catalog_only":true,"cost":{"input":0.25,"output":0.747},"deprecated":false,"extra":{"attachment":true,"description":"Qwen2.5-VL is a powerful vision-language model with advanced capabilities in visual understanding, long video reasoning, and structured output generation.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"qwen2.5-vl-72b-instruct","knowledge":null,"last_updated":"2025-01-27","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"qwen2.5-vl-72b-instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-01-27","retired":false,"tags":null},"qwen3-235b-a22b-instruct-2507":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.018,"input":0.069,"output":0.455},"deprecated":false,"extra":{"attachment":false,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"qwen3-235b-a22b-instruct-2507","knowledge":null,"last_updated":"2025-07-21","lifecycle":null,"limits":{"context":262000,"output":131000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 235B-A22B Instruct 2507","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-21","retired":false,"tags":null},"qwen3-30b-a3b-instruct-2507":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.099,"output":0.299},"deprecated":false,"extra":{"attachment":false,"description":"Qwen3-30B-A3B-Instruct-2507 is an advanced Mixture-of-Experts model optimized for reasoning, coding, and multilingual instruction following.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"qwen3-30b-a3b-instruct-2507","knowledge":null,"last_updated":"2025-07-28","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"qwen3-30b-a3b-instruct-2507","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-07-28","retired":false,"tags":null},"qwen3-32b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.099,"output":0.299},"deprecated":false,"extra":{"attachment":false,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"qwen3-32b","knowledge":"2025-04","last_updated":"2025-04","lifecycle":null,"limits":{"context":40000,"output":40000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 32B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04","retired":false,"tags":null},"qwen3-coder-30b-a3b-instruct":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.014,"input":0.067,"output":0.245},"deprecated":false,"extra":{"attachment":false,"description":"Smaller Qwen coder for efficient local agents and repo-level fixes","family":"qwen","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"qwen3-coder-30b-a3b-instruct","knowledge":"2025-04","last_updated":"2025-04","lifecycle":null,"limits":{"context":262144,"output":262000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3-Coder 30B-A3B Instruct","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-04","retired":false,"tags":null},"qwen3-coder-next":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.167,"output":0.891},"deprecated":false,"extra":{"attachment":false,"description":"Qwen coding model for software agents, repository edits, and code reasoning","family":"qwen","open_weights":true,"structured_output":true,"temperature":true},"family":null,"id":"qwen3-coder-next","knowledge":"2025-09","last_updated":"2026-02-03","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3 Coder Next","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-03","retired":false,"tags":null},"qwen3-next-80b-a3b-thinking":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.149,"output":1.195},"deprecated":false,"extra":{"attachment":false,"description":"Reasoning model for deliberate analysis, multi-step problem solving, and tool use","family":"qwen","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"qwen3-next-80b-a3b-thinking","knowledge":"2025-04","last_updated":"2025-09","lifecycle":null,"limits":{"context":128000,"output":128000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3-Next 80B-A3B (Thinking)","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2025-09","retired":false,"tags":null},"qwen3-vl-235b-a22b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.052,"input":0.617,"output":3.119},"deprecated":false,"extra":{"attachment":true,"description":"Qwen3 VL 235B A22B is a 235B-parameter MoE vision-language flagship model (≈22B active) designed for frontier-level multimodal understanding across text, images, documents, and long videos.","open_weights":false,"reasoning_options":[],"structured_output":true,"temperature":false},"family":null,"id":"qwen3-vl-235b-a22b","knowledge":null,"last_updated":"2026-01-13","lifecycle":null,"limits":{"context":256000,"output":256000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"qwen3-vl-235b-a22b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-01-13","retired":false,"tags":null},"qwen3.5-122b-a10b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.124,"input":0.495,"output":3.46},"deprecated":false,"extra":{"attachment":false,"description":"Qwen vision-language model for visual reasoning, documents, and agent tasks","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3.5-122b-a10b","knowledge":null,"last_updated":"2026-02-23","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3.5 122B-A10B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-23","retired":false,"tags":null},"qwen3.5-397b-a17b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.668,"output":4.01},"deprecated":false,"extra":{"attachment":false,"description":"Large open Qwen multimodal MoE for visual agents and long technical tasks","family":"qwen","open_weights":true,"reasoning_options":[],"temperature":true},"family":null,"id":"qwen3.5-397b-a17b","knowledge":null,"last_updated":"2026-02-15","lifecycle":null,"limits":{"context":262000,"output":250000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3.5 397B-A17B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-15","retired":false,"tags":null},"qwen3.5-9b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.111,"output":0.167},"deprecated":false,"extra":{"attachment":true,"description":"Qwen instruction model for multilingual chat, reasoning, and tool use","family":"qwen","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"qwen3.5-9b","knowledge":null,"last_updated":"2026-02-23","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Qwen3.5 9B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-23","retired":false,"tags":null},"qwen3.6-27b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.446,"output":3.008},"deprecated":false,"extra":{"attachment":true,"description":"Qwen vision-language model for visual reasoning, documents, and agent tasks","family":"qwen","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"qwen3.6-27b","knowledge":null,"last_updated":"2026-04-22","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Qwen3.6 27B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-22","retired":false,"tags":null},"qwen3.6-35b-a3b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"input":0.167,"output":0.557},"deprecated":false,"extra":{"attachment":true,"description":"Open multimodal Qwen MoE for local agents that need vision, audio, and code","family":"qwen","open_weights":true,"reasoning_options":[],"structured_output":true,"temperature":true},"family":null,"id":"qwen3.6-35b-a3b","knowledge":null,"last_updated":"2026-04-17","lifecycle":null,"limits":{"context":262000,"output":262000},"modalities":{"input":["text","image"],"output":["text"]},"model":null,"name":"Qwen3.6 35B-A3B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-04-17","retired":false,"tags":null},"qwen3.8-2.4t-a95b":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":true},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.625,"input":2.5,"output":6},"deprecated":false,"extra":{"attachment":false,"description":"Open-weight sparse MoE (2.4T total, 95B active), the open-weight twin of Qwen3.8 Max for coding, research, complex reasoning, and agentic workflows","family":"qwen","open_weights":true,"reasoning_options":[{"type":"effort","values":["low","medium","xhigh"]}],"structured_output":true,"temperature":true},"family":null,"id":"qwen3.8-2.4t-a95b","knowledge":null,"last_updated":"2026-08-12","lifecycle":null,"limits":{"context":262144,"output":262144},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"Qwen3.8 2.4T A95B","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-08-12","retired":false,"tags":null},"qwen3guard-gen-0.6b":{"aliases":[],"base_url":null,"capabilities":null,"catalog_only":true,"cost":{"input":0,"output":0},"deprecated":false,"extra":{"attachment":false,"description":"Qwen3Guard-Gen-0.6B is a lightweight multilingual safety moderation model that classifies prompts and responses into safe, controversial, or unsafe categories.","open_weights":false,"structured_output":false,"temperature":false},"family":null,"id":"qwen3guard-gen-0.6b","knowledge":null,"last_updated":"2026-02-04","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"qwen3guard-gen-0.6b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-04","retired":false,"tags":null},"qwen3guard-gen-8b":{"aliases":[],"base_url":null,"capabilities":null,"catalog_only":true,"cost":{"input":0,"output":0},"deprecated":false,"extra":{"attachment":false,"description":"Qwen3Guard-Gen-8B is a large-scale multilingual safety moderation model designed for high-accuracy prompt and response classification.","open_weights":false,"structured_output":false,"temperature":false},"family":null,"id":"qwen3guard-gen-8b","knowledge":null,"last_updated":"2026-02-04","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text"],"output":["text"]},"model":null,"name":"qwen3guard-gen-8b","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-04","retired":false,"tags":null},"voxtral-small-2507":{"aliases":[],"base_url":null,"capabilities":{"chat":true,"embeddings":false,"json":{"native":false,"schema":false,"strict":false},"reasoning":{"enabled":false},"rerank":false,"streaming":{"tool_calls":true},"tools":{"enabled":true}},"catalog_only":true,"cost":{"cache_read":0.011,"input":0.111,"output":0.334},"deprecated":false,"extra":{"attachment":true,"description":"Voxtral Small is a multimodal model with audio input, combining advanced speech capabilities with strong text performance for transcription, translation, and audio understanding.","open_weights":false,"structured_output":true,"temperature":false},"family":null,"id":"voxtral-small-2507","knowledge":null,"last_updated":"2026-02-02","lifecycle":null,"limits":{"context":32000,"output":32000},"modalities":{"input":["text","audio"],"output":["text"]},"model":null,"name":"voxtral-small-2507","pricing":null,"provider":"cortecs","provider_model_id":null,"release_date":"2026-02-02","retired":false,"tags":null}},"name":"Cortecs","pricing_defaults":null};
4
4
  const provider = /* @__PURE__ */ createProviderCatalog(data);
5
5
  export default provider;