cct-cli 0.7.9.0__py3-none-any.whl

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (325) hide show
  1. calc_terminal/__init__.py +14 -0
  2. calc_terminal/__main__.py +14 -0
  3. calc_terminal/activity.py +1334 -0
  4. calc_terminal/agent.py +3387 -0
  5. calc_terminal/agent_runtime.py +519 -0
  6. calc_terminal/ai_context.py +447 -0
  7. calc_terminal/ai_modes.py +752 -0
  8. calc_terminal/ai_personalization.py +286 -0
  9. calc_terminal/ai_preview_feedback.py +213 -0
  10. calc_terminal/aicore.py +2572 -0
  11. calc_terminal/anim.py +367 -0
  12. calc_terminal/app.py +3685 -0
  13. calc_terminal/art.py +639 -0
  14. calc_terminal/atomsim.py +368 -0
  15. calc_terminal/attachments.py +743 -0
  16. calc_terminal/benchmark_system.py +414 -0
  17. calc_terminal/browser/__init__.py +36 -0
  18. calc_terminal/browser/browser_state.py +346 -0
  19. calc_terminal/browser/devserver.py +176 -0
  20. calc_terminal/browser/engine.py +494 -0
  21. calc_terminal/browser/navigation.py +84 -0
  22. calc_terminal/browser/preview.py +429 -0
  23. calc_terminal/browser/preview_entry.py +95 -0
  24. calc_terminal/browser/project_detector.py +144 -0
  25. calc_terminal/browser/server.py +449 -0
  26. calc_terminal/browser/state.py +75 -0
  27. calc_terminal/browser/watcher.py +99 -0
  28. calc_terminal/browser_gui/__init__.py +1 -0
  29. calc_terminal/browser_gui/__main__.py +3 -0
  30. calc_terminal/browser_gui/launcher.py +173 -0
  31. calc_terminal/browser_gui/playwright_browser.py +117 -0
  32. calc_terminal/browser_gui/qt_browser.py +1501 -0
  33. calc_terminal/browser_gui/webview_browser.py +57 -0
  34. calc_terminal/capabilities/__init__.py +35 -0
  35. calc_terminal/capabilities/adapters/__init__.py +33 -0
  36. calc_terminal/capabilities/adapters/bioinformatics.py +204 -0
  37. calc_terminal/capabilities/adapters/browser_adapter.py +205 -0
  38. calc_terminal/capabilities/adapters/filesystem.py +206 -0
  39. calc_terminal/capabilities/adapters/git_adapter.py +202 -0
  40. calc_terminal/capabilities/adapters/jupyter_adapter.py +138 -0
  41. calc_terminal/capabilities/adapters/ml_frameworks.py +158 -0
  42. calc_terminal/capabilities/adapters/platforms.py +200 -0
  43. calc_terminal/capabilities/adapters/python_exec.py +93 -0
  44. calc_terminal/capabilities/adapters/quantum_adapter.py +150 -0
  45. calc_terminal/capabilities/adapters/scientific_comp.py +123 -0
  46. calc_terminal/capabilities/adapters/structural_bio.py +161 -0
  47. calc_terminal/capabilities/adapters/terminal.py +99 -0
  48. calc_terminal/capabilities/bus.py +178 -0
  49. calc_terminal/capabilities/discovery.py +207 -0
  50. calc_terminal/capabilities/schema.py +221 -0
  51. calc_terminal/cat.ico +0 -0
  52. calc_terminal/cat_browser.py +2018 -0
  53. calc_terminal/chat_store.py +703 -0
  54. calc_terminal/cli.py +1178 -0
  55. calc_terminal/code_editor.py +640 -0
  56. calc_terminal/collaboration.py +723 -0
  57. calc_terminal/commands_data.py +139 -0
  58. calc_terminal/compatibility_engine.py +352 -0
  59. calc_terminal/compute/__init__.py +31 -0
  60. calc_terminal/compute/fabric.py +350 -0
  61. calc_terminal/config.py +227 -0
  62. calc_terminal/core/__init__.py +41 -0
  63. calc_terminal/core/checkpoint.py +156 -0
  64. calc_terminal/core/input/__init__.py +45 -0
  65. calc_terminal/core/mode_registry.py +300 -0
  66. calc_terminal/core/project_graph.py +172 -0
  67. calc_terminal/core/recovery.py +129 -0
  68. calc_terminal/core/security_layer.py +112 -0
  69. calc_terminal/core/task_graph.py +202 -0
  70. calc_terminal/core/unified_runtime.py +184 -0
  71. calc_terminal/core/verification.py +257 -0
  72. calc_terminal/customization.py +1566 -0
  73. calc_terminal/derivations.py +153 -0
  74. calc_terminal/device_control.py +263 -0
  75. calc_terminal/diagnostics/__init__.py +27 -0
  76. calc_terminal/diagnostics/doctor_engine.py +382 -0
  77. calc_terminal/diagnostics/self_test.py +247 -0
  78. calc_terminal/doctor.py +519 -0
  79. calc_terminal/easter_eggs.py +274 -0
  80. calc_terminal/editor/__init__.py +1 -0
  81. calc_terminal/editor/actions.py +263 -0
  82. calc_terminal/editor/commands.py +160 -0
  83. calc_terminal/editor/shortcuts.py +226 -0
  84. calc_terminal/engine.py +259 -0
  85. calc_terminal/errors.py +120 -0
  86. calc_terminal/event_stream.py +146 -0
  87. calc_terminal/eventbus.py +133 -0
  88. calc_terminal/extensions.py +733 -0
  89. calc_terminal/fallback_cli.py +1321 -0
  90. calc_terminal/first_run.py +265 -0
  91. calc_terminal/fomoji_auth.py +1043 -0
  92. calc_terminal/formulas.py +82 -0
  93. calc_terminal/fs_cache.py +121 -0
  94. calc_terminal/fs_watcher.py +277 -0
  95. calc_terminal/game.py +193 -0
  96. calc_terminal/gen1.py +5 -0
  97. calc_terminal/generators.py +245 -0
  98. calc_terminal/gestures/__init__.py +42 -0
  99. calc_terminal/gestures/bindings.py +175 -0
  100. calc_terminal/gestures/manager.py +477 -0
  101. calc_terminal/goodbye.py +363 -0
  102. calc_terminal/gpu3d.py +290 -0
  103. calc_terminal/graphs.py +358 -0
  104. calc_terminal/hardware_analyzer.py +440 -0
  105. calc_terminal/host/__init__.py +30 -0
  106. calc_terminal/host/browser_manager.py +187 -0
  107. calc_terminal/host/desktop.py +1386 -0
  108. calc_terminal/host/launcher.py +395 -0
  109. calc_terminal/host/terminal.py +279 -0
  110. calc_terminal/identity.py +216 -0
  111. calc_terminal/input/__init__.py +54 -0
  112. calc_terminal/input/capabilities.py +258 -0
  113. calc_terminal/input/focus.py +87 -0
  114. calc_terminal/input/gestures.py +64 -0
  115. calc_terminal/input/pointer.py +114 -0
  116. calc_terminal/input/touch.py +345 -0
  117. calc_terminal/keys.py +84 -0
  118. calc_terminal/live_automation.py +165 -0
  119. calc_terminal/mathtext.py +433 -0
  120. calc_terminal/mcp.py +386 -0
  121. calc_terminal/memory.py +337 -0
  122. calc_terminal/memory_v2.py +479 -0
  123. calc_terminal/metrics.py +333 -0
  124. calc_terminal/mode_detection.py +146 -0
  125. calc_terminal/model.py +2431 -0
  126. calc_terminal/model_router.py +665 -0
  127. calc_terminal/models/__init__.py +0 -0
  128. calc_terminal/models/active_state.py +187 -0
  129. calc_terminal/models/dynamic_registry.py +584 -0
  130. calc_terminal/models/manager.py +781 -0
  131. calc_terminal/models/model_metadata.json +3526 -0
  132. calc_terminal/models/profiles.py +194 -0
  133. calc_terminal/models/registry.py +265 -0
  134. calc_terminal/models/schema.py +197 -0
  135. calc_terminal/models/validator.py +287 -0
  136. calc_terminal/models/verification_engine.py +368 -0
  137. calc_terminal/native_picker.py +215 -0
  138. calc_terminal/ollama_catalog.py +279 -0
  139. calc_terminal/ollama_download.py +233 -0
  140. calc_terminal/orchestrator.py +304 -0
  141. calc_terminal/package_research.py +322 -0
  142. calc_terminal/packages.py +1024 -0
  143. calc_terminal/pc_specs.py +116 -0
  144. calc_terminal/permissions.py +334 -0
  145. calc_terminal/pet.py +106 -0
  146. calc_terminal/pipeline.py +505 -0
  147. calc_terminal/platform/__init__.py +491 -0
  148. calc_terminal/platform/desktop.py +491 -0
  149. calc_terminal/platform/web.py +781 -0
  150. calc_terminal/preview/__init__.py +1 -0
  151. calc_terminal/preview/dev_server.py +303 -0
  152. calc_terminal/preview/diagnostics.py +131 -0
  153. calc_terminal/preview/live_reload.py +66 -0
  154. calc_terminal/preview/manager.py +129 -0
  155. calc_terminal/project_stats.py +209 -0
  156. calc_terminal/projects.py +328 -0
  157. calc_terminal/providers/__init__.py +0 -0
  158. calc_terminal/providers/adapters/__init__.py +80 -0
  159. calc_terminal/providers/adapters/anthropic_adapter.py +127 -0
  160. calc_terminal/providers/adapters/base.py +106 -0
  161. calc_terminal/providers/adapters/chinese_adapters.py +420 -0
  162. calc_terminal/providers/adapters/gemini_adapter.py +101 -0
  163. calc_terminal/providers/adapters/ollama_adapter.py +83 -0
  164. calc_terminal/providers/adapters/openai_adapter.py +159 -0
  165. calc_terminal/providers/adapters/other_adapters.py +246 -0
  166. calc_terminal/providers/anthropic_provider.py +172 -0
  167. calc_terminal/providers/auto_update.py +416 -0
  168. calc_terminal/providers/base_provider.py +105 -0
  169. calc_terminal/providers/discovery_manager.py +207 -0
  170. calc_terminal/providers/gemini_provider.py +178 -0
  171. calc_terminal/providers/lifecycle.py +767 -0
  172. calc_terminal/providers/ollama_adapter.py +707 -0
  173. calc_terminal/providers/openai_provider.py +254 -0
  174. calc_terminal/providers/provider_manager.py +1827 -0
  175. calc_terminal/providers/providers.json +4075 -0
  176. calc_terminal/reactionsim.py +279 -0
  177. calc_terminal/registry.py +337 -0
  178. calc_terminal/report.py +162 -0
  179. calc_terminal/research/__init__.py +45 -0
  180. calc_terminal/research/artifact_intel.py +126 -0
  181. calc_terminal/research/data_lineage.py +123 -0
  182. calc_terminal/research/experiment_ledger.py +303 -0
  183. calc_terminal/research/reproducibility.py +131 -0
  184. calc_terminal/resilience/__init__.py +47 -0
  185. calc_terminal/resilience/agent_state.py +121 -0
  186. calc_terminal/resilience/capability_matcher.py +174 -0
  187. calc_terminal/resilience/circuit_breaker.py +158 -0
  188. calc_terminal/resilience/failover_engine.py +230 -0
  189. calc_terminal/resilience/health_monitor.py +192 -0
  190. calc_terminal/resilience/ollama_adapter.py +125 -0
  191. calc_terminal/resilience/orchestrator.py +312 -0
  192. calc_terminal/resilience/types.py +134 -0
  193. calc_terminal/sandbox.py +98 -0
  194. calc_terminal/scires.py +558 -0
  195. calc_terminal/security_scanner.py +126 -0
  196. calc_terminal/session.py +294 -0
  197. calc_terminal/sim3d.py +206 -0
  198. calc_terminal/solver.py +276 -0
  199. calc_terminal/sound.py +127 -0
  200. calc_terminal/task_reports.py +287 -0
  201. calc_terminal/terminal_host.py +201 -0
  202. calc_terminal/terminal_identity.py +411 -0
  203. calc_terminal/test_ai_mode_reliability.py +344 -0
  204. calc_terminal/test_browser.py +368 -0
  205. calc_terminal/test_code_editor_upgrade.py +485 -0
  206. calc_terminal/test_customization.py +1148 -0
  207. calc_terminal/test_customization_ui.py +612 -0
  208. calc_terminal/test_dynamic_registry.py +304 -0
  209. calc_terminal/test_extensions.py +436 -0
  210. calc_terminal/test_overhaul.py +557 -0
  211. calc_terminal/test_project_detect.py +255 -0
  212. calc_terminal/test_root_cause_fix.py +527 -0
  213. calc_terminal/test_stability.py +532 -0
  214. calc_terminal/test_terminal_identity.py +132 -0
  215. calc_terminal/test_v079_speed.py +460 -0
  216. calc_terminal/theme.py +1107 -0
  217. calc_terminal/timeline.py +139 -0
  218. calc_terminal/todos.py +246 -0
  219. calc_terminal/tool_call_normalizer.py +419 -0
  220. calc_terminal/tui.py +104 -0
  221. calc_terminal/ui/__init__.py +8 -0
  222. calc_terminal/ui/activity_panel.py +231 -0
  223. calc_terminal/ui/activity_stream_panel.py +238 -0
  224. calc_terminal/ui/animations.py +122 -0
  225. calc_terminal/ui/app.py +7271 -0
  226. calc_terminal/ui/attach_panel.py +597 -0
  227. calc_terminal/ui/attachments.py +424 -0
  228. calc_terminal/ui/backup_panel.py +810 -0
  229. calc_terminal/ui/browser_shell.py +887 -0
  230. calc_terminal/ui/cat_agent.py +357 -0
  231. calc_terminal/ui/chats_panel.py +899 -0
  232. calc_terminal/ui/command_palette.py +125 -0
  233. calc_terminal/ui/command_palette_modal.py +166 -0
  234. calc_terminal/ui/composer.py +1141 -0
  235. calc_terminal/ui/context_menu.py +197 -0
  236. calc_terminal/ui/conversation.py +1435 -0
  237. calc_terminal/ui/customization_panel.py +1229 -0
  238. calc_terminal/ui/dashboard.py +404 -0
  239. calc_terminal/ui/design_system.py +557 -0
  240. calc_terminal/ui/diff_panel.py +213 -0
  241. calc_terminal/ui/editor.py +2102 -0
  242. calc_terminal/ui/empty_state.py +302 -0
  243. calc_terminal/ui/events.py +487 -0
  244. calc_terminal/ui/extensions_panel.py +815 -0
  245. calc_terminal/ui/footer.py +166 -0
  246. calc_terminal/ui/gestures_panel.py +383 -0
  247. calc_terminal/ui/goodbye_screen.py +100 -0
  248. calc_terminal/ui/header.py +1034 -0
  249. calc_terminal/ui/help_panel.py +254 -0
  250. calc_terminal/ui/live_activities.py +914 -0
  251. calc_terminal/ui/mcp_panel.py +570 -0
  252. calc_terminal/ui/memory_center.py +524 -0
  253. calc_terminal/ui/mode_colors_panel.py +525 -0
  254. calc_terminal/ui/nav_screens.py +747 -0
  255. calc_terminal/ui/ollama_panel.py +536 -0
  256. calc_terminal/ui/palette.py +221 -0
  257. calc_terminal/ui/permission_panel.py +269 -0
  258. calc_terminal/ui/personalization_panel.py +517 -0
  259. calc_terminal/ui/personalize_center.py +1568 -0
  260. calc_terminal/ui/preview_panel.py +441 -0
  261. calc_terminal/ui/resizers.py +402 -0
  262. calc_terminal/ui/sidebar.py +1285 -0
  263. calc_terminal/ui/statusbar.py +168 -0
  264. calc_terminal/ui/theme_css.py +1396 -0
  265. calc_terminal/ui/thinking.py +226 -0
  266. calc_terminal/ui/timeline_panel.py +102 -0
  267. calc_terminal/ui/todo_panel.py +193 -0
  268. calc_terminal/ui/viewport.py +136 -0
  269. calc_terminal/ui/vision_panel.py +489 -0
  270. calc_terminal/ui/welcome_modal.py +343 -0
  271. calc_terminal/ui/widgets.py +160 -0
  272. calc_terminal/ui/workspace.py +831 -0
  273. calc_terminal/viewers/__init__.py +1 -0
  274. calc_terminal/viewers/document_viewer.py +252 -0
  275. calc_terminal/viewers/image_viewer.py +241 -0
  276. calc_terminal/viewers/pdf_viewer.py +203 -0
  277. calc_terminal/viewers/presentation_viewer.py +164 -0
  278. calc_terminal/viewers/registry.py +120 -0
  279. calc_terminal/viewers/spreadsheet_viewer.py +204 -0
  280. calc_terminal/vision/__init__.py +89 -0
  281. calc_terminal/vision/analysis.py +194 -0
  282. calc_terminal/vision/annotations.py +297 -0
  283. calc_terminal/vision/capture.py +171 -0
  284. calc_terminal/vision/context.py +231 -0
  285. calc_terminal/vision/correlation.py +169 -0
  286. calc_terminal/vision/cursor.py +258 -0
  287. calc_terminal/vision/events.py +66 -0
  288. calc_terminal/vision/frame_pipeline.py +259 -0
  289. calc_terminal/vision/priority.py +218 -0
  290. calc_terminal/vision/provider.py +180 -0
  291. calc_terminal/vision/safety.py +149 -0
  292. calc_terminal/vision/session.py +281 -0
  293. calc_terminal/vision/verify.py +162 -0
  294. calc_terminal/vision.py +514 -0
  295. calc_terminal/vscode_integration.py +113 -0
  296. calc_terminal/web/__init__.py +8 -0
  297. calc_terminal/web/cat_runtime.py +710 -0
  298. calc_terminal/web/server.py +2891 -0
  299. calc_terminal/web/static/css/app.css +3152 -0
  300. calc_terminal/web/static/icons/badge-72.png +0 -0
  301. calc_terminal/web/static/icons/cat.ico +0 -0
  302. calc_terminal/web/static/icons/icon-128.png +0 -0
  303. calc_terminal/web/static/icons/icon-144.png +0 -0
  304. calc_terminal/web/static/icons/icon-152.png +0 -0
  305. calc_terminal/web/static/icons/icon-192.png +0 -0
  306. calc_terminal/web/static/icons/icon-384.png +0 -0
  307. calc_terminal/web/static/icons/icon-512.png +0 -0
  308. calc_terminal/web/static/icons/icon-72.png +0 -0
  309. calc_terminal/web/static/icons/icon-96.png +0 -0
  310. calc_terminal/web/static/icons/icon.svg +34 -0
  311. calc_terminal/web/static/icons/new-project.png +0 -0
  312. calc_terminal/web/static/icons/open-project.png +0 -0
  313. calc_terminal/web/static/index.html +734 -0
  314. calc_terminal/web/static/js/app.js +2403 -0
  315. calc_terminal/web/static/manifest.json +88 -0
  316. calc_terminal/web/static/sw.js +230 -0
  317. calc_terminal/workflow_engine.py +769 -0
  318. calc_terminal/workspace.py +593 -0
  319. calc_terminal/workspace_index.py +385 -0
  320. cct_cli-0.7.9.0.dist-info/METADATA +210 -0
  321. cct_cli-0.7.9.0.dist-info/RECORD +325 -0
  322. cct_cli-0.7.9.0.dist-info/WHEEL +5 -0
  323. cct_cli-0.7.9.0.dist-info/entry_points.txt +4 -0
  324. cct_cli-0.7.9.0.dist-info/licenses/LICENSE +21 -0
  325. cct_cli-0.7.9.0.dist-info/top_level.txt +1 -0
@@ -0,0 +1,4075 @@
1
+ {
2
+ "providers": [
3
+ {
4
+ "id": "yi",
5
+ "name": "01.AI (Yi)",
6
+ "country": "China",
7
+ "api_endpoint": "https://api.lingyiwanwu.com/v1",
8
+ "api_style": "openai",
9
+ "supports_model_listing": true,
10
+ "supports_streaming": true,
11
+ "supports_vision": true,
12
+ "supports_embeddings": false,
13
+ "supports_audio": false,
14
+ "supports_reasoning": false,
15
+ "openai_compatible": true,
16
+ "free_models_available": true,
17
+ "paid_models_available": true,
18
+ "needs_key": true,
19
+ "description": "Yi series — 01.AI's flagship model family",
20
+ "last_updated": "2026-07-31T05:17:14+00:00",
21
+ "fallback_models": [
22
+ "yi-large",
23
+ "yi-medium",
24
+ "yi-vision",
25
+ "yi-lightning",
26
+ "yi-spark",
27
+ "yi-34b-chat"
28
+ ],
29
+ "default_model": "yi-large"
30
+ },
31
+ {
32
+ "id": "abacus",
33
+ "name": "Abacus AI",
34
+ "country": "United States",
35
+ "api_endpoint": "https://api.abacus.ai/v1",
36
+ "api_style": "openai",
37
+ "supports_model_listing": true,
38
+ "supports_streaming": true,
39
+ "supports_vision": false,
40
+ "supports_embeddings": false,
41
+ "supports_audio": false,
42
+ "supports_reasoning": false,
43
+ "openai_compatible": true,
44
+ "free_models_available": true,
45
+ "paid_models_available": true,
46
+ "needs_key": true,
47
+ "description": "Enterprise AI platform",
48
+ "last_updated": "2026-07-31T05:17:14+00:00",
49
+ "fallback_models": [
50
+ "claude-3.5",
51
+ "gpt-4o"
52
+ ],
53
+ "default_model": "gpt-4o"
54
+ },
55
+ {
56
+ "id": "ai21",
57
+ "name": "AI21 Labs",
58
+ "country": "Israel",
59
+ "api_endpoint": "https://api.ai21.com/studio/v1",
60
+ "api_style": "openai",
61
+ "supports_model_listing": true,
62
+ "supports_streaming": true,
63
+ "supports_vision": false,
64
+ "supports_embeddings": false,
65
+ "supports_audio": false,
66
+ "supports_reasoning": false,
67
+ "openai_compatible": true,
68
+ "free_models_available": true,
69
+ "paid_models_available": true,
70
+ "needs_key": true,
71
+ "description": "AI21 Jamba — Hybrid SSM-Transformer",
72
+ "last_updated": "2026-09-03T00:00:00+00:00",
73
+ "fallback_models": [
74
+ "jamba-1.5-large",
75
+ "jamba-1.5-mini",
76
+ "jamba-instruct"
77
+ ],
78
+ "default_model": "jamba-1.5-large"
79
+ },
80
+ {
81
+ "id": "airtop",
82
+ "name": "Airtop",
83
+ "country": "United States",
84
+ "api_endpoint": "https://api.airtop.ai/v1",
85
+ "api_style": "openai",
86
+ "supports_model_listing": true,
87
+ "supports_streaming": true,
88
+ "supports_vision": false,
89
+ "supports_embeddings": false,
90
+ "supports_audio": false,
91
+ "supports_reasoning": false,
92
+ "openai_compatible": true,
93
+ "free_models_available": true,
94
+ "paid_models_available": true,
95
+ "needs_key": true,
96
+ "description": "AI browser automation",
97
+ "last_updated": "2026-07-31T05:17:14+00:00",
98
+ "fallback_models": [
99
+ "default"
100
+ ],
101
+ "default_model": "default"
102
+ },
103
+ {
104
+ "id": "aisera",
105
+ "name": "Aisera",
106
+ "country": "United States",
107
+ "api_endpoint": "https://api.aisera.com/v1",
108
+ "api_style": "openai",
109
+ "supports_model_listing": true,
110
+ "supports_streaming": true,
111
+ "supports_vision": false,
112
+ "supports_embeddings": false,
113
+ "supports_audio": false,
114
+ "supports_reasoning": false,
115
+ "openai_compatible": true,
116
+ "free_models_available": true,
117
+ "paid_models_available": true,
118
+ "needs_key": true,
119
+ "description": "Enterprise AI service desk",
120
+ "last_updated": "2026-07-31T05:17:14+00:00",
121
+ "fallback_models": [
122
+ "aisera-llm"
123
+ ],
124
+ "default_model": "aisera-llm"
125
+ },
126
+ {
127
+ "id": "akamai",
128
+ "name": "Akamai (Cloudlet)",
129
+ "country": "United States",
130
+ "api_endpoint": "https://api.cloudlet.ai/v1",
131
+ "api_style": "openai",
132
+ "supports_model_listing": true,
133
+ "supports_streaming": true,
134
+ "supports_vision": false,
135
+ "supports_embeddings": false,
136
+ "supports_audio": false,
137
+ "supports_reasoning": false,
138
+ "openai_compatible": true,
139
+ "free_models_available": true,
140
+ "paid_models_available": true,
141
+ "needs_key": true,
142
+ "description": "Edge AI inference",
143
+ "last_updated": "2026-07-31T05:17:14+00:00",
144
+ "fallback_models": [
145
+ "default"
146
+ ],
147
+ "default_model": "default"
148
+ },
149
+ {
150
+ "id": "aleph-alpha",
151
+ "name": "Aleph Alpha",
152
+ "country": "Germany",
153
+ "api_endpoint": "https://api.aleph-alpha.com",
154
+ "api_style": "openai",
155
+ "supports_model_listing": true,
156
+ "supports_streaming": true,
157
+ "supports_vision": false,
158
+ "supports_embeddings": true,
159
+ "supports_audio": false,
160
+ "supports_reasoning": false,
161
+ "openai_compatible": true,
162
+ "free_models_available": true,
163
+ "paid_models_available": true,
164
+ "needs_key": true,
165
+ "description": "Luminous series — European LLM",
166
+ "last_updated": "2026-07-31T05:17:14+00:00",
167
+ "fallback_models": [
168
+ "luminous-base",
169
+ "luminous-extended",
170
+ "luminous-supreme"
171
+ ],
172
+ "default_model": "luminous-base"
173
+ },
174
+ {
175
+ "id": "alibaba-qwen",
176
+ "name": "Alibaba Qwen",
177
+ "country": "China",
178
+ "api_endpoint": "https://dashscope.aliyuncs.com/compatible-mode/v1",
179
+ "api_style": "openai",
180
+ "supports_model_listing": true,
181
+ "supports_streaming": true,
182
+ "supports_vision": true,
183
+ "supports_embeddings": true,
184
+ "supports_audio": false,
185
+ "supports_reasoning": true,
186
+ "openai_compatible": true,
187
+ "free_models_available": true,
188
+ "paid_models_available": true,
189
+ "needs_key": true,
190
+ "description": "Qwen 2.5 & 3 series — Alibaba's flagship LLMs",
191
+ "last_updated": "2026-09-03T00:00:00+00:00",
192
+ "fallback_models": [
193
+ "qwen-max",
194
+ "qwen-plus",
195
+ "qwen-turbo",
196
+ "qwen2.5-72b-instruct",
197
+ "qwen2.5-32b-instruct",
198
+ "qwen2.5-14b-instruct",
199
+ "qwen-vl-max",
200
+ "qwen3-235b-a22b"
201
+ ],
202
+ "default_model": "qwen-max"
203
+ },
204
+ {
205
+ "id": "ai2",
206
+ "name": "Allen Institute for AI (OLMo)",
207
+ "country": "United States",
208
+ "api_endpoint": "https://api.allenai.org/v1",
209
+ "api_style": "openai",
210
+ "supports_model_listing": true,
211
+ "supports_streaming": true,
212
+ "supports_vision": true,
213
+ "supports_embeddings": true,
214
+ "supports_audio": false,
215
+ "supports_reasoning": true,
216
+ "openai_compatible": true,
217
+ "free_models_available": true,
218
+ "paid_models_available": false,
219
+ "needs_key": true,
220
+ "description": "Open Language Model (OLMo) and Molmo multimodal open-science models",
221
+ "last_updated": "2026-09-01T00:00:00+00:00",
222
+ "fallback_models": [
223
+ "olmo-2-13b-instruct",
224
+ "olmo-7b-instruct",
225
+ "molmo-72b-instruct"
226
+ ],
227
+ "default_model": "olmo-2-13b-instruct"
228
+ },
229
+ {
230
+ "id": "aws-bedrock",
231
+ "name": "Amazon Web Services (Bedrock)",
232
+ "country": "United States",
233
+ "api_endpoint": "https://bedrock-runtime.us-east-1.amazonaws.com",
234
+ "api_style": "openai",
235
+ "supports_model_listing": true,
236
+ "supports_streaming": true,
237
+ "supports_vision": false,
238
+ "supports_embeddings": true,
239
+ "supports_audio": false,
240
+ "supports_reasoning": false,
241
+ "openai_compatible": true,
242
+ "free_models_available": true,
243
+ "paid_models_available": true,
244
+ "needs_key": true,
245
+ "description": "AWS Bedrock — managed foundation models",
246
+ "last_updated": "2026-07-31T05:17:14+00:00",
247
+ "fallback_models": [
248
+ "claude-3-5-sonnet-20241022",
249
+ "claude-3-haiku",
250
+ "llama3-70b",
251
+ "mistral-large",
252
+ "anthropic.claude-3-5-sonnet-20241022-v2:0",
253
+ "anthropic.claude-3-opus-20240229-v1:0",
254
+ "meta.llama3-1-70b-instruct-v1:0",
255
+ "amazon.titan-text-express-v1"
256
+ ],
257
+ "default_model": "claude-3-5-sonnet-20241022"
258
+ },
259
+ {
260
+ "id": "anthropic",
261
+ "name": "Anthropic (Claude)",
262
+ "country": "United States",
263
+ "api_endpoint": "https://api.anthropic.com/v1",
264
+ "api_style": "anthropic",
265
+ "supports_model_listing": true,
266
+ "supports_streaming": true,
267
+ "supports_vision": true,
268
+ "supports_embeddings": false,
269
+ "supports_audio": false,
270
+ "supports_reasoning": true,
271
+ "openai_compatible": false,
272
+ "free_models_available": true,
273
+ "paid_models_available": true,
274
+ "needs_key": true,
275
+ "description": "Claude 4 Fable & Sonnet — latest AI models",
276
+ "last_updated": "2026-09-03T00:00:00+00:00",
277
+ "fallback_models": [
278
+ "claude-2.1",
279
+ "claude-3-5-haiku-20241022",
280
+ "claude-3-5-sonnet-20240620",
281
+ "claude-3-5-sonnet-20241022",
282
+ "claude-3-haiku-20240307",
283
+ "claude-3-opus-20240229",
284
+ "claude-3-sonnet-20240229",
285
+ "claude-opus-4",
286
+ "claude-opus-4-20250514",
287
+ "claude-sonnet-4",
288
+ "claude-sonnet-4-20250514",
289
+ "claude-3-7-sonnet-20250219",
290
+ "claude-3-5-haiku-latest",
291
+ "claude-4-fable",
292
+ "claude-4-fable-5.1"
293
+ ],
294
+ "default_model": "claude-4-fable-5.1"
295
+ },
296
+ {
297
+ "id": "anyscale",
298
+ "name": "Anyscale",
299
+ "country": "United States",
300
+ "api_endpoint": "https://api.endpoints.anyscale.com/v1",
301
+ "api_style": "openai",
302
+ "supports_model_listing": true,
303
+ "supports_streaming": true,
304
+ "supports_vision": false,
305
+ "supports_embeddings": false,
306
+ "supports_audio": false,
307
+ "supports_reasoning": false,
308
+ "openai_compatible": true,
309
+ "free_models_available": true,
310
+ "paid_models_available": true,
311
+ "needs_key": true,
312
+ "description": "Ray-based model serving",
313
+ "last_updated": "2026-07-31T05:17:14+00:00",
314
+ "fallback_models": [
315
+ "llama3-70b",
316
+ "mixtral-8x7b",
317
+ "meta-llama/Meta-Llama-3.1-70B-Instruct",
318
+ "mistralai/Mixtral-8x22B-Instruct-v0.1"
319
+ ],
320
+ "default_model": "llama3-70b"
321
+ },
322
+ {
323
+ "id": "aphrodite_hosted",
324
+ "name": "Aphrodite Engine (Local / Remote)",
325
+ "country": "International",
326
+ "api_endpoint": "http://localhost:2242/v1",
327
+ "api_style": "openai",
328
+ "supports_model_listing": true,
329
+ "supports_streaming": true,
330
+ "supports_vision": false,
331
+ "supports_embeddings": false,
332
+ "supports_audio": false,
333
+ "supports_reasoning": true,
334
+ "openai_compatible": true,
335
+ "free_models_available": true,
336
+ "paid_models_available": false,
337
+ "needs_key": false,
338
+ "description": "High-throughput LLM serving engine based on vLLM, optimized for creative writing",
339
+ "last_updated": "2026-09-01T00:00:00+00:00",
340
+ "fallback_models": [
341
+ "default",
342
+ "pygmalion-2-13b",
343
+ "mythomax-l2-13b"
344
+ ],
345
+ "default_model": "default"
346
+ },
347
+ {
348
+ "id": "assemblyai",
349
+ "name": "AssemblyAI",
350
+ "country": "United States",
351
+ "api_endpoint": "https://api.assemblyai.com/v2",
352
+ "api_style": "openai",
353
+ "supports_model_listing": true,
354
+ "supports_streaming": false,
355
+ "supports_vision": false,
356
+ "supports_embeddings": false,
357
+ "supports_audio": true,
358
+ "supports_reasoning": false,
359
+ "openai_compatible": false,
360
+ "free_models_available": true,
361
+ "paid_models_available": true,
362
+ "needs_key": true,
363
+ "description": "Speech-to-text API",
364
+ "last_updated": "2026-07-31T05:17:14+00:00",
365
+ "fallback_models": [
366
+ "default"
367
+ ],
368
+ "default_model": "default"
369
+ },
370
+ {
371
+ "id": "azure-openai",
372
+ "name": "Azure OpenAI",
373
+ "country": "United States",
374
+ "api_endpoint": "https://api.openai.azure.com",
375
+ "api_style": "openai",
376
+ "supports_model_listing": true,
377
+ "supports_streaming": true,
378
+ "supports_vision": true,
379
+ "supports_embeddings": true,
380
+ "supports_audio": true,
381
+ "supports_reasoning": true,
382
+ "openai_compatible": true,
383
+ "free_models_available": false,
384
+ "paid_models_available": true,
385
+ "needs_key": true,
386
+ "description": "Azure OpenAI — Enterprise OpenAI models",
387
+ "last_updated": "2026-09-03T00:00:00+00:00",
388
+ "fallback_models": [
389
+ "gpt-4o",
390
+ "gpt-4o-mini",
391
+ "gpt-4-turbo",
392
+ "o3-mini"
393
+ ],
394
+ "default_model": "gpt-4o"
395
+ },
396
+ {
397
+ "id": "baichuan",
398
+ "name": "Baichuan",
399
+ "country": "China",
400
+ "api_endpoint": "https://api.baichuan-ai.com/v1",
401
+ "api_style": "openai",
402
+ "supports_model_listing": true,
403
+ "supports_streaming": true,
404
+ "supports_vision": false,
405
+ "supports_embeddings": false,
406
+ "supports_audio": false,
407
+ "supports_reasoning": false,
408
+ "openai_compatible": true,
409
+ "free_models_available": true,
410
+ "paid_models_available": true,
411
+ "needs_key": true,
412
+ "description": "Baichuan 4 — Baichuan's flagship LLM",
413
+ "last_updated": "2026-09-03T00:00:00+00:00",
414
+ "fallback_models": [
415
+ "baichuan4",
416
+ "baichuan3-turbo",
417
+ "baichuan2-turbo"
418
+ ],
419
+ "default_model": "baichuan4"
420
+ },
421
+ {
422
+ "id": "baidu",
423
+ "name": "Baidu (ERNIE)",
424
+ "country": "China",
425
+ "api_endpoint": "https://aip.baidubce.com/rpc/2.0/ai_custom/v1/wenxinworkshop",
426
+ "api_style": "openai",
427
+ "supports_model_listing": true,
428
+ "supports_streaming": true,
429
+ "supports_vision": false,
430
+ "supports_embeddings": false,
431
+ "supports_audio": false,
432
+ "supports_reasoning": false,
433
+ "openai_compatible": true,
434
+ "free_models_available": true,
435
+ "paid_models_available": true,
436
+ "needs_key": true,
437
+ "description": "ERNIE series — Baidu's flagship LLM",
438
+ "last_updated": "2026-07-31T05:17:14+00:00",
439
+ "fallback_models": [
440
+ "ernie-3.5",
441
+ "ernie-4.0",
442
+ "ernie-speed",
443
+ "ernie-4.0-8k",
444
+ "ernie-3.5-8k",
445
+ "ernie-speed-128k",
446
+ "ernie-lite-8k"
447
+ ],
448
+ "default_model": "ernie-3.5"
449
+ },
450
+ {
451
+ "id": "qianfan",
452
+ "name": "Baidu Qianfan",
453
+ "country": "China",
454
+ "api_endpoint": "https://qianfan.baidubce.com/v2",
455
+ "api_style": "openai",
456
+ "supports_model_listing": true,
457
+ "supports_streaming": true,
458
+ "supports_vision": false,
459
+ "supports_embeddings": false,
460
+ "supports_audio": false,
461
+ "supports_reasoning": false,
462
+ "openai_compatible": true,
463
+ "free_models_available": true,
464
+ "paid_models_available": true,
465
+ "needs_key": true,
466
+ "description": "Baidu's enterprise AI platform",
467
+ "last_updated": "2026-07-31T05:17:14+00:00",
468
+ "fallback_models": [
469
+ "ernie-4.0",
470
+ "ernie-speed"
471
+ ],
472
+ "default_model": "ernie-4.0"
473
+ },
474
+ {
475
+ "id": "banana",
476
+ "name": "Banana",
477
+ "country": "United States",
478
+ "api_endpoint": "https://api.banana.dev/v1",
479
+ "api_style": "openai",
480
+ "supports_model_listing": true,
481
+ "supports_streaming": true,
482
+ "supports_vision": false,
483
+ "supports_embeddings": false,
484
+ "supports_audio": false,
485
+ "supports_reasoning": false,
486
+ "openai_compatible": true,
487
+ "free_models_available": true,
488
+ "paid_models_available": true,
489
+ "needs_key": true,
490
+ "description": "Serverless GPU inference",
491
+ "last_updated": "2026-07-31T05:17:14+00:00",
492
+ "fallback_models": [
493
+ "llama3-70b"
494
+ ],
495
+ "default_model": "llama3-70b"
496
+ },
497
+ {
498
+ "id": "baseten",
499
+ "name": "Baseten",
500
+ "country": "United States",
501
+ "api_endpoint": "https://bridge.baseten.co/v1",
502
+ "api_style": "openai",
503
+ "supports_model_listing": true,
504
+ "supports_streaming": true,
505
+ "supports_vision": false,
506
+ "supports_embeddings": false,
507
+ "supports_audio": false,
508
+ "supports_reasoning": false,
509
+ "openai_compatible": true,
510
+ "free_models_available": true,
511
+ "paid_models_available": true,
512
+ "needs_key": true,
513
+ "description": "Serverless GPU inference",
514
+ "last_updated": "2026-07-31T05:17:14+00:00",
515
+ "fallback_models": [
516
+ "llama3-70b",
517
+ "mixtral-8x7b"
518
+ ],
519
+ "default_model": "llama3-70b"
520
+ },
521
+ {
522
+ "id": "beam",
523
+ "name": "Beam",
524
+ "country": "United States",
525
+ "api_endpoint": "https://api.beam.cloud/v1",
526
+ "api_style": "openai",
527
+ "supports_model_listing": true,
528
+ "supports_streaming": true,
529
+ "supports_vision": false,
530
+ "supports_embeddings": false,
531
+ "supports_audio": false,
532
+ "supports_reasoning": false,
533
+ "openai_compatible": true,
534
+ "free_models_available": true,
535
+ "paid_models_available": true,
536
+ "needs_key": true,
537
+ "description": "Serverless GPU cloud",
538
+ "last_updated": "2026-07-31T05:17:14+00:00",
539
+ "fallback_models": [
540
+ "llama3-70b"
541
+ ],
542
+ "default_model": "llama3-70b"
543
+ },
544
+ {
545
+ "id": "bento_ml",
546
+ "name": "BentoCloud (BentoML)",
547
+ "country": "United States",
548
+ "api_endpoint": "https://cloud.bentoml.com/v1",
549
+ "api_style": "openai",
550
+ "supports_model_listing": true,
551
+ "supports_streaming": true,
552
+ "supports_vision": true,
553
+ "supports_embeddings": true,
554
+ "supports_audio": false,
555
+ "supports_reasoning": true,
556
+ "openai_compatible": true,
557
+ "free_models_available": false,
558
+ "paid_models_available": true,
559
+ "needs_key": true,
560
+ "description": "Production inference platform for fast deployment of open LLMs and custom models",
561
+ "last_updated": "2026-09-01T00:00:00+00:00",
562
+ "fallback_models": [
563
+ "llama3-70b-instruct",
564
+ "mistral-7b-instruct",
565
+ "qwen2.5-32b-instruct"
566
+ ],
567
+ "default_model": "llama3-70b-instruct"
568
+ },
569
+ {
570
+ "id": "bioage",
571
+ "name": "BioAge Labs",
572
+ "country": "United States",
573
+ "api_endpoint": "https://api.bioage.ai/v1",
574
+ "api_style": "openai",
575
+ "supports_model_listing": true,
576
+ "supports_streaming": true,
577
+ "supports_vision": false,
578
+ "supports_embeddings": false,
579
+ "supports_audio": false,
580
+ "supports_reasoning": false,
581
+ "openai_compatible": true,
582
+ "free_models_available": true,
583
+ "paid_models_available": true,
584
+ "needs_key": true,
585
+ "description": "Biology AI research",
586
+ "last_updated": "2026-07-31T05:17:14+00:00",
587
+ "fallback_models": [
588
+ "default"
589
+ ],
590
+ "default_model": "default"
591
+ },
592
+ {
593
+ "id": "black-forest-labs",
594
+ "name": "Black Forest Labs",
595
+ "country": "Germany",
596
+ "api_endpoint": "https://api.black-forest-labs.com/v1",
597
+ "api_style": "openai",
598
+ "supports_model_listing": true,
599
+ "supports_streaming": false,
600
+ "supports_vision": true,
601
+ "supports_embeddings": false,
602
+ "supports_audio": false,
603
+ "supports_reasoning": false,
604
+ "openai_compatible": false,
605
+ "free_models_available": true,
606
+ "paid_models_available": true,
607
+ "needs_key": true,
608
+ "description": "FLUX series — image generation",
609
+ "last_updated": "2026-07-31T05:17:14+00:00",
610
+ "fallback_models": [
611
+ "flux-dev",
612
+ "flux-pro"
613
+ ],
614
+ "default_model": "flux-dev"
615
+ },
616
+ {
617
+ "id": "bytedance-doubao",
618
+ "name": "ByteDance Doubao",
619
+ "country": "China",
620
+ "api_endpoint": "https://ark.cn-beijing.volces.com/api/v3",
621
+ "api_style": "openai",
622
+ "supports_model_listing": true,
623
+ "supports_streaming": true,
624
+ "supports_vision": true,
625
+ "supports_embeddings": false,
626
+ "supports_audio": false,
627
+ "supports_reasoning": false,
628
+ "openai_compatible": true,
629
+ "free_models_available": true,
630
+ "paid_models_available": true,
631
+ "needs_key": true,
632
+ "description": "Doubao Pro — ByteDance's LLM platform",
633
+ "last_updated": "2026-09-03T00:00:00+00:00",
634
+ "fallback_models": [
635
+ "doubao-pro-256k",
636
+ "doubao-pro-128k",
637
+ "doubao-lite-128k",
638
+ "doubao-lite-32k",
639
+ "doubao-pro-32k"
640
+ ],
641
+ "default_model": "doubao-pro-256k"
642
+ },
643
+ {
644
+ "id": "cartesia",
645
+ "name": "Cartesia Sonic",
646
+ "country": "United States",
647
+ "api_endpoint": "https://api.cartesia.ai/v1",
648
+ "api_style": "openai",
649
+ "supports_model_listing": true,
650
+ "supports_streaming": true,
651
+ "supports_vision": false,
652
+ "supports_embeddings": false,
653
+ "supports_audio": true,
654
+ "supports_reasoning": false,
655
+ "openai_compatible": false,
656
+ "free_models_available": false,
657
+ "paid_models_available": true,
658
+ "needs_key": true,
659
+ "description": "Sub-100ms ultra-realistic streaming voice generation",
660
+ "last_updated": "2026-09-01T00:00:00+00:00",
661
+ "fallback_models": [
662
+ "sonic-english",
663
+ "sonic-multilingual"
664
+ ],
665
+ "default_model": "sonic-multilingual"
666
+ },
667
+ {
668
+ "id": "cerebras",
669
+ "name": "Cerebras",
670
+ "country": "United States",
671
+ "api_endpoint": "https://api.cerebras.ai/v1",
672
+ "api_style": "openai",
673
+ "supports_model_listing": true,
674
+ "supports_streaming": true,
675
+ "supports_vision": false,
676
+ "supports_embeddings": false,
677
+ "supports_audio": false,
678
+ "supports_reasoning": false,
679
+ "openai_compatible": true,
680
+ "free_models_available": true,
681
+ "paid_models_available": true,
682
+ "needs_key": true,
683
+ "description": "Cerebras CS-3 — Wafer-scale AI inference",
684
+ "last_updated": "2026-09-03T00:00:00+00:00",
685
+ "fallback_models": [
686
+ "llama3.1-8b",
687
+ "llama3.1-70b",
688
+ "llama3.3-70b"
689
+ ],
690
+ "default_model": "llama3.3-70b"
691
+ },
692
+ {
693
+ "id": "cerebrium",
694
+ "name": "Cerebrium",
695
+ "country": "United States",
696
+ "api_endpoint": "https://api.cortex.cerebrium.ai/v4",
697
+ "api_style": "openai",
698
+ "supports_model_listing": true,
699
+ "supports_streaming": true,
700
+ "supports_vision": true,
701
+ "supports_embeddings": true,
702
+ "supports_audio": false,
703
+ "supports_reasoning": true,
704
+ "openai_compatible": true,
705
+ "free_models_available": true,
706
+ "paid_models_available": true,
707
+ "needs_key": true,
708
+ "description": "Serverless GPU infrastructure and production inference pipelines for open models",
709
+ "last_updated": "2026-09-01T00:00:00+00:00",
710
+ "fallback_models": [
711
+ "cerebrium-llama-3.1-70b",
712
+ "cerebrium-mistral-large",
713
+ "cerebrium-qwen-2.5-72b"
714
+ ],
715
+ "default_model": "cerebrium-llama-3.1-70b"
716
+ },
717
+ {
718
+ "id": "chutes",
719
+ "name": "Chutes.ai",
720
+ "country": "United States",
721
+ "api_endpoint": "https://api.chutes.ai/v1",
722
+ "api_style": "openai",
723
+ "supports_model_listing": true,
724
+ "supports_streaming": true,
725
+ "supports_vision": true,
726
+ "supports_embeddings": false,
727
+ "supports_audio": false,
728
+ "supports_reasoning": true,
729
+ "openai_compatible": true,
730
+ "free_models_available": true,
731
+ "paid_models_available": true,
732
+ "needs_key": true,
733
+ "description": "Decentralized serverless compute for open weights and fine-tuned models",
734
+ "last_updated": "2026-09-01T00:00:00+00:00",
735
+ "fallback_models": [
736
+ "deepseek-r1-distill-llama-70b",
737
+ "llama-3.3-70b-instruct",
738
+ "qwen-2.5-72b-instruct"
739
+ ],
740
+ "default_model": "deepseek-r1-distill-llama-70b"
741
+ },
742
+ {
743
+ "id": "clarifai",
744
+ "name": "Clarifai",
745
+ "country": "United States",
746
+ "api_endpoint": "https://api.clarifai.com/v2",
747
+ "api_style": "openai",
748
+ "supports_model_listing": true,
749
+ "supports_streaming": true,
750
+ "supports_vision": true,
751
+ "supports_embeddings": false,
752
+ "supports_audio": false,
753
+ "supports_reasoning": false,
754
+ "openai_compatible": true,
755
+ "free_models_available": true,
756
+ "paid_models_available": true,
757
+ "needs_key": true,
758
+ "description": "Computer vision & LLM platform",
759
+ "last_updated": "2026-07-31T05:17:14+00:00",
760
+ "fallback_models": [
761
+ "default"
762
+ ],
763
+ "default_model": "default"
764
+ },
765
+ {
766
+ "id": "cloudflare_workers_ai",
767
+ "name": "Cloudflare Workers AI",
768
+ "country": "United States",
769
+ "api_endpoint": "https://api.cloudflare.com/client/v4/accounts",
770
+ "api_style": "openai",
771
+ "supports_model_listing": true,
772
+ "supports_streaming": true,
773
+ "supports_vision": true,
774
+ "supports_embeddings": true,
775
+ "supports_audio": true,
776
+ "supports_reasoning": true,
777
+ "openai_compatible": true,
778
+ "free_models_available": true,
779
+ "paid_models_available": true,
780
+ "needs_key": true,
781
+ "description": "Serverless GPU global inference running at Cloudflare edge network",
782
+ "last_updated": "2026-09-01T00:00:00+00:00",
783
+ "fallback_models": [
784
+ "@cf/meta/llama-3.1-70b-instruct",
785
+ "@cf/deepseek-ai/deepseek-r1-distill-qwen-32b",
786
+ "@cf/mistral/mistral-7b-instruct-v0.2"
787
+ ],
788
+ "default_model": "@cf/meta/llama-3.1-70b-instruct"
789
+ },
790
+ {
791
+ "id": "codeium",
792
+ "name": "Codeium (Windsurf)",
793
+ "country": "United States",
794
+ "api_endpoint": "https://api.codeium.com/v1",
795
+ "api_style": "openai",
796
+ "supports_model_listing": true,
797
+ "supports_streaming": true,
798
+ "supports_vision": false,
799
+ "supports_embeddings": false,
800
+ "supports_audio": false,
801
+ "supports_reasoning": true,
802
+ "openai_compatible": true,
803
+ "free_models_available": true,
804
+ "paid_models_available": true,
805
+ "needs_key": true,
806
+ "description": "Windsurf Cascade and ultra-fast developer reasoning",
807
+ "last_updated": "2026-09-01T00:00:00+00:00",
808
+ "fallback_models": [
809
+ "cascade-base",
810
+ "codeium-chat",
811
+ "codeium-flow"
812
+ ],
813
+ "default_model": "cascade-base"
814
+ },
815
+ {
816
+ "id": "cohere",
817
+ "name": "Cohere",
818
+ "country": "Canada",
819
+ "api_endpoint": "https://api.cohere.ai/v1",
820
+ "api_style": "openai",
821
+ "supports_model_listing": true,
822
+ "supports_streaming": true,
823
+ "supports_vision": false,
824
+ "supports_embeddings": true,
825
+ "supports_audio": false,
826
+ "supports_reasoning": false,
827
+ "openai_compatible": true,
828
+ "free_models_available": true,
829
+ "paid_models_available": true,
830
+ "needs_key": true,
831
+ "description": "Command-R series & embedding models",
832
+ "last_updated": "2026-07-31T05:17:14+00:00",
833
+ "fallback_models": [
834
+ "command",
835
+ "command-light",
836
+ "command-nightly",
837
+ "command-r",
838
+ "command-r-plus",
839
+ "embed-english-v3.0"
840
+ ],
841
+ "default_model": "command-r-plus"
842
+ },
843
+ {
844
+ "id": "coreweave",
845
+ "name": "CoreWeave",
846
+ "country": "United States",
847
+ "api_endpoint": "https://api.coreweave.com/v1",
848
+ "api_style": "openai",
849
+ "supports_model_listing": true,
850
+ "supports_streaming": true,
851
+ "supports_vision": false,
852
+ "supports_embeddings": false,
853
+ "supports_audio": false,
854
+ "supports_reasoning": false,
855
+ "openai_compatible": true,
856
+ "free_models_available": true,
857
+ "paid_models_available": true,
858
+ "needs_key": true,
859
+ "description": "GPU cloud for AI",
860
+ "last_updated": "2026-07-31T05:17:14+00:00",
861
+ "fallback_models": [
862
+ "llama3-70b"
863
+ ],
864
+ "default_model": "llama3-70b"
865
+ },
866
+ {
867
+ "id": "cortex",
868
+ "name": "Cortex AI",
869
+ "country": "United States",
870
+ "api_endpoint": "https://api.cortex.ai/v1",
871
+ "api_style": "openai",
872
+ "supports_model_listing": true,
873
+ "supports_streaming": true,
874
+ "supports_vision": false,
875
+ "supports_embeddings": false,
876
+ "supports_audio": false,
877
+ "supports_reasoning": false,
878
+ "openai_compatible": true,
879
+ "free_models_available": true,
880
+ "paid_models_available": true,
881
+ "needs_key": true,
882
+ "description": "Enterprise AI platform",
883
+ "last_updated": "2026-07-31T05:17:14+00:00",
884
+ "fallback_models": [
885
+ "default"
886
+ ],
887
+ "default_model": "default"
888
+ },
889
+ {
890
+ "id": "cortex_so",
891
+ "name": "Cortex.so (Local Engine)",
892
+ "country": "International",
893
+ "api_endpoint": "http://localhost:39281/v1",
894
+ "api_style": "openai",
895
+ "supports_model_listing": true,
896
+ "supports_streaming": true,
897
+ "supports_vision": true,
898
+ "supports_embeddings": true,
899
+ "supports_audio": false,
900
+ "supports_reasoning": true,
901
+ "openai_compatible": true,
902
+ "free_models_available": true,
903
+ "paid_models_available": false,
904
+ "needs_key": false,
905
+ "description": "Hardware-accelerated local AI runtime supporting multi-engine model orchestration",
906
+ "last_updated": "2026-09-01T00:00:00+00:00",
907
+ "fallback_models": [
908
+ "llama3.1:8b",
909
+ "mistral:7b",
910
+ "gemma2:9b"
911
+ ],
912
+ "default_model": "llama3.1:8b"
913
+ },
914
+ {
915
+ "id": "cursor",
916
+ "name": "Cursor (Anysphere)",
917
+ "country": "United States",
918
+ "api_endpoint": "https://api.cursor.sh/v1",
919
+ "api_style": "openai",
920
+ "supports_model_listing": true,
921
+ "supports_streaming": true,
922
+ "supports_vision": true,
923
+ "supports_embeddings": false,
924
+ "supports_audio": false,
925
+ "supports_reasoning": true,
926
+ "openai_compatible": true,
927
+ "free_models_available": false,
928
+ "paid_models_available": true,
929
+ "needs_key": true,
930
+ "description": "AI-first code generation and repo-level agent inference",
931
+ "last_updated": "2026-09-01T00:00:00+00:00",
932
+ "fallback_models": [
933
+ "cursor-small",
934
+ "cursor-fast",
935
+ "cursor-large"
936
+ ],
937
+ "default_model": "cursor-large"
938
+ },
939
+ {
940
+ "id": "databricks",
941
+ "name": "Databricks",
942
+ "country": "United States",
943
+ "api_endpoint": "https://dbc-*.cloud.databricks.com/serving-endpoints",
944
+ "api_style": "openai",
945
+ "supports_model_listing": true,
946
+ "supports_streaming": true,
947
+ "supports_vision": false,
948
+ "supports_embeddings": true,
949
+ "supports_audio": false,
950
+ "supports_reasoning": false,
951
+ "openai_compatible": true,
952
+ "free_models_available": false,
953
+ "paid_models_available": true,
954
+ "needs_key": true,
955
+ "description": "Databricks DBRX — Enterprise AI platform",
956
+ "last_updated": "2026-09-03T00:00:00+00:00",
957
+ "fallback_models": [
958
+ "databricks-dbrx-instruct",
959
+ "llama-3-70b-chat"
960
+ ],
961
+ "default_model": "databricks-dbrx-instruct"
962
+ },
963
+ {
964
+ "id": "deepgram",
965
+ "name": "Deepgram",
966
+ "country": "United States",
967
+ "api_endpoint": "https://api.deepgram.com/v1",
968
+ "api_style": "openai",
969
+ "supports_model_listing": true,
970
+ "supports_streaming": false,
971
+ "supports_vision": false,
972
+ "supports_embeddings": false,
973
+ "supports_audio": true,
974
+ "supports_reasoning": false,
975
+ "openai_compatible": false,
976
+ "free_models_available": true,
977
+ "paid_models_available": true,
978
+ "needs_key": true,
979
+ "description": "Speech recognition & generation",
980
+ "last_updated": "2026-07-31T05:17:14+00:00",
981
+ "fallback_models": [
982
+ "nova-2"
983
+ ],
984
+ "default_model": "nova-2"
985
+ },
986
+ {
987
+ "id": "deepinfra",
988
+ "name": "DeepInfra",
989
+ "country": "United States",
990
+ "api_endpoint": "https://api.deepinfra.com/v1/openai",
991
+ "api_style": "openai",
992
+ "supports_model_listing": true,
993
+ "supports_streaming": true,
994
+ "supports_vision": false,
995
+ "supports_embeddings": false,
996
+ "supports_audio": false,
997
+ "supports_reasoning": false,
998
+ "openai_compatible": true,
999
+ "free_models_available": true,
1000
+ "paid_models_available": true,
1001
+ "needs_key": true,
1002
+ "description": "DeepInfra — Serverless AI inference",
1003
+ "last_updated": "2026-09-03T00:00:00+00:00",
1004
+ "fallback_models": [
1005
+ "meta-llama/Meta-Llama-3.1-405B-Instruct",
1006
+ "meta-llama/Meta-Llama-3.1-70B-Instruct",
1007
+ "mistralai/Mixtral-8x22B-Instruct-v0.1",
1008
+ "Qwen/Qwen2.5-72B-Instruct"
1009
+ ],
1010
+ "default_model": "meta-llama/Meta-Llama-3.1-70B-Instruct"
1011
+ },
1012
+ {
1013
+ "id": "deepseek",
1014
+ "name": "DeepSeek",
1015
+ "country": "China",
1016
+ "api_endpoint": "https://api.deepseek.com/v1",
1017
+ "api_style": "openai",
1018
+ "supports_model_listing": true,
1019
+ "supports_streaming": true,
1020
+ "supports_vision": false,
1021
+ "supports_embeddings": false,
1022
+ "supports_audio": false,
1023
+ "supports_reasoning": true,
1024
+ "openai_compatible": true,
1025
+ "free_models_available": true,
1026
+ "paid_models_available": true,
1027
+ "needs_key": true,
1028
+ "description": "DeepSeek-V3 & DeepSeek-R1 — strong reasoning",
1029
+ "last_updated": "2026-09-03T00:00:00+00:00",
1030
+ "fallback_models": [
1031
+ "deepseek-chat",
1032
+ "deepseek-coder",
1033
+ "deepseek-reasoner",
1034
+ "deepseek-v3",
1035
+ "deepseek-r1"
1036
+ ],
1037
+ "default_model": "deepseek-chat"
1038
+ },
1039
+ {
1040
+ "id": "deployai",
1041
+ "name": "Deploy AI",
1042
+ "country": "United States",
1043
+ "api_endpoint": "https://api.deploy.ai/v1",
1044
+ "api_style": "openai",
1045
+ "supports_model_listing": true,
1046
+ "supports_streaming": true,
1047
+ "supports_vision": false,
1048
+ "supports_embeddings": false,
1049
+ "supports_audio": false,
1050
+ "supports_reasoning": false,
1051
+ "openai_compatible": true,
1052
+ "free_models_available": true,
1053
+ "paid_models_available": true,
1054
+ "needs_key": true,
1055
+ "description": "Model deployment platform",
1056
+ "last_updated": "2026-07-31T05:17:14+00:00",
1057
+ "fallback_models": [
1058
+ "default"
1059
+ ],
1060
+ "default_model": "default"
1061
+ },
1062
+ {
1063
+ "id": "elastic",
1064
+ "name": "Elastic (ELSER)",
1065
+ "country": "United States",
1066
+ "api_endpoint": "https://api.elastic-cloud.com/v1",
1067
+ "api_style": "openai",
1068
+ "supports_model_listing": true,
1069
+ "supports_streaming": false,
1070
+ "supports_vision": false,
1071
+ "supports_embeddings": true,
1072
+ "supports_audio": false,
1073
+ "supports_reasoning": false,
1074
+ "openai_compatible": true,
1075
+ "free_models_available": true,
1076
+ "paid_models_available": true,
1077
+ "needs_key": true,
1078
+ "description": "Elasticsearch AI search",
1079
+ "last_updated": "2026-07-31T05:17:14+00:00",
1080
+ "fallback_models": [
1081
+ "elser-v2"
1082
+ ],
1083
+ "default_model": "elser-v2"
1084
+ },
1085
+ {
1086
+ "id": "eleutherai",
1087
+ "name": "EleutherAI",
1088
+ "country": "International",
1089
+ "api_endpoint": "https://api.eleuther.ai/v1",
1090
+ "api_style": "openai",
1091
+ "supports_model_listing": true,
1092
+ "supports_streaming": true,
1093
+ "supports_vision": false,
1094
+ "supports_embeddings": false,
1095
+ "supports_audio": false,
1096
+ "supports_reasoning": false,
1097
+ "openai_compatible": true,
1098
+ "free_models_available": true,
1099
+ "paid_models_available": false,
1100
+ "needs_key": false,
1101
+ "description": "Open-source AI research",
1102
+ "last_updated": "2026-07-31T05:17:14+00:00",
1103
+ "fallback_models": [
1104
+ "pythia-12b"
1105
+ ],
1106
+ "default_model": "pythia-12b"
1107
+ },
1108
+ {
1109
+ "id": "elevenlabs",
1110
+ "name": "ElevenLabs AI",
1111
+ "country": "United States",
1112
+ "api_endpoint": "https://api.elevenlabs.io/v1",
1113
+ "api_style": "openai",
1114
+ "supports_model_listing": true,
1115
+ "supports_streaming": true,
1116
+ "supports_vision": false,
1117
+ "supports_embeddings": false,
1118
+ "supports_audio": true,
1119
+ "supports_reasoning": false,
1120
+ "openai_compatible": false,
1121
+ "free_models_available": true,
1122
+ "paid_models_available": true,
1123
+ "needs_key": true,
1124
+ "description": "Industry leader in voice synthesis, voice cloning, and audio AI",
1125
+ "last_updated": "2026-09-01T00:00:00+00:00",
1126
+ "fallback_models": [
1127
+ "eleven_multilingual_v2",
1128
+ "eleven_turbo_v2_5",
1129
+ "eleven_flash_v2"
1130
+ ],
1131
+ "default_model": "eleven_multilingual_v2"
1132
+ },
1133
+ {
1134
+ "id": "exllamav2_hosted",
1135
+ "name": "ExLlamaV2 Server (Local)",
1136
+ "country": "International",
1137
+ "api_endpoint": "http://localhost:5000/v1",
1138
+ "api_style": "openai",
1139
+ "supports_model_listing": true,
1140
+ "supports_streaming": true,
1141
+ "supports_vision": false,
1142
+ "supports_embeddings": false,
1143
+ "supports_audio": false,
1144
+ "supports_reasoning": true,
1145
+ "openai_compatible": true,
1146
+ "free_models_available": true,
1147
+ "paid_models_available": false,
1148
+ "needs_key": false,
1149
+ "description": "Ultra-fast inference library designed from scratch for local modern GPUs",
1150
+ "last_updated": "2026-09-01T00:00:00+00:00",
1151
+ "fallback_models": [
1152
+ "exllamav2",
1153
+ "llama-3.1-70b-exl2"
1154
+ ],
1155
+ "default_model": "exllamav2"
1156
+ },
1157
+ {
1158
+ "id": "fal_ai",
1159
+ "name": "fal.ai",
1160
+ "country": "United States",
1161
+ "api_endpoint": "https://queue.fal.run",
1162
+ "api_style": "openai",
1163
+ "supports_model_listing": true,
1164
+ "supports_streaming": true,
1165
+ "supports_vision": true,
1166
+ "supports_embeddings": false,
1167
+ "supports_audio": true,
1168
+ "supports_reasoning": false,
1169
+ "openai_compatible": true,
1170
+ "free_models_available": true,
1171
+ "paid_models_available": true,
1172
+ "needs_key": true,
1173
+ "description": "Ultra-low latency serverless media, vision, and multimodal inference platform",
1174
+ "last_updated": "2026-09-01T00:00:00+00:00",
1175
+ "fallback_models": [
1176
+ "flux-pro-1.1",
1177
+ "fast-sdxl",
1178
+ "whisper-large-v3"
1179
+ ],
1180
+ "default_model": "flux-pro-1.1"
1181
+ },
1182
+ {
1183
+ "id": "featherless",
1184
+ "name": "Featherless AI",
1185
+ "country": "United States",
1186
+ "api_endpoint": "https://api.featherless.ai/v1",
1187
+ "api_style": "openai",
1188
+ "supports_model_listing": true,
1189
+ "supports_streaming": true,
1190
+ "supports_vision": false,
1191
+ "supports_embeddings": false,
1192
+ "supports_audio": false,
1193
+ "supports_reasoning": false,
1194
+ "openai_compatible": true,
1195
+ "free_models_available": true,
1196
+ "paid_models_available": true,
1197
+ "needs_key": true,
1198
+ "description": "Simple model API",
1199
+ "last_updated": "2026-07-31T05:17:14+00:00",
1200
+ "fallback_models": [
1201
+ "llama3-70b"
1202
+ ],
1203
+ "default_model": "llama3-70b"
1204
+ },
1205
+ {
1206
+ "id": "fireworks",
1207
+ "name": "Fireworks AI",
1208
+ "country": "United States",
1209
+ "api_endpoint": "https://api.fireworks.ai/inference/v1",
1210
+ "api_style": "openai",
1211
+ "supports_model_listing": true,
1212
+ "supports_streaming": true,
1213
+ "supports_vision": true,
1214
+ "supports_embeddings": true,
1215
+ "supports_audio": false,
1216
+ "supports_reasoning": false,
1217
+ "openai_compatible": true,
1218
+ "free_models_available": true,
1219
+ "paid_models_available": true,
1220
+ "needs_key": true,
1221
+ "description": "Fireworks AI — Fast inference platform",
1222
+ "last_updated": "2026-09-03T00:00:00+00:00",
1223
+ "fallback_models": [
1224
+ "accounts/fireworks/models/llama-v3p3-70b-instruct",
1225
+ "accounts/fireworks/models/mixtral-8x22b-instruct",
1226
+ "accounts/fireworks/models/qwen2-72b-instruct",
1227
+ "accounts/fireworks/models/deepseek-v3"
1228
+ ],
1229
+ "default_model": "accounts/fireworks/models/llama-v3p3-70b-instruct"
1230
+ },
1231
+ {
1232
+ "id": "forefront",
1233
+ "name": "Forefront AI",
1234
+ "country": "United States",
1235
+ "api_endpoint": "https://api.forefront.ai/v1",
1236
+ "api_style": "openai",
1237
+ "supports_model_listing": true,
1238
+ "supports_streaming": true,
1239
+ "supports_vision": false,
1240
+ "supports_embeddings": false,
1241
+ "supports_audio": false,
1242
+ "supports_reasoning": false,
1243
+ "openai_compatible": true,
1244
+ "free_models_available": true,
1245
+ "paid_models_available": true,
1246
+ "needs_key": true,
1247
+ "description": "Multi-provider API gateway",
1248
+ "last_updated": "2026-07-31T05:17:14+00:00",
1249
+ "fallback_models": [
1250
+ "claude-3",
1251
+ "gpt-4o"
1252
+ ],
1253
+ "default_model": "gpt-4o"
1254
+ },
1255
+ {
1256
+ "id": "friendli",
1257
+ "name": "FriendliAI",
1258
+ "country": "South Korea",
1259
+ "api_endpoint": "https://inference.friendli.ai/v1",
1260
+ "api_style": "openai",
1261
+ "supports_model_listing": true,
1262
+ "supports_streaming": true,
1263
+ "supports_vision": false,
1264
+ "supports_embeddings": false,
1265
+ "supports_audio": false,
1266
+ "supports_reasoning": false,
1267
+ "openai_compatible": true,
1268
+ "free_models_available": true,
1269
+ "paid_models_available": true,
1270
+ "needs_key": true,
1271
+ "description": "FriendliAI — Fast inference platform",
1272
+ "last_updated": "2026-09-03T00:00:00+00:00",
1273
+ "fallback_models": [
1274
+ "meta-llama-3.1-8b-instruct",
1275
+ "meta-llama-3.1-70b-instruct",
1276
+ "mixtral-8x7b-instruct-v0.1"
1277
+ ],
1278
+ "default_model": "meta-llama-3.1-70b-instruct"
1279
+ },
1280
+ {
1281
+ "id": "gcore",
1282
+ "name": "Gcore",
1283
+ "country": "Luxembourg",
1284
+ "api_endpoint": "https://api.gcore.com/v1",
1285
+ "api_style": "openai",
1286
+ "supports_model_listing": true,
1287
+ "supports_streaming": true,
1288
+ "supports_vision": false,
1289
+ "supports_embeddings": false,
1290
+ "supports_audio": false,
1291
+ "supports_reasoning": false,
1292
+ "openai_compatible": true,
1293
+ "free_models_available": true,
1294
+ "paid_models_available": true,
1295
+ "needs_key": true,
1296
+ "description": "Edge & cloud AI inference",
1297
+ "last_updated": "2026-07-31T05:17:14+00:00",
1298
+ "fallback_models": [
1299
+ "llama3-70b"
1300
+ ],
1301
+ "default_model": "llama3-70b"
1302
+ },
1303
+ {
1304
+ "id": "github-copilot",
1305
+ "name": "GitHub Copilot API",
1306
+ "country": "United States",
1307
+ "api_endpoint": "https://api.githubcopilot.com",
1308
+ "api_style": "openai",
1309
+ "supports_model_listing": true,
1310
+ "supports_streaming": true,
1311
+ "supports_vision": true,
1312
+ "supports_embeddings": true,
1313
+ "supports_audio": false,
1314
+ "supports_reasoning": true,
1315
+ "openai_compatible": true,
1316
+ "free_models_available": false,
1317
+ "paid_models_available": true,
1318
+ "needs_key": true,
1319
+ "description": "GitHub Copilot developer models — GPT-4o, Claude 3.5, Gemini 1.5",
1320
+ "last_updated": "2026-09-01T00:00:00+00:00",
1321
+ "fallback_models": [
1322
+ "gpt-4o",
1323
+ "claude-3.5-sonnet",
1324
+ "o1-preview"
1325
+ ],
1326
+ "default_model": "gpt-4o"
1327
+ },
1328
+ {
1329
+ "id": "gemini",
1330
+ "name": "Google Gemini",
1331
+ "country": "United States",
1332
+ "api_endpoint": "https://generativelanguage.googleapis.com/v1beta",
1333
+ "api_style": "gemini",
1334
+ "supports_model_listing": true,
1335
+ "supports_streaming": true,
1336
+ "supports_vision": true,
1337
+ "supports_embeddings": true,
1338
+ "supports_audio": false,
1339
+ "supports_reasoning": true,
1340
+ "openai_compatible": false,
1341
+ "free_models_available": true,
1342
+ "paid_models_available": true,
1343
+ "needs_key": true,
1344
+ "description": "Gemini 2.0 & 2.5 — Google DeepMind",
1345
+ "last_updated": "2026-09-03T00:00:00+00:00",
1346
+ "fallback_models": [
1347
+ "gemini-1.5-flash",
1348
+ "gemini-1.5-flash-8b",
1349
+ "gemini-1.5-flash-latest",
1350
+ "gemini-1.5-pro",
1351
+ "gemini-1.5-pro-latest",
1352
+ "gemini-2.0-flash",
1353
+ "gemini-2.0-flash-exp",
1354
+ "gemini-2.5-flash",
1355
+ "gemini-2.5-flash-preview-04-17",
1356
+ "gemini-2.5-pro",
1357
+ "gemini-2.5-pro-exp-03-25",
1358
+ "gemini-2.5-pro-preview-05-06",
1359
+ "gemma-2-27b-it",
1360
+ "gemma-2-9b-it"
1361
+ ],
1362
+ "default_model": "gemini-2.5-flash"
1363
+ },
1364
+ {
1365
+ "id": "palm",
1366
+ "name": "Google PaLM API",
1367
+ "country": "United States",
1368
+ "api_endpoint": "https://generativelanguage.googleapis.com/v1beta2",
1369
+ "api_style": "gemini",
1370
+ "supports_model_listing": true,
1371
+ "supports_streaming": true,
1372
+ "supports_vision": false,
1373
+ "supports_embeddings": false,
1374
+ "supports_audio": false,
1375
+ "supports_reasoning": false,
1376
+ "openai_compatible": false,
1377
+ "free_models_available": true,
1378
+ "paid_models_available": true,
1379
+ "needs_key": true,
1380
+ "description": "PaLM 2 — Google's previous-gen LLM",
1381
+ "last_updated": "2026-07-31T05:17:14+00:00",
1382
+ "fallback_models": [
1383
+ "chat-bison-001",
1384
+ "text-bison-001"
1385
+ ],
1386
+ "default_model": "chat-bison-001"
1387
+ },
1388
+ {
1389
+ "id": "google-vertex",
1390
+ "name": "Google Vertex AI",
1391
+ "country": "United States",
1392
+ "api_endpoint": "https://us-central1-aiplatform.googleapis.com/v1",
1393
+ "api_style": "gemini",
1394
+ "supports_model_listing": true,
1395
+ "supports_streaming": true,
1396
+ "supports_vision": true,
1397
+ "supports_embeddings": true,
1398
+ "supports_audio": false,
1399
+ "supports_reasoning": true,
1400
+ "openai_compatible": false,
1401
+ "free_models_available": false,
1402
+ "paid_models_available": true,
1403
+ "needs_key": true,
1404
+ "description": "Google Vertex AI — Enterprise Gemini",
1405
+ "last_updated": "2026-09-03T00:00:00+00:00",
1406
+ "fallback_models": [
1407
+ "gemini-2.5-pro",
1408
+ "gemini-2.5-flash",
1409
+ "gemini-2.0-flash",
1410
+ "text-embedding-004",
1411
+ "claude-3-5-sonnet",
1412
+ "gemini-1.5-pro",
1413
+ "llama3-70b"
1414
+ ],
1415
+ "default_model": "gemini-2.5-pro"
1416
+ },
1417
+ {
1418
+ "id": "gpt4all",
1419
+ "name": "GPT4All",
1420
+ "country": "International",
1421
+ "api_endpoint": "http://localhost:4891",
1422
+ "api_style": "ollama",
1423
+ "supports_model_listing": true,
1424
+ "supports_streaming": true,
1425
+ "supports_vision": false,
1426
+ "supports_embeddings": true,
1427
+ "supports_audio": false,
1428
+ "supports_reasoning": false,
1429
+ "openai_compatible": true,
1430
+ "free_models_available": true,
1431
+ "paid_models_available": false,
1432
+ "needs_key": false,
1433
+ "description": "GPT4All — Local AI assistant",
1434
+ "last_updated": "2026-09-03T00:00:00+00:00",
1435
+ "fallback_models": [
1436
+ "gpt4all-falcon",
1437
+ "gpt4all-llama3-8b",
1438
+ "gpt4all-mistral-7b"
1439
+ ],
1440
+ "default_model": "gpt4all-llama3-8b"
1441
+ },
1442
+ {
1443
+ "id": "gradient",
1444
+ "name": "Gradient AI",
1445
+ "country": "Canada",
1446
+ "api_endpoint": "https://api.gradient.ai/v1",
1447
+ "api_style": "openai",
1448
+ "supports_model_listing": true,
1449
+ "supports_streaming": true,
1450
+ "supports_vision": false,
1451
+ "supports_embeddings": false,
1452
+ "supports_audio": false,
1453
+ "supports_reasoning": false,
1454
+ "openai_compatible": true,
1455
+ "free_models_available": true,
1456
+ "paid_models_available": true,
1457
+ "needs_key": true,
1458
+ "description": "Serverless fine-tuning & inference",
1459
+ "last_updated": "2026-07-31T05:17:14+00:00",
1460
+ "fallback_models": [
1461
+ "llama3-70b"
1462
+ ],
1463
+ "default_model": "llama3-70b"
1464
+ },
1465
+ {
1466
+ "id": "groq",
1467
+ "name": "Groq",
1468
+ "country": "United States",
1469
+ "api_endpoint": "https://api.groq.com/openai/v1",
1470
+ "api_style": "openai",
1471
+ "supports_model_listing": true,
1472
+ "supports_streaming": true,
1473
+ "supports_vision": false,
1474
+ "supports_embeddings": false,
1475
+ "supports_audio": false,
1476
+ "supports_reasoning": false,
1477
+ "openai_compatible": true,
1478
+ "free_models_available": true,
1479
+ "paid_models_available": true,
1480
+ "needs_key": true,
1481
+ "description": "LPU inference — exceptionally fast",
1482
+ "last_updated": "2026-07-31T05:17:14+00:00",
1483
+ "fallback_models": [
1484
+ "deepseek-r1-distill-llama-70b",
1485
+ "gemma2-9b-it",
1486
+ "llama-3.1-70b-versatile",
1487
+ "llama-3.1-8b-instant",
1488
+ "llama-3.3-70b-versatile",
1489
+ "llama-4-scout-17b-16e-instruct",
1490
+ "llama3-70b-8192",
1491
+ "llama3.1-70b-versatile",
1492
+ "mixtral-8x7b-32768"
1493
+ ],
1494
+ "default_model": "llama-3.3-70b-versatile"
1495
+ },
1496
+ {
1497
+ "id": "h2o",
1498
+ "name": "H2O.ai",
1499
+ "country": "United States",
1500
+ "api_endpoint": "https://api.h2o.ai/v1",
1501
+ "api_style": "openai",
1502
+ "supports_model_listing": true,
1503
+ "supports_streaming": true,
1504
+ "supports_vision": false,
1505
+ "supports_embeddings": false,
1506
+ "supports_audio": false,
1507
+ "supports_reasoning": false,
1508
+ "openai_compatible": true,
1509
+ "free_models_available": true,
1510
+ "paid_models_available": true,
1511
+ "needs_key": true,
1512
+ "description": "Enterprise AI & h2oGPT",
1513
+ "last_updated": "2026-07-31T05:17:14+00:00",
1514
+ "fallback_models": [
1515
+ "h2ogpt-gm"
1516
+ ],
1517
+ "default_model": "h2ogpt-gm"
1518
+ },
1519
+ {
1520
+ "id": "hotpot",
1521
+ "name": "Hotpot AI",
1522
+ "country": "United States",
1523
+ "api_endpoint": "https://api.hotpot.ai/v1",
1524
+ "api_style": "openai",
1525
+ "supports_model_listing": true,
1526
+ "supports_streaming": true,
1527
+ "supports_vision": true,
1528
+ "supports_embeddings": false,
1529
+ "supports_audio": false,
1530
+ "supports_reasoning": false,
1531
+ "openai_compatible": true,
1532
+ "free_models_available": true,
1533
+ "paid_models_available": true,
1534
+ "needs_key": true,
1535
+ "description": "AI image editing & generation",
1536
+ "last_updated": "2026-07-31T05:17:14+00:00",
1537
+ "fallback_models": [
1538
+ "default"
1539
+ ],
1540
+ "default_model": "default"
1541
+ },
1542
+ {
1543
+ "id": "huggingface",
1544
+ "name": "Hugging Face",
1545
+ "country": "United States",
1546
+ "api_endpoint": "https://api-inference.huggingface.co",
1547
+ "api_style": "openai",
1548
+ "supports_model_listing": true,
1549
+ "supports_streaming": true,
1550
+ "supports_vision": true,
1551
+ "supports_embeddings": true,
1552
+ "supports_audio": false,
1553
+ "supports_reasoning": false,
1554
+ "openai_compatible": true,
1555
+ "free_models_available": true,
1556
+ "paid_models_available": true,
1557
+ "needs_key": true,
1558
+ "description": "Hugging Face Inference — Open-source model hub",
1559
+ "last_updated": "2026-09-03T00:00:00+00:00",
1560
+ "fallback_models": [
1561
+ "meta-llama/Llama-3.1-405B-Instruct",
1562
+ "meta-llama/Llama-3.1-70B-Instruct",
1563
+ "mistralai/Mixtral-8x22B-Instruct-v0.1",
1564
+ "Qwen/Qwen2.5-72B-Instruct"
1565
+ ],
1566
+ "default_model": "meta-llama/Llama-3.1-70B-Instruct"
1567
+ },
1568
+ {
1569
+ "id": "tgi_hosted",
1570
+ "name": "Hugging Face TGI (Text Generation Inference)",
1571
+ "country": "United States",
1572
+ "api_endpoint": "http://localhost:8080/v1",
1573
+ "api_style": "openai",
1574
+ "supports_model_listing": true,
1575
+ "supports_streaming": true,
1576
+ "supports_vision": true,
1577
+ "supports_embeddings": false,
1578
+ "supports_audio": false,
1579
+ "supports_reasoning": true,
1580
+ "openai_compatible": true,
1581
+ "free_models_available": true,
1582
+ "paid_models_available": false,
1583
+ "needs_key": false,
1584
+ "description": "Purpose-built solution for deploying and serving large language models",
1585
+ "last_updated": "2026-09-01T00:00:00+00:00",
1586
+ "fallback_models": [
1587
+ "tgi",
1588
+ "mistralai/Mistral-7B-Instruct-v0.3",
1589
+ "meta-llama/Meta-Llama-3-8B-Instruct"
1590
+ ],
1591
+ "default_model": "tgi"
1592
+ },
1593
+ {
1594
+ "id": "hyperbolic",
1595
+ "name": "Hyperbolic",
1596
+ "country": "United States",
1597
+ "api_endpoint": "https://api.hyperbolic.xyz/v1",
1598
+ "api_style": "openai",
1599
+ "supports_model_listing": true,
1600
+ "supports_streaming": true,
1601
+ "supports_vision": false,
1602
+ "supports_embeddings": false,
1603
+ "supports_audio": false,
1604
+ "supports_reasoning": false,
1605
+ "openai_compatible": true,
1606
+ "free_models_available": true,
1607
+ "paid_models_available": true,
1608
+ "needs_key": true,
1609
+ "description": "Open model inference",
1610
+ "last_updated": "2026-07-31T05:17:14+00:00",
1611
+ "fallback_models": [
1612
+ "llama3-70b",
1613
+ "qwen2-72b",
1614
+ "meta-llama/Llama-3.3-70B-Instruct",
1615
+ "deepseek-ai/DeepSeek-V3"
1616
+ ],
1617
+ "default_model": "llama3-70b"
1618
+ },
1619
+ {
1620
+ "id": "ibm-watsonx",
1621
+ "name": "IBM watsonx",
1622
+ "country": "United States",
1623
+ "api_endpoint": "https://us-south.ml.cloud.ibm.com/ml/v1",
1624
+ "api_style": "openai",
1625
+ "supports_model_listing": true,
1626
+ "supports_streaming": true,
1627
+ "supports_vision": false,
1628
+ "supports_embeddings": true,
1629
+ "supports_audio": false,
1630
+ "supports_reasoning": false,
1631
+ "openai_compatible": true,
1632
+ "free_models_available": false,
1633
+ "paid_models_available": true,
1634
+ "needs_key": true,
1635
+ "description": "IBM watsonx — Enterprise AI platform",
1636
+ "last_updated": "2026-09-03T00:00:00+00:00",
1637
+ "fallback_models": [
1638
+ "ibm/granite-3-8b-instruct",
1639
+ "ibm/granite-3-2b-instruct",
1640
+ "meta-llama/llama-3-1-70b-instruct"
1641
+ ],
1642
+ "default_model": "ibm/granite-3-8b-instruct"
1643
+ },
1644
+ {
1645
+ "id": "watson",
1646
+ "name": "IBM Watsonx.ai",
1647
+ "country": "United States",
1648
+ "api_endpoint": "https://api.watsonx.ai/v1",
1649
+ "api_style": "openai",
1650
+ "supports_model_listing": true,
1651
+ "supports_streaming": true,
1652
+ "supports_vision": false,
1653
+ "supports_embeddings": false,
1654
+ "supports_audio": false,
1655
+ "supports_reasoning": false,
1656
+ "openai_compatible": true,
1657
+ "free_models_available": true,
1658
+ "paid_models_available": true,
1659
+ "needs_key": true,
1660
+ "description": "IBM enterprise AI platform",
1661
+ "last_updated": "2026-07-31T05:17:14+00:00",
1662
+ "fallback_models": [
1663
+ "granite-13b",
1664
+ "llama3-70b"
1665
+ ],
1666
+ "default_model": "granite-13b"
1667
+ },
1668
+ {
1669
+ "id": "idiomatic",
1670
+ "name": "Idiomatic AI",
1671
+ "country": "Canada",
1672
+ "api_endpoint": "https://api.idiomatic.ai/v1",
1673
+ "api_style": "openai",
1674
+ "supports_model_listing": true,
1675
+ "supports_streaming": true,
1676
+ "supports_vision": false,
1677
+ "supports_embeddings": false,
1678
+ "supports_audio": false,
1679
+ "supports_reasoning": false,
1680
+ "openai_compatible": true,
1681
+ "free_models_available": true,
1682
+ "paid_models_available": true,
1683
+ "needs_key": true,
1684
+ "description": "NLP for customer experience",
1685
+ "last_updated": "2026-07-31T05:17:14+00:00",
1686
+ "fallback_models": [
1687
+ "default"
1688
+ ],
1689
+ "default_model": "default"
1690
+ },
1691
+ {
1692
+ "id": "xunfei-spark",
1693
+ "name": "iFlytek Spark",
1694
+ "country": "China",
1695
+ "api_endpoint": "https://spark-api-open.xf-yun.com/v1",
1696
+ "api_style": "openai",
1697
+ "supports_model_listing": true,
1698
+ "supports_streaming": true,
1699
+ "supports_vision": true,
1700
+ "supports_embeddings": false,
1701
+ "supports_audio": false,
1702
+ "supports_reasoning": false,
1703
+ "openai_compatible": true,
1704
+ "free_models_available": true,
1705
+ "paid_models_available": true,
1706
+ "needs_key": true,
1707
+ "description": "Spark 4.0 — iFlytek's multimodal LLM",
1708
+ "last_updated": "2026-09-03T00:00:00+00:00",
1709
+ "fallback_models": [
1710
+ "spark4.0-ultra",
1711
+ "spark-max",
1712
+ "spark-pro",
1713
+ "spark-lite"
1714
+ ],
1715
+ "default_model": "spark4.0-ultra"
1716
+ },
1717
+ {
1718
+ "id": "infermatic",
1719
+ "name": "Infermatic.ai",
1720
+ "country": "United States",
1721
+ "api_endpoint": "https://api.infermatic.ai/v1",
1722
+ "api_style": "openai",
1723
+ "supports_model_listing": true,
1724
+ "supports_streaming": true,
1725
+ "supports_vision": true,
1726
+ "supports_embeddings": false,
1727
+ "supports_audio": false,
1728
+ "supports_reasoning": true,
1729
+ "openai_compatible": true,
1730
+ "free_models_available": false,
1731
+ "paid_models_available": true,
1732
+ "needs_key": true,
1733
+ "description": "Platform specialized in open-weights models, storytelling, roleplay, and uncensored LLMs",
1734
+ "last_updated": "2026-09-01T00:00:00+00:00",
1735
+ "fallback_models": [
1736
+ "magnum-v4-72b",
1737
+ "command-r-plus",
1738
+ "llama-3.1-nemotron-70b"
1739
+ ],
1740
+ "default_model": "magnum-v4-72b"
1741
+ },
1742
+ {
1743
+ "id": "inflection",
1744
+ "name": "Inflection AI",
1745
+ "country": "United States",
1746
+ "api_endpoint": "https://api.inflection.ai/v1",
1747
+ "api_style": "openai",
1748
+ "supports_model_listing": true,
1749
+ "supports_streaming": true,
1750
+ "supports_vision": false,
1751
+ "supports_embeddings": false,
1752
+ "supports_audio": false,
1753
+ "supports_reasoning": false,
1754
+ "openai_compatible": true,
1755
+ "free_models_available": true,
1756
+ "paid_models_available": true,
1757
+ "needs_key": true,
1758
+ "description": "Pi — personal AI",
1759
+ "last_updated": "2026-07-31T05:17:14+00:00",
1760
+ "fallback_models": [
1761
+ "inflection-2.5"
1762
+ ],
1763
+ "default_model": "inflection-2.5"
1764
+ },
1765
+ {
1766
+ "id": "internlm",
1767
+ "name": "InternLM (Shanghai AI Lab)",
1768
+ "country": "China",
1769
+ "api_endpoint": "https://internlm-chat.intern-ai.org.cn/api/v1",
1770
+ "api_style": "openai",
1771
+ "supports_model_listing": true,
1772
+ "supports_streaming": true,
1773
+ "supports_vision": true,
1774
+ "supports_embeddings": false,
1775
+ "supports_audio": false,
1776
+ "supports_reasoning": true,
1777
+ "openai_compatible": true,
1778
+ "free_models_available": true,
1779
+ "paid_models_available": true,
1780
+ "needs_key": true,
1781
+ "description": "Open-source bilingual foundation models with strong reasoning and coding capabilities",
1782
+ "last_updated": "2026-09-01T00:00:00+00:00",
1783
+ "fallback_models": [
1784
+ "internlm2_5-20b-chat",
1785
+ "internlm2_5-7b-chat",
1786
+ "internlm-xcomposer2d5-7b"
1787
+ ],
1788
+ "default_model": "internlm2_5-20b-chat"
1789
+ },
1790
+ {
1791
+ "id": "jan_ai",
1792
+ "name": "Jan.ai (Local Server)",
1793
+ "country": "Singapore",
1794
+ "api_endpoint": "http://localhost:1337/v1",
1795
+ "api_style": "openai",
1796
+ "supports_model_listing": true,
1797
+ "supports_streaming": true,
1798
+ "supports_vision": true,
1799
+ "supports_embeddings": false,
1800
+ "supports_audio": false,
1801
+ "supports_reasoning": true,
1802
+ "openai_compatible": true,
1803
+ "free_models_available": true,
1804
+ "paid_models_available": false,
1805
+ "needs_key": false,
1806
+ "description": "Open source, local-first alternative to ChatGPT running 100% offline on your device",
1807
+ "last_updated": "2026-09-01T00:00:00+00:00",
1808
+ "fallback_models": [
1809
+ "llama3.1-8b-instruct",
1810
+ "deepseek-coder-6.7b",
1811
+ "mistral-ins-7b-q4"
1812
+ ],
1813
+ "default_model": "llama3.1-8b-instruct"
1814
+ },
1815
+ {
1816
+ "id": "jasper",
1817
+ "name": "Jasper AI",
1818
+ "country": "United States",
1819
+ "api_endpoint": "https://api.jasper.ai/v1",
1820
+ "api_style": "openai",
1821
+ "supports_model_listing": true,
1822
+ "supports_streaming": true,
1823
+ "supports_vision": false,
1824
+ "supports_embeddings": false,
1825
+ "supports_audio": false,
1826
+ "supports_reasoning": false,
1827
+ "openai_compatible": true,
1828
+ "free_models_available": true,
1829
+ "paid_models_available": true,
1830
+ "needs_key": true,
1831
+ "description": "AI content platform",
1832
+ "last_updated": "2026-07-31T05:17:14+00:00",
1833
+ "fallback_models": [
1834
+ "default"
1835
+ ],
1836
+ "default_model": "default"
1837
+ },
1838
+ {
1839
+ "id": "jina",
1840
+ "name": "Jina AI",
1841
+ "country": "Germany",
1842
+ "api_endpoint": "https://api.jina.ai/v1",
1843
+ "api_style": "openai",
1844
+ "supports_model_listing": true,
1845
+ "supports_streaming": false,
1846
+ "supports_vision": false,
1847
+ "supports_embeddings": true,
1848
+ "supports_audio": false,
1849
+ "supports_reasoning": false,
1850
+ "openai_compatible": true,
1851
+ "free_models_available": true,
1852
+ "paid_models_available": true,
1853
+ "needs_key": true,
1854
+ "description": "Embeddings & neural search",
1855
+ "last_updated": "2026-07-31T05:17:14+00:00",
1856
+ "fallback_models": [
1857
+ "jina-embeddings-v3",
1858
+ "reader-lm"
1859
+ ],
1860
+ "default_model": "jina-embeddings-v3"
1861
+ },
1862
+ {
1863
+ "id": "kakaocorp",
1864
+ "name": "Kakao Corp (HyperCLOVA)",
1865
+ "country": "South Korea",
1866
+ "api_endpoint": "https://api.kakaocloud.com/v1",
1867
+ "api_style": "openai",
1868
+ "supports_model_listing": true,
1869
+ "supports_streaming": true,
1870
+ "supports_vision": false,
1871
+ "supports_embeddings": false,
1872
+ "supports_audio": false,
1873
+ "supports_reasoning": false,
1874
+ "openai_compatible": true,
1875
+ "free_models_available": true,
1876
+ "paid_models_available": true,
1877
+ "needs_key": true,
1878
+ "description": "HyperCLOVA X — Kakao's Korean-optimized LLM",
1879
+ "last_updated": "2026-09-03T00:00:00+00:00",
1880
+ "fallback_models": [
1881
+ "hyperclova-x-seed-llama-3.1-70b",
1882
+ "hyperclova-x-seed-llama-3.1-8b"
1883
+ ],
1884
+ "default_model": "hyperclova-x-seed-llama-3.1-70b"
1885
+ },
1886
+ {
1887
+ "id": "koala",
1888
+ "name": "Koala AI",
1889
+ "country": "United States",
1890
+ "api_endpoint": "https://api.koala.sh/v1",
1891
+ "api_style": "openai",
1892
+ "supports_model_listing": true,
1893
+ "supports_streaming": true,
1894
+ "supports_vision": false,
1895
+ "supports_embeddings": false,
1896
+ "supports_audio": false,
1897
+ "supports_reasoning": false,
1898
+ "openai_compatible": true,
1899
+ "free_models_available": true,
1900
+ "paid_models_available": true,
1901
+ "needs_key": true,
1902
+ "description": "Fine-tuned chat model",
1903
+ "last_updated": "2026-07-31T05:17:14+00:00",
1904
+ "fallback_models": [
1905
+ "koala-13b"
1906
+ ],
1907
+ "default_model": "koala-13b"
1908
+ },
1909
+ {
1910
+ "id": "koboldai",
1911
+ "name": "KoboldAI",
1912
+ "country": "International",
1913
+ "api_endpoint": "https://lite.koboldai.net/v1",
1914
+ "api_style": "openai",
1915
+ "supports_model_listing": true,
1916
+ "supports_streaming": true,
1917
+ "supports_vision": false,
1918
+ "supports_embeddings": false,
1919
+ "supports_audio": false,
1920
+ "supports_reasoning": false,
1921
+ "openai_compatible": true,
1922
+ "free_models_available": true,
1923
+ "paid_models_available": false,
1924
+ "needs_key": false,
1925
+ "description": "Community-run open model inference",
1926
+ "last_updated": "2026-07-31T05:17:14+00:00",
1927
+ "fallback_models": [
1928
+ "mythomax-l2-13b",
1929
+ "pygmalion-6b"
1930
+ ],
1931
+ "default_model": "mythomax-l2-13b"
1932
+ },
1933
+ {
1934
+ "id": "koboldcpp",
1935
+ "name": "KoboldCpp",
1936
+ "country": "International",
1937
+ "api_endpoint": "http://localhost:5001/v1",
1938
+ "api_style": "openai",
1939
+ "supports_model_listing": true,
1940
+ "supports_streaming": true,
1941
+ "supports_vision": false,
1942
+ "supports_embeddings": false,
1943
+ "supports_audio": false,
1944
+ "supports_reasoning": false,
1945
+ "openai_compatible": true,
1946
+ "free_models_available": true,
1947
+ "paid_models_available": false,
1948
+ "needs_key": false,
1949
+ "description": "KoboldCpp — Local LLM runner",
1950
+ "last_updated": "2026-09-03T00:00:00+00:00",
1951
+ "fallback_models": [
1952
+ "default",
1953
+ "kobold-model",
1954
+ "fimbulvetr-11b-v2"
1955
+ ],
1956
+ "default_model": "default"
1957
+ },
1958
+ {
1959
+ "id": "kyutai",
1960
+ "name": "Kyutai (Moshi)",
1961
+ "country": "France",
1962
+ "api_endpoint": "https://api.kyutai.org/v1",
1963
+ "api_style": "openai",
1964
+ "supports_model_listing": true,
1965
+ "supports_streaming": true,
1966
+ "supports_vision": false,
1967
+ "supports_embeddings": false,
1968
+ "supports_audio": true,
1969
+ "supports_reasoning": true,
1970
+ "openai_compatible": true,
1971
+ "free_models_available": true,
1972
+ "paid_models_available": false,
1973
+ "needs_key": false,
1974
+ "description": "Full-duplex real-time multimodal voice and conversational AI foundation lab",
1975
+ "last_updated": "2026-09-01T00:00:00+00:00",
1976
+ "fallback_models": [
1977
+ "moshi-v1",
1978
+ "mimi-audio-codec",
1979
+ "helium-1"
1980
+ ],
1981
+ "default_model": "moshi-v1"
1982
+ },
1983
+ {
1984
+ "id": "lambda",
1985
+ "name": "Lambda Labs",
1986
+ "country": "United States",
1987
+ "api_endpoint": "https://api.lambdalabs.com/v1",
1988
+ "api_style": "openai",
1989
+ "supports_model_listing": true,
1990
+ "supports_streaming": true,
1991
+ "supports_vision": false,
1992
+ "supports_embeddings": false,
1993
+ "supports_audio": false,
1994
+ "supports_reasoning": false,
1995
+ "openai_compatible": true,
1996
+ "free_models_available": false,
1997
+ "paid_models_available": true,
1998
+ "needs_key": true,
1999
+ "description": "Lambda Labs — GPU cloud for AI",
2000
+ "last_updated": "2026-09-03T00:00:00+00:00",
2001
+ "fallback_models": [
2002
+ "llama-3.1-405b-instruct-fp8",
2003
+ "llama-3.1-70b-instruct",
2004
+ "llama-3.1-8b-instruct",
2005
+ "hermes-2-pro",
2006
+ "llama3-70b"
2007
+ ],
2008
+ "default_model": "llama-3.1-70b-instruct"
2009
+ },
2010
+ {
2011
+ "id": "lepton",
2012
+ "name": "Lepton AI",
2013
+ "country": "United States",
2014
+ "api_endpoint": "https://lepton.ai/api/v1",
2015
+ "api_style": "openai",
2016
+ "supports_model_listing": true,
2017
+ "supports_streaming": true,
2018
+ "supports_vision": false,
2019
+ "supports_embeddings": false,
2020
+ "supports_audio": false,
2021
+ "supports_reasoning": false,
2022
+ "openai_compatible": true,
2023
+ "free_models_available": true,
2024
+ "paid_models_available": true,
2025
+ "needs_key": true,
2026
+ "description": "Serverless GPU cloud",
2027
+ "last_updated": "2026-07-31T05:17:14+00:00",
2028
+ "fallback_models": [
2029
+ "llama3-70b",
2030
+ "mixtral-8x22b"
2031
+ ],
2032
+ "default_model": "llama3-70b"
2033
+ },
2034
+ {
2035
+ "id": "liquid_ai",
2036
+ "name": "Liquid AI",
2037
+ "country": "United States",
2038
+ "api_endpoint": "https://api.liquid.ai/v1",
2039
+ "api_style": "openai",
2040
+ "supports_model_listing": true,
2041
+ "supports_streaming": true,
2042
+ "supports_vision": false,
2043
+ "supports_embeddings": false,
2044
+ "supports_audio": false,
2045
+ "supports_reasoning": true,
2046
+ "openai_compatible": true,
2047
+ "free_models_available": true,
2048
+ "paid_models_available": true,
2049
+ "needs_key": true,
2050
+ "description": "Liquid Neural Networks (LNN) and Liquid Foundation Models (LFM-40B, LFM-7B)",
2051
+ "last_updated": "2026-09-01T00:00:00+00:00",
2052
+ "fallback_models": [
2053
+ "lfm-40b-chat",
2054
+ "lfm-7b-chat",
2055
+ "lfm-3b-chat"
2056
+ ],
2057
+ "default_model": "lfm-40b-chat"
2058
+ },
2059
+ {
2060
+ "id": "llama-api",
2061
+ "name": "Llama API",
2062
+ "country": "United States",
2063
+ "api_endpoint": "https://api.llama-api.com/v1",
2064
+ "api_style": "openai",
2065
+ "supports_model_listing": true,
2066
+ "supports_streaming": true,
2067
+ "supports_vision": false,
2068
+ "supports_embeddings": false,
2069
+ "supports_audio": false,
2070
+ "supports_reasoning": false,
2071
+ "openai_compatible": true,
2072
+ "free_models_available": true,
2073
+ "paid_models_available": true,
2074
+ "needs_key": true,
2075
+ "description": "Meta Llama hosted API",
2076
+ "last_updated": "2026-07-31T05:17:14+00:00",
2077
+ "fallback_models": [
2078
+ "llama3-70b",
2079
+ "llama3-8b"
2080
+ ],
2081
+ "default_model": "llama3-70b"
2082
+ },
2083
+ {
2084
+ "id": "llamacpp_server",
2085
+ "name": "llama.cpp Server (Local)",
2086
+ "country": "International",
2087
+ "api_endpoint": "http://localhost:8080/v1",
2088
+ "api_style": "openai",
2089
+ "supports_model_listing": true,
2090
+ "supports_streaming": true,
2091
+ "supports_vision": true,
2092
+ "supports_embeddings": true,
2093
+ "supports_audio": false,
2094
+ "supports_reasoning": true,
2095
+ "openai_compatible": true,
2096
+ "free_models_available": true,
2097
+ "paid_models_available": false,
2098
+ "needs_key": false,
2099
+ "description": "Lightweight, dependency-free local inference server running GGUF models on CPU and GPU",
2100
+ "last_updated": "2026-09-01T00:00:00+00:00",
2101
+ "fallback_models": [
2102
+ "gpt-3.5-turbo",
2103
+ "llama-3.1-8b-instruct",
2104
+ "qwen2.5-coder-7b"
2105
+ ],
2106
+ "default_model": "gpt-3.5-turbo"
2107
+ },
2108
+ {
2109
+ "id": "lmstudio",
2110
+ "name": "LM Studio",
2111
+ "country": "International",
2112
+ "api_endpoint": "http://localhost:1234/v1",
2113
+ "api_style": "openai",
2114
+ "supports_model_listing": true,
2115
+ "supports_streaming": true,
2116
+ "supports_vision": false,
2117
+ "supports_embeddings": true,
2118
+ "supports_audio": false,
2119
+ "supports_reasoning": false,
2120
+ "openai_compatible": true,
2121
+ "free_models_available": true,
2122
+ "paid_models_available": false,
2123
+ "needs_key": false,
2124
+ "description": "LM Studio — Local model runner",
2125
+ "last_updated": "2026-09-03T00:00:00+00:00",
2126
+ "fallback_models": [
2127
+ "local-model"
2128
+ ],
2129
+ "default_model": "local-model"
2130
+ },
2131
+ {
2132
+ "id": "lmdeploy_hosted",
2133
+ "name": "LMDeploy Server (Shanghai AI Lab)",
2134
+ "country": "China",
2135
+ "api_endpoint": "http://localhost:23333/v1",
2136
+ "api_style": "openai",
2137
+ "supports_model_listing": true,
2138
+ "supports_streaming": true,
2139
+ "supports_vision": true,
2140
+ "supports_embeddings": false,
2141
+ "supports_audio": false,
2142
+ "supports_reasoning": true,
2143
+ "openai_compatible": true,
2144
+ "free_models_available": true,
2145
+ "paid_models_available": false,
2146
+ "needs_key": false,
2147
+ "description": "High-performance toolkit for compressing, deploying, and serving LLM and VLM models",
2148
+ "last_updated": "2026-09-01T00:00:00+00:00",
2149
+ "fallback_models": [
2150
+ "internlm2_5-20b-chat",
2151
+ "internlm-xcomposer2d5-7b",
2152
+ "qwen2.5-72b-instruct"
2153
+ ],
2154
+ "default_model": "internlm2_5-20b-chat"
2155
+ },
2156
+ {
2157
+ "id": "lmsys",
2158
+ "name": "LMSYS (FastChat)",
2159
+ "country": "United States",
2160
+ "api_endpoint": "https://api.lmsys.org/v1",
2161
+ "api_style": "openai",
2162
+ "supports_model_listing": true,
2163
+ "supports_streaming": true,
2164
+ "supports_vision": false,
2165
+ "supports_embeddings": false,
2166
+ "supports_audio": false,
2167
+ "supports_reasoning": false,
2168
+ "openai_compatible": true,
2169
+ "free_models_available": true,
2170
+ "paid_models_available": false,
2171
+ "needs_key": false,
2172
+ "description": "Open chat model arena",
2173
+ "last_updated": "2026-07-31T05:17:14+00:00",
2174
+ "fallback_models": [
2175
+ "vicuna-13b"
2176
+ ],
2177
+ "default_model": "vicuna-13b"
2178
+ },
2179
+ {
2180
+ "id": "localai",
2181
+ "name": "LocalAI",
2182
+ "country": "International",
2183
+ "api_endpoint": "http://localhost:8080/v1",
2184
+ "api_style": "openai",
2185
+ "supports_model_listing": true,
2186
+ "supports_streaming": true,
2187
+ "supports_vision": false,
2188
+ "supports_embeddings": false,
2189
+ "supports_audio": false,
2190
+ "supports_reasoning": false,
2191
+ "openai_compatible": true,
2192
+ "free_models_available": true,
2193
+ "paid_models_available": false,
2194
+ "needs_key": false,
2195
+ "description": "Self-hosted OpenAI-compatible API",
2196
+ "last_updated": "2026-07-31T05:17:14+00:00",
2197
+ "fallback_models": [
2198
+ "llama3",
2199
+ "phi-3"
2200
+ ],
2201
+ "default_model": "llama3"
2202
+ },
2203
+ {
2204
+ "id": "luma",
2205
+ "name": "Luma Dream Machine",
2206
+ "country": "United States",
2207
+ "api_endpoint": "https://api.lumalabs.ai/v1",
2208
+ "api_style": "openai",
2209
+ "supports_model_listing": true,
2210
+ "supports_streaming": false,
2211
+ "supports_vision": true,
2212
+ "supports_embeddings": false,
2213
+ "supports_audio": false,
2214
+ "supports_reasoning": false,
2215
+ "openai_compatible": false,
2216
+ "free_models_available": false,
2217
+ "paid_models_available": true,
2218
+ "needs_key": true,
2219
+ "description": "Dream Machine next-generation generative video and 3D NeRFs",
2220
+ "last_updated": "2026-09-01T00:00:00+00:00",
2221
+ "fallback_models": [
2222
+ "dream-machine",
2223
+ "ray-1-6"
2224
+ ],
2225
+ "default_model": "dream-machine"
2226
+ },
2227
+ {
2228
+ "id": "maritaca",
2229
+ "name": "Maritaca AI (Sabia)",
2230
+ "country": "Brazil",
2231
+ "api_endpoint": "https://api.maritaca.ai/v1",
2232
+ "api_style": "openai",
2233
+ "supports_model_listing": true,
2234
+ "supports_streaming": true,
2235
+ "supports_vision": false,
2236
+ "supports_embeddings": false,
2237
+ "supports_audio": false,
2238
+ "supports_reasoning": false,
2239
+ "openai_compatible": true,
2240
+ "free_models_available": true,
2241
+ "paid_models_available": true,
2242
+ "needs_key": true,
2243
+ "description": "Portuguese-language LLM",
2244
+ "last_updated": "2026-07-31T05:17:14+00:00",
2245
+ "fallback_models": [
2246
+ "sabia-3"
2247
+ ],
2248
+ "default_model": "sabia-3"
2249
+ },
2250
+ {
2251
+ "id": "microsoft-azure",
2252
+ "name": "Microsoft Azure OpenAI",
2253
+ "country": "United States",
2254
+ "api_endpoint": "https://YOUR_RESOURCE.openai.azure.com/v1",
2255
+ "api_style": "openai",
2256
+ "supports_model_listing": true,
2257
+ "supports_streaming": true,
2258
+ "supports_vision": false,
2259
+ "supports_embeddings": true,
2260
+ "supports_audio": false,
2261
+ "supports_reasoning": false,
2262
+ "openai_compatible": true,
2263
+ "free_models_available": true,
2264
+ "paid_models_available": true,
2265
+ "needs_key": true,
2266
+ "description": "Azure OpenAI Service",
2267
+ "last_updated": "2026-07-31T05:17:14+00:00",
2268
+ "fallback_models": [
2269
+ "gpt-4o",
2270
+ "gpt-4o-mini"
2271
+ ],
2272
+ "default_model": "gpt-4o-mini"
2273
+ },
2274
+ {
2275
+ "id": "milvus",
2276
+ "name": "Milvus (Zilliz Cloud)",
2277
+ "country": "United States",
2278
+ "api_endpoint": "https://controller.api.cloud.zilliz.com/v1",
2279
+ "api_style": "openai",
2280
+ "supports_model_listing": true,
2281
+ "supports_streaming": false,
2282
+ "supports_vision": false,
2283
+ "supports_embeddings": true,
2284
+ "supports_audio": false,
2285
+ "supports_reasoning": false,
2286
+ "openai_compatible": false,
2287
+ "free_models_available": true,
2288
+ "paid_models_available": true,
2289
+ "needs_key": true,
2290
+ "description": "Massive-scale distributed vector database for enterprise AI",
2291
+ "last_updated": "2026-09-01T00:00:00+00:00",
2292
+ "fallback_models": [
2293
+ "bge-m3",
2294
+ "voyage-code-2"
2295
+ ],
2296
+ "default_model": "bge-m3"
2297
+ },
2298
+ {
2299
+ "id": "minimax",
2300
+ "name": "MiniMax",
2301
+ "country": "China",
2302
+ "api_endpoint": "https://api.minimax.chat/v1",
2303
+ "api_style": "openai",
2304
+ "supports_model_listing": true,
2305
+ "supports_streaming": true,
2306
+ "supports_vision": false,
2307
+ "supports_embeddings": false,
2308
+ "supports_audio": false,
2309
+ "supports_reasoning": false,
2310
+ "openai_compatible": true,
2311
+ "free_models_available": true,
2312
+ "paid_models_available": true,
2313
+ "needs_key": true,
2314
+ "description": "MiniMax — Chinese LLM with long context",
2315
+ "last_updated": "2026-09-03T00:00:00+00:00",
2316
+ "fallback_models": [
2317
+ "abab6.5-chat",
2318
+ "abab5.5-chat",
2319
+ "abab4-chat"
2320
+ ],
2321
+ "default_model": "abab6.5-chat"
2322
+ },
2323
+ {
2324
+ "id": "mistral",
2325
+ "name": "Mistral AI",
2326
+ "country": "France",
2327
+ "api_endpoint": "https://api.mistral.ai/v1",
2328
+ "api_style": "openai",
2329
+ "supports_model_listing": true,
2330
+ "supports_streaming": true,
2331
+ "supports_vision": false,
2332
+ "supports_embeddings": true,
2333
+ "supports_audio": false,
2334
+ "supports_reasoning": true,
2335
+ "openai_compatible": true,
2336
+ "free_models_available": true,
2337
+ "paid_models_available": true,
2338
+ "needs_key": true,
2339
+ "description": "Mistral Large & Codestral — European leader",
2340
+ "last_updated": "2026-07-31T05:17:14+00:00",
2341
+ "fallback_models": [
2342
+ "codestral-2405",
2343
+ "mistral-large-2407",
2344
+ "mistral-large-2411",
2345
+ "mistral-large-latest",
2346
+ "mistral-medium",
2347
+ "mistral-medium-latest",
2348
+ "mistral-small-2402",
2349
+ "mistral-small-latest",
2350
+ "mistral-tiny",
2351
+ "open-mistral-nemo",
2352
+ "open-mixtral-8x22b",
2353
+ "open-mixtral-8x7b"
2354
+ ],
2355
+ "default_model": "open-mistral-nemo"
2356
+ },
2357
+ {
2358
+ "id": "modal",
2359
+ "name": "Modal",
2360
+ "country": "United States",
2361
+ "api_endpoint": "https://api.modal.com/v1",
2362
+ "api_style": "openai",
2363
+ "supports_model_listing": true,
2364
+ "supports_streaming": true,
2365
+ "supports_vision": false,
2366
+ "supports_embeddings": false,
2367
+ "supports_audio": false,
2368
+ "supports_reasoning": false,
2369
+ "openai_compatible": true,
2370
+ "free_models_available": true,
2371
+ "paid_models_available": true,
2372
+ "needs_key": true,
2373
+ "description": "Serverless cloud for AI",
2374
+ "last_updated": "2026-07-31T05:17:14+00:00",
2375
+ "fallback_models": [
2376
+ "llama3-70b"
2377
+ ],
2378
+ "default_model": "llama3-70b"
2379
+ },
2380
+ {
2381
+ "id": "monsterapi",
2382
+ "name": "Monster API",
2383
+ "country": "United States",
2384
+ "api_endpoint": "https://api.monsterapi.ai/v1",
2385
+ "api_style": "openai",
2386
+ "supports_model_listing": true,
2387
+ "supports_streaming": true,
2388
+ "supports_vision": true,
2389
+ "supports_embeddings": false,
2390
+ "supports_audio": false,
2391
+ "supports_reasoning": false,
2392
+ "openai_compatible": true,
2393
+ "free_models_available": true,
2394
+ "paid_models_available": true,
2395
+ "needs_key": true,
2396
+ "description": "Affordable model API",
2397
+ "last_updated": "2026-07-31T05:17:14+00:00",
2398
+ "fallback_models": [
2399
+ "llama3-70b",
2400
+ "mixtral-8x7b",
2401
+ "monster-llama-3.1-8b",
2402
+ "monster-mistral-7b",
2403
+ "monster-whisper-large-v3"
2404
+ ],
2405
+ "default_model": "llama3-70b"
2406
+ },
2407
+ {
2408
+ "id": "moonshot",
2409
+ "name": "Moonshot AI (Kimi)",
2410
+ "country": "China",
2411
+ "api_endpoint": "https://api.moonshot.cn/v1",
2412
+ "api_style": "openai",
2413
+ "supports_model_listing": true,
2414
+ "supports_streaming": true,
2415
+ "supports_vision": false,
2416
+ "supports_embeddings": false,
2417
+ "supports_audio": false,
2418
+ "supports_reasoning": true,
2419
+ "openai_compatible": true,
2420
+ "free_models_available": true,
2421
+ "paid_models_available": true,
2422
+ "needs_key": true,
2423
+ "description": "Kimi K3 — latest long-context Chinese LLM",
2424
+ "last_updated": "2026-09-03T00:00:00+00:00",
2425
+ "fallback_models": [
2426
+ "kimi-k3",
2427
+ "moonshot-v1-128k",
2428
+ "moonshot-v1-32k",
2429
+ "moonshot-v1-8k"
2430
+ ],
2431
+ "default_model": "kimi-k3"
2432
+ },
2433
+ {
2434
+ "id": "ncloud",
2435
+ "name": "Naver Cloud (HyperCLOVA X)",
2436
+ "country": "South Korea",
2437
+ "api_endpoint": "https://clovastudio.apigw.ntruss.com/v1",
2438
+ "api_style": "openai",
2439
+ "supports_model_listing": true,
2440
+ "supports_streaming": true,
2441
+ "supports_vision": false,
2442
+ "supports_embeddings": false,
2443
+ "supports_audio": false,
2444
+ "supports_reasoning": false,
2445
+ "openai_compatible": true,
2446
+ "free_models_available": true,
2447
+ "paid_models_available": true,
2448
+ "needs_key": true,
2449
+ "description": "HyperCLOVA X — Naver's enterprise LLM",
2450
+ "last_updated": "2026-09-03T00:00:00+00:00",
2451
+ "fallback_models": [
2452
+ "HCX-005",
2453
+ "HCX-003"
2454
+ ],
2455
+ "default_model": "HCX-005"
2456
+ },
2457
+ {
2458
+ "id": "nebius",
2459
+ "name": "Nebius AI",
2460
+ "country": "Netherlands",
2461
+ "api_endpoint": "https://api.nebius.ai/v1",
2462
+ "api_style": "openai",
2463
+ "supports_model_listing": true,
2464
+ "supports_streaming": true,
2465
+ "supports_vision": false,
2466
+ "supports_embeddings": false,
2467
+ "supports_audio": false,
2468
+ "supports_reasoning": false,
2469
+ "openai_compatible": true,
2470
+ "free_models_available": true,
2471
+ "paid_models_available": true,
2472
+ "needs_key": true,
2473
+ "description": "European GPU cloud",
2474
+ "last_updated": "2026-07-31T05:17:14+00:00",
2475
+ "fallback_models": [
2476
+ "llama3-70b",
2477
+ "qwen2-72b"
2478
+ ],
2479
+ "default_model": "llama3-70b"
2480
+ },
2481
+ {
2482
+ "id": "neets",
2483
+ "name": "Neets AI",
2484
+ "country": "United States",
2485
+ "api_endpoint": "https://api.neets.ai/v1",
2486
+ "api_style": "openai",
2487
+ "supports_model_listing": true,
2488
+ "supports_streaming": true,
2489
+ "supports_vision": false,
2490
+ "supports_embeddings": false,
2491
+ "supports_audio": false,
2492
+ "supports_reasoning": false,
2493
+ "openai_compatible": true,
2494
+ "free_models_available": true,
2495
+ "paid_models_available": true,
2496
+ "needs_key": true,
2497
+ "description": "API gateway with multiple backends",
2498
+ "last_updated": "2026-07-31T05:17:14+00:00",
2499
+ "fallback_models": [
2500
+ "gpt-4o",
2501
+ "llama3-70b"
2502
+ ],
2503
+ "default_model": "gpt-4o"
2504
+ },
2505
+ {
2506
+ "id": "neuralmagic",
2507
+ "name": "Neural Magic",
2508
+ "country": "United States",
2509
+ "api_endpoint": "https://api.neuralmagic.com/v1",
2510
+ "api_style": "openai",
2511
+ "supports_model_listing": true,
2512
+ "supports_streaming": true,
2513
+ "supports_vision": false,
2514
+ "supports_embeddings": false,
2515
+ "supports_audio": false,
2516
+ "supports_reasoning": false,
2517
+ "openai_compatible": true,
2518
+ "free_models_available": true,
2519
+ "paid_models_available": true,
2520
+ "needs_key": true,
2521
+ "description": "CPU-optimized inference",
2522
+ "last_updated": "2026-07-31T05:17:14+00:00",
2523
+ "fallback_models": [
2524
+ "llama3-70b"
2525
+ ],
2526
+ "default_model": "llama3-70b"
2527
+ },
2528
+ {
2529
+ "id": "nomic",
2530
+ "name": "Nomic AI",
2531
+ "country": "United States",
2532
+ "api_endpoint": "https://api.nomic.ai/v1",
2533
+ "api_style": "openai",
2534
+ "supports_model_listing": true,
2535
+ "supports_streaming": false,
2536
+ "supports_vision": false,
2537
+ "supports_embeddings": true,
2538
+ "supports_audio": false,
2539
+ "supports_reasoning": false,
2540
+ "openai_compatible": true,
2541
+ "free_models_available": true,
2542
+ "paid_models_available": true,
2543
+ "needs_key": true,
2544
+ "description": "Embedding models & GPT4All",
2545
+ "last_updated": "2026-07-31T05:17:14+00:00",
2546
+ "fallback_models": [
2547
+ "nomic-embed-text-v1"
2548
+ ],
2549
+ "default_model": "nomic-embed-text-v1"
2550
+ },
2551
+ {
2552
+ "id": "notebooklm",
2553
+ "name": "NotebookLM",
2554
+ "country": "United States",
2555
+ "api_endpoint": "https://notebooklm.google.com",
2556
+ "api_style": "openai",
2557
+ "supports_model_listing": true,
2558
+ "supports_streaming": true,
2559
+ "supports_vision": false,
2560
+ "supports_embeddings": false,
2561
+ "supports_audio": false,
2562
+ "supports_reasoning": false,
2563
+ "openai_compatible": false,
2564
+ "free_models_available": true,
2565
+ "paid_models_available": false,
2566
+ "needs_key": false,
2567
+ "description": "Google research assistant",
2568
+ "last_updated": "2026-07-31T05:17:14+00:00",
2569
+ "fallback_models": [
2570
+ "default"
2571
+ ],
2572
+ "default_model": "default"
2573
+ },
2574
+ {
2575
+ "id": "novita",
2576
+ "name": "Novita AI",
2577
+ "country": "United States",
2578
+ "api_endpoint": "https://api.novita.ai/v3/openai",
2579
+ "api_style": "openai",
2580
+ "supports_model_listing": true,
2581
+ "supports_streaming": true,
2582
+ "supports_vision": false,
2583
+ "supports_embeddings": false,
2584
+ "supports_audio": false,
2585
+ "supports_reasoning": false,
2586
+ "openai_compatible": true,
2587
+ "free_models_available": true,
2588
+ "paid_models_available": true,
2589
+ "needs_key": true,
2590
+ "description": "Novita AI — Cost-effective inference",
2591
+ "last_updated": "2026-09-03T00:00:00+00:00",
2592
+ "fallback_models": [
2593
+ "meta-llama/llama-3.1-70b-instruct",
2594
+ "meta-llama/llama-3.1-8b-instruct",
2595
+ "mistralai/mixtral-8x7b-instruct"
2596
+ ],
2597
+ "default_model": "meta-llama/llama-3.1-70b-instruct"
2598
+ },
2599
+ {
2600
+ "id": "nvidia",
2601
+ "name": "NVIDIA NIM",
2602
+ "country": "United States",
2603
+ "api_endpoint": "https://integrate.api.nvidia.com/v1",
2604
+ "api_style": "openai",
2605
+ "supports_model_listing": true,
2606
+ "supports_streaming": true,
2607
+ "supports_vision": false,
2608
+ "supports_embeddings": false,
2609
+ "supports_audio": false,
2610
+ "supports_reasoning": false,
2611
+ "openai_compatible": true,
2612
+ "free_models_available": true,
2613
+ "paid_models_available": true,
2614
+ "needs_key": true,
2615
+ "description": "NVIDIA Inference Microservices",
2616
+ "last_updated": "2026-07-31T05:17:14+00:00",
2617
+ "fallback_models": [
2618
+ "meta/llama-3.2-11b-vision-instruct",
2619
+ "mistralai/mistral-large-2-instruct",
2620
+ "nvidia/llama-3.1-nemotron-70b-instruct",
2621
+ "ibm/granite-3.0-8b-instruct"
2622
+ ],
2623
+ "default_model": "meta/llama-3.2-11b-vision-instruct"
2624
+ },
2625
+ {
2626
+ "id": "oci-genai",
2627
+ "name": "OCI Generative AI",
2628
+ "country": "United States",
2629
+ "api_endpoint": "https://inference.generativeai.us-chicago-1.oci.oraclecloud.com",
2630
+ "api_style": "openai",
2631
+ "supports_model_listing": true,
2632
+ "supports_streaming": true,
2633
+ "supports_vision": false,
2634
+ "supports_embeddings": false,
2635
+ "supports_audio": false,
2636
+ "supports_reasoning": false,
2637
+ "openai_compatible": true,
2638
+ "free_models_available": true,
2639
+ "paid_models_available": true,
2640
+ "needs_key": true,
2641
+ "description": "Oracle Cloud — managed AI service",
2642
+ "last_updated": "2026-07-31T05:17:14+00:00",
2643
+ "fallback_models": [
2644
+ "cohere.command-r-plus",
2645
+ "meta.llama-3-70b-instruct"
2646
+ ],
2647
+ "default_model": "cohere.command-r-plus"
2648
+ },
2649
+ {
2650
+ "id": "octoai",
2651
+ "name": "OctoAI",
2652
+ "country": "United States",
2653
+ "api_endpoint": "https://text.octoai.run/v1",
2654
+ "api_style": "openai",
2655
+ "supports_model_listing": true,
2656
+ "supports_streaming": true,
2657
+ "supports_vision": true,
2658
+ "supports_embeddings": true,
2659
+ "supports_audio": false,
2660
+ "supports_reasoning": true,
2661
+ "openai_compatible": true,
2662
+ "free_models_available": false,
2663
+ "paid_models_available": true,
2664
+ "needs_key": true,
2665
+ "description": "High-performance inference engine for text, vision, and custom open-source models",
2666
+ "last_updated": "2026-09-01T00:00:00+00:00",
2667
+ "fallback_models": [
2668
+ "meta-llama-3.1-70b-instruct",
2669
+ "mixtral-8x22b-instruct",
2670
+ "hermes-2-pro-mistral-7b"
2671
+ ],
2672
+ "default_model": "meta-llama-3.1-70b-instruct"
2673
+ },
2674
+ {
2675
+ "id": "ollama",
2676
+ "name": "Ollama (Local)",
2677
+ "country": "International",
2678
+ "api_endpoint": "http://localhost:11434",
2679
+ "api_style": "ollama",
2680
+ "supports_model_listing": true,
2681
+ "supports_streaming": true,
2682
+ "supports_vision": false,
2683
+ "supports_embeddings": false,
2684
+ "supports_audio": false,
2685
+ "supports_reasoning": false,
2686
+ "openai_compatible": true,
2687
+ "free_models_available": true,
2688
+ "paid_models_available": false,
2689
+ "needs_key": false,
2690
+ "description": "Local open-source model runner with auto-discovery",
2691
+ "last_updated": "2026-09-03T00:00:00+00:00",
2692
+ "supports_custom_models": true,
2693
+ "auto_discover_models": true,
2694
+ "fallback_models": [
2695
+ "codellama",
2696
+ "deepseek-r1",
2697
+ "gemma2",
2698
+ "gemma2:latest",
2699
+ "llama3",
2700
+ "llama3.1",
2701
+ "llama3.2",
2702
+ "llama3.3",
2703
+ "llama3.3:latest",
2704
+ "mistral",
2705
+ "mixtral",
2706
+ "phi3",
2707
+ "phi4",
2708
+ "qwen2.5",
2709
+ "qwen2.5:latest",
2710
+ "qwen3",
2711
+ "command-r",
2712
+ "dolphin-mistral",
2713
+ "neural-chat",
2714
+ "starling-lm",
2715
+ "vicuna",
2716
+ "yi",
2717
+ "openhermes",
2718
+ "custom-model"
2719
+ ],
2720
+ "default_model": "llama3.3"
2721
+ },
2722
+ {
2723
+ "id": "ollama_cloud",
2724
+ "name": "Ollama Cloud (Hosted)",
2725
+ "country": "International",
2726
+ "api_endpoint": "https://cloud.ollama.ai/v1",
2727
+ "api_style": "ollama",
2728
+ "supports_model_listing": true,
2729
+ "supports_streaming": true,
2730
+ "supports_vision": true,
2731
+ "supports_embeddings": true,
2732
+ "supports_audio": false,
2733
+ "supports_reasoning": true,
2734
+ "openai_compatible": true,
2735
+ "free_models_available": false,
2736
+ "paid_models_available": true,
2737
+ "needs_key": true,
2738
+ "description": "Cloud-managed Ollama instances providing high-availability GPU endpoints",
2739
+ "last_updated": "2026-09-01T00:00:00+00:00",
2740
+ "fallback_models": [
2741
+ "llama3.3:70b",
2742
+ "deepseek-r1:70b",
2743
+ "qwen2.5:72b"
2744
+ ],
2745
+ "default_model": "llama3.3:70b"
2746
+ },
2747
+ {
2748
+ "id": "openai",
2749
+ "name": "OpenAI",
2750
+ "country": "United States",
2751
+ "api_endpoint": "https://api.openai.com/v1",
2752
+ "api_style": "openai",
2753
+ "supports_model_listing": true,
2754
+ "supports_streaming": true,
2755
+ "supports_vision": true,
2756
+ "supports_embeddings": true,
2757
+ "supports_audio": true,
2758
+ "supports_reasoning": true,
2759
+ "openai_compatible": true,
2760
+ "free_models_available": true,
2761
+ "paid_models_available": true,
2762
+ "needs_key": true,
2763
+ "description": "GPT-4o & o1 — industry standard",
2764
+ "last_updated": "2026-09-03T00:00:00+00:00",
2765
+ "fallback_models": [
2766
+ "dall-e-3",
2767
+ "gpt-3.5-turbo",
2768
+ "gpt-4",
2769
+ "gpt-4-turbo",
2770
+ "gpt-4.1",
2771
+ "gpt-4.1-mini",
2772
+ "gpt-4.1-nano",
2773
+ "gpt-4o",
2774
+ "gpt-4o-2024-08-06",
2775
+ "gpt-4o-mini",
2776
+ "gpt-4o-mini-2024-07-18",
2777
+ "o1",
2778
+ "o1-mini",
2779
+ "o3",
2780
+ "o3-mini",
2781
+ "o4-mini",
2782
+ "text-embedding-3-large",
2783
+ "text-embedding-3-small",
2784
+ "tts-1",
2785
+ "whisper-1"
2786
+ ],
2787
+ "default_model": "gpt-4o"
2788
+ },
2789
+ {
2790
+ "id": "openpipe",
2791
+ "name": "OpenPipe",
2792
+ "country": "United States",
2793
+ "api_endpoint": "https://app.openpipe.ai/api/v1",
2794
+ "api_style": "openai",
2795
+ "supports_model_listing": true,
2796
+ "supports_streaming": true,
2797
+ "supports_vision": false,
2798
+ "supports_embeddings": false,
2799
+ "supports_audio": false,
2800
+ "supports_reasoning": true,
2801
+ "openai_compatible": true,
2802
+ "free_models_available": false,
2803
+ "paid_models_available": true,
2804
+ "needs_key": true,
2805
+ "description": "Continuous model fine-tuning, distillation, and optimized inference for production apps",
2806
+ "last_updated": "2026-09-01T00:00:00+00:00",
2807
+ "fallback_models": [
2808
+ "openpipe:llama-3-8b-instruct",
2809
+ "openpipe:mistral-7b-instruct"
2810
+ ],
2811
+ "default_model": "openpipe:llama-3-8b-instruct"
2812
+ },
2813
+ {
2814
+ "id": "openrouter",
2815
+ "name": "OpenRouter",
2816
+ "country": "United States",
2817
+ "api_endpoint": "https://openrouter.ai/api/v1",
2818
+ "api_style": "openai",
2819
+ "supports_model_listing": true,
2820
+ "supports_streaming": true,
2821
+ "supports_vision": true,
2822
+ "supports_embeddings": false,
2823
+ "supports_audio": false,
2824
+ "supports_reasoning": false,
2825
+ "openai_compatible": true,
2826
+ "free_models_available": true,
2827
+ "paid_models_available": true,
2828
+ "needs_key": true,
2829
+ "description": "Multi-provider router up to 200+ models",
2830
+ "last_updated": "2026-07-31T05:17:14+00:00",
2831
+ "fallback_models": [
2832
+ "anthropic/claude-3.5-sonnet",
2833
+ "anthropic/claude-sonnet-4",
2834
+ "deepseek/deepseek-chat",
2835
+ "google/gemini-2.0-flash",
2836
+ "meta-llama/llama-3.1-70b-instruct",
2837
+ "meta-llama/llama-3.3-70b-instruct",
2838
+ "mistralai/mistral-large",
2839
+ "openai/gpt-4o",
2840
+ "openai/gpt-4o-mini"
2841
+ ],
2842
+ "default_model": "google/gemini-2.0-flash"
2843
+ },
2844
+ {
2845
+ "id": "oracle-oci",
2846
+ "name": "Oracle OCI Generative AI",
2847
+ "country": "United States",
2848
+ "api_endpoint": "https://inference.generativeai.us-ashburn-1.oci.oraclecloud.com",
2849
+ "api_style": "openai",
2850
+ "supports_model_listing": true,
2851
+ "supports_streaming": true,
2852
+ "supports_vision": false,
2853
+ "supports_embeddings": false,
2854
+ "supports_audio": false,
2855
+ "supports_reasoning": false,
2856
+ "openai_compatible": true,
2857
+ "free_models_available": false,
2858
+ "paid_models_available": true,
2859
+ "needs_key": true,
2860
+ "description": "Oracle OCI — Enterprise AI inference",
2861
+ "last_updated": "2026-09-03T00:00:00+00:00",
2862
+ "fallback_models": [
2863
+ "meta-llama-3.1-70b-instruct",
2864
+ "meta-llama-3.1-8b-instruct"
2865
+ ],
2866
+ "default_model": "meta-llama-3.1-70b-instruct"
2867
+ },
2868
+ {
2869
+ "id": "ovhcloud",
2870
+ "name": "OVHcloud AI Endpoints",
2871
+ "country": "France",
2872
+ "api_endpoint": "https://endpoints.ai.cloud.ovh.net/v1",
2873
+ "api_style": "openai",
2874
+ "supports_model_listing": true,
2875
+ "supports_streaming": true,
2876
+ "supports_vision": true,
2877
+ "supports_embeddings": true,
2878
+ "supports_audio": false,
2879
+ "supports_reasoning": true,
2880
+ "openai_compatible": true,
2881
+ "free_models_available": true,
2882
+ "paid_models_available": true,
2883
+ "needs_key": true,
2884
+ "description": "European sovereign cloud AI endpoints with complete data privacy and GDPR compliance",
2885
+ "last_updated": "2026-09-01T00:00:00+00:00",
2886
+ "fallback_models": [
2887
+ "Meta-Llama-3-1-70B-Instruct",
2888
+ "Mistral-7B-Instruct-v0.3",
2889
+ "bge-large-en-v1.5"
2890
+ ],
2891
+ "default_model": "Meta-Llama-3-1-70B-Instruct"
2892
+ },
2893
+ {
2894
+ "id": "perplexity",
2895
+ "name": "Perplexity AI",
2896
+ "country": "United States",
2897
+ "api_endpoint": "https://api.perplexity.ai",
2898
+ "api_style": "openai",
2899
+ "supports_model_listing": true,
2900
+ "supports_streaming": true,
2901
+ "supports_vision": false,
2902
+ "supports_embeddings": false,
2903
+ "supports_audio": false,
2904
+ "supports_reasoning": false,
2905
+ "openai_compatible": true,
2906
+ "free_models_available": true,
2907
+ "paid_models_available": true,
2908
+ "needs_key": true,
2909
+ "description": "Sonar — AI search & chat",
2910
+ "last_updated": "2026-07-31T05:17:14+00:00",
2911
+ "fallback_models": [
2912
+ "llama-3.1-sonar-huge-128k",
2913
+ "llama-3.1-sonar-large-128k",
2914
+ "llama-3.1-sonar-small-128k",
2915
+ "sonar",
2916
+ "sonar-pro",
2917
+ "sonar-reasoning",
2918
+ "sonar-reasoning-pro",
2919
+ "sonar-small-chat"
2920
+ ],
2921
+ "default_model": "sonar-reasoning-pro"
2922
+ },
2923
+ {
2924
+ "id": "petals",
2925
+ "name": "Petals",
2926
+ "country": "International",
2927
+ "api_endpoint": "https://petals.dev/v1",
2928
+ "api_style": "openai",
2929
+ "supports_model_listing": true,
2930
+ "supports_streaming": true,
2931
+ "supports_vision": false,
2932
+ "supports_embeddings": false,
2933
+ "supports_audio": false,
2934
+ "supports_reasoning": false,
2935
+ "openai_compatible": true,
2936
+ "free_models_available": true,
2937
+ "paid_models_available": false,
2938
+ "needs_key": false,
2939
+ "description": "Decentralized inference network",
2940
+ "last_updated": "2026-07-31T05:17:14+00:00",
2941
+ "fallback_models": [
2942
+ "llama3-70b"
2943
+ ],
2944
+ "default_model": "llama3-70b"
2945
+ },
2946
+ {
2947
+ "id": "phind",
2948
+ "name": "Phind",
2949
+ "country": "United States",
2950
+ "api_endpoint": "https://api.phind.com/v1",
2951
+ "api_style": "openai",
2952
+ "supports_model_listing": true,
2953
+ "supports_streaming": true,
2954
+ "supports_vision": false,
2955
+ "supports_embeddings": false,
2956
+ "supports_audio": false,
2957
+ "supports_reasoning": false,
2958
+ "openai_compatible": true,
2959
+ "free_models_available": true,
2960
+ "paid_models_available": true,
2961
+ "needs_key": true,
2962
+ "description": "AI search & code assistant",
2963
+ "last_updated": "2026-07-31T05:17:14+00:00",
2964
+ "fallback_models": [
2965
+ "phind-coder",
2966
+ "phind-v2"
2967
+ ],
2968
+ "default_model": "phind-coder"
2969
+ },
2970
+ {
2971
+ "id": "pinecone",
2972
+ "name": "Pinecone AI",
2973
+ "country": "United States",
2974
+ "api_endpoint": "https://api.pinecone.io",
2975
+ "api_style": "openai",
2976
+ "supports_model_listing": true,
2977
+ "supports_streaming": false,
2978
+ "supports_vision": false,
2979
+ "supports_embeddings": true,
2980
+ "supports_audio": false,
2981
+ "supports_reasoning": false,
2982
+ "openai_compatible": false,
2983
+ "free_models_available": true,
2984
+ "paid_models_available": true,
2985
+ "needs_key": true,
2986
+ "description": "Managed vector database and semantic inference service",
2987
+ "last_updated": "2026-09-01T00:00:00+00:00",
2988
+ "fallback_models": [
2989
+ "multilingual-e5-large",
2990
+ "pinecone-rerank-v0"
2991
+ ],
2992
+ "default_model": "multilingual-e5-large"
2993
+ },
2994
+ {
2995
+ "id": "playht",
2996
+ "name": "PlayHT Voice",
2997
+ "country": "United States",
2998
+ "api_endpoint": "https://api.play.ht/api/v2",
2999
+ "api_style": "openai",
3000
+ "supports_model_listing": true,
3001
+ "supports_streaming": true,
3002
+ "supports_vision": false,
3003
+ "supports_embeddings": false,
3004
+ "supports_audio": true,
3005
+ "supports_reasoning": false,
3006
+ "openai_compatible": false,
3007
+ "free_models_available": true,
3008
+ "paid_models_available": true,
3009
+ "needs_key": true,
3010
+ "description": "Conversational AI voice generation and real-time TTS",
3011
+ "last_updated": "2026-09-01T00:00:00+00:00",
3012
+ "fallback_models": [
3013
+ "Play3.0-mini",
3014
+ "PlayDialog"
3015
+ ],
3016
+ "default_model": "Play3.0-mini"
3017
+ },
3018
+ {
3019
+ "id": "poe",
3020
+ "name": "Poe (Quora)",
3021
+ "country": "United States",
3022
+ "api_endpoint": "https://api.poe.com/v1",
3023
+ "api_style": "openai",
3024
+ "supports_model_listing": true,
3025
+ "supports_streaming": true,
3026
+ "supports_vision": true,
3027
+ "supports_embeddings": false,
3028
+ "supports_audio": false,
3029
+ "supports_reasoning": true,
3030
+ "openai_compatible": true,
3031
+ "free_models_available": true,
3032
+ "paid_models_available": true,
3033
+ "needs_key": true,
3034
+ "description": "Quora Poe API protocol for frontier and specialized AI bot interaction",
3035
+ "last_updated": "2026-09-01T00:00:00+00:00",
3036
+ "fallback_models": [
3037
+ "Claude-3.5-Sonnet",
3038
+ "GPT-4o",
3039
+ "DeepSeek-R1"
3040
+ ],
3041
+ "default_model": "Claude-3.5-Sonnet"
3042
+ },
3043
+ {
3044
+ "id": "predibase",
3045
+ "name": "Predibase",
3046
+ "country": "United States",
3047
+ "api_endpoint": "https://api.predibase.com/v1",
3048
+ "api_style": "openai",
3049
+ "supports_model_listing": true,
3050
+ "supports_streaming": true,
3051
+ "supports_vision": false,
3052
+ "supports_embeddings": false,
3053
+ "supports_audio": false,
3054
+ "supports_reasoning": true,
3055
+ "openai_compatible": true,
3056
+ "free_models_available": true,
3057
+ "paid_models_available": true,
3058
+ "needs_key": true,
3059
+ "description": "Serverless fine-tuned & open-weight LLM inference with LoRAX engine",
3060
+ "last_updated": "2026-09-01T00:00:00+00:00",
3061
+ "fallback_models": [
3062
+ "llama-3.3-70b-instruct",
3063
+ "mistral-7b-instruct",
3064
+ "qwen-2.5-coder-32b"
3065
+ ],
3066
+ "default_model": "llama-3.3-70b-instruct"
3067
+ },
3068
+ {
3069
+ "id": "premai",
3070
+ "name": "Prem AI",
3071
+ "country": "United States",
3072
+ "api_endpoint": "https://api.premai.io/v1",
3073
+ "api_style": "openai",
3074
+ "supports_model_listing": true,
3075
+ "supports_streaming": true,
3076
+ "supports_vision": false,
3077
+ "supports_embeddings": false,
3078
+ "supports_audio": false,
3079
+ "supports_reasoning": false,
3080
+ "openai_compatible": true,
3081
+ "free_models_available": true,
3082
+ "paid_models_available": true,
3083
+ "needs_key": true,
3084
+ "description": "Enterprise AI gateway",
3085
+ "last_updated": "2026-07-31T05:17:14+00:00",
3086
+ "fallback_models": [
3087
+ "llama3-70b"
3088
+ ],
3089
+ "default_model": "llama3-70b"
3090
+ },
3091
+ {
3092
+ "id": "qdrant",
3093
+ "name": "Qdrant Vector AI",
3094
+ "country": "Germany",
3095
+ "api_endpoint": "https://api.qdrant.io",
3096
+ "api_style": "openai",
3097
+ "supports_model_listing": true,
3098
+ "supports_streaming": false,
3099
+ "supports_vision": false,
3100
+ "supports_embeddings": true,
3101
+ "supports_audio": false,
3102
+ "supports_reasoning": false,
3103
+ "openai_compatible": false,
3104
+ "free_models_available": true,
3105
+ "paid_models_available": true,
3106
+ "needs_key": true,
3107
+ "description": "High-performance vector search engine and hybrid neural search",
3108
+ "last_updated": "2026-09-01T00:00:00+00:00",
3109
+ "fallback_models": [
3110
+ "bge-large-en-v1.5",
3111
+ "fastembed-hybrid"
3112
+ ],
3113
+ "default_model": "bge-large-en-v1.5"
3114
+ },
3115
+ {
3116
+ "id": "reka",
3117
+ "name": "Reka AI",
3118
+ "country": "United Kingdom",
3119
+ "api_endpoint": "https://api.reka.ai/v1",
3120
+ "api_style": "openai",
3121
+ "supports_model_listing": true,
3122
+ "supports_streaming": true,
3123
+ "supports_vision": true,
3124
+ "supports_embeddings": false,
3125
+ "supports_audio": false,
3126
+ "supports_reasoning": false,
3127
+ "openai_compatible": true,
3128
+ "free_models_available": true,
3129
+ "paid_models_available": true,
3130
+ "needs_key": true,
3131
+ "description": "Reka Core — multimodal AI",
3132
+ "last_updated": "2026-07-31T05:17:14+00:00",
3133
+ "fallback_models": [
3134
+ "reka-core",
3135
+ "reka-flash"
3136
+ ],
3137
+ "default_model": "reka-core"
3138
+ },
3139
+ {
3140
+ "id": "replicate",
3141
+ "name": "Replicate",
3142
+ "country": "United States",
3143
+ "api_endpoint": "https://api.replicate.com/v1",
3144
+ "api_style": "openai",
3145
+ "supports_model_listing": true,
3146
+ "supports_streaming": true,
3147
+ "supports_vision": true,
3148
+ "supports_embeddings": false,
3149
+ "supports_audio": false,
3150
+ "supports_reasoning": false,
3151
+ "openai_compatible": true,
3152
+ "free_models_available": false,
3153
+ "paid_models_available": true,
3154
+ "needs_key": true,
3155
+ "description": "Replicate — Run open-source models",
3156
+ "last_updated": "2026-09-03T00:00:00+00:00",
3157
+ "fallback_models": [
3158
+ "meta/meta-llama-3.1-405b-instruct",
3159
+ "meta/meta-llama-3.1-70b-instruct",
3160
+ "mistralai/mixtral-8x22b-instruct-v0.1"
3161
+ ],
3162
+ "default_model": "meta/meta-llama-3.1-70b-instruct"
3163
+ },
3164
+ {
3165
+ "id": "runpod",
3166
+ "name": "RunPod",
3167
+ "country": "United States",
3168
+ "api_endpoint": "https://api.runpod.ai/v1",
3169
+ "api_style": "openai",
3170
+ "supports_model_listing": true,
3171
+ "supports_streaming": true,
3172
+ "supports_vision": false,
3173
+ "supports_embeddings": false,
3174
+ "supports_audio": false,
3175
+ "supports_reasoning": false,
3176
+ "openai_compatible": true,
3177
+ "free_models_available": true,
3178
+ "paid_models_available": true,
3179
+ "needs_key": true,
3180
+ "description": "Serverless GPU cloud",
3181
+ "last_updated": "2026-07-31T05:17:14+00:00",
3182
+ "fallback_models": [
3183
+ "llama3-70b"
3184
+ ],
3185
+ "default_model": "llama3-70b"
3186
+ },
3187
+ {
3188
+ "id": "runway",
3189
+ "name": "Runway Gen-3",
3190
+ "country": "United States",
3191
+ "api_endpoint": "https://api.runwayml.com/v1",
3192
+ "api_style": "openai",
3193
+ "supports_model_listing": true,
3194
+ "supports_streaming": false,
3195
+ "supports_vision": true,
3196
+ "supports_embeddings": false,
3197
+ "supports_audio": false,
3198
+ "supports_reasoning": false,
3199
+ "openai_compatible": false,
3200
+ "free_models_available": false,
3201
+ "paid_models_available": true,
3202
+ "needs_key": true,
3203
+ "description": "Gen-3 Alpha photorealistic video generation and world models",
3204
+ "last_updated": "2026-09-01T00:00:00+00:00",
3205
+ "fallback_models": [
3206
+ "gen3a_turbo",
3207
+ "gen2"
3208
+ ],
3209
+ "default_model": "gen3a_turbo"
3210
+ },
3211
+ {
3212
+ "id": "sagify",
3213
+ "name": "Sagify AI",
3214
+ "country": "United States",
3215
+ "api_endpoint": "https://api.sagify.ai/v1",
3216
+ "api_style": "openai",
3217
+ "supports_model_listing": true,
3218
+ "supports_streaming": true,
3219
+ "supports_vision": false,
3220
+ "supports_embeddings": false,
3221
+ "supports_audio": false,
3222
+ "supports_reasoning": false,
3223
+ "openai_compatible": true,
3224
+ "free_models_available": true,
3225
+ "paid_models_available": true,
3226
+ "needs_key": true,
3227
+ "description": "ML training & deployment",
3228
+ "last_updated": "2026-07-31T05:17:14+00:00",
3229
+ "fallback_models": [
3230
+ "default"
3231
+ ],
3232
+ "default_model": "default"
3233
+ },
3234
+ {
3235
+ "id": "sakana_ai",
3236
+ "name": "Sakana AI",
3237
+ "country": "Japan",
3238
+ "api_endpoint": "https://api.sakana.ai/v1",
3239
+ "api_style": "openai",
3240
+ "supports_model_listing": true,
3241
+ "supports_streaming": true,
3242
+ "supports_vision": true,
3243
+ "supports_embeddings": false,
3244
+ "supports_audio": false,
3245
+ "supports_reasoning": true,
3246
+ "openai_compatible": true,
3247
+ "free_models_available": false,
3248
+ "paid_models_available": true,
3249
+ "needs_key": true,
3250
+ "description": "Evolutionary model merging and nature-inspired foundation intelligence (EvoLLM, The AI Scientist)",
3251
+ "last_updated": "2026-09-01T00:00:00+00:00",
3252
+ "fallback_models": [
3253
+ "evo-llm-jp",
3254
+ "ai-scientist-v1",
3255
+ "disco-llm"
3256
+ ],
3257
+ "default_model": "evo-llm-jp"
3258
+ },
3259
+ {
3260
+ "id": "sambanova",
3261
+ "name": "SambaNova",
3262
+ "country": "United States",
3263
+ "api_endpoint": "https://api.sambanova.ai/v1",
3264
+ "api_style": "openai",
3265
+ "supports_model_listing": true,
3266
+ "supports_streaming": true,
3267
+ "supports_vision": false,
3268
+ "supports_embeddings": false,
3269
+ "supports_audio": false,
3270
+ "supports_reasoning": false,
3271
+ "openai_compatible": true,
3272
+ "free_models_available": true,
3273
+ "paid_models_available": true,
3274
+ "needs_key": true,
3275
+ "description": "SN40L Reconfigurable Dataflow — fast inference",
3276
+ "last_updated": "2026-07-31T05:17:14+00:00",
3277
+ "fallback_models": [
3278
+ "llama3.1-70b",
3279
+ "llama3.1-8b",
3280
+ "qwen2.5-72b"
3281
+ ],
3282
+ "default_model": "llama3.1-8b"
3283
+ },
3284
+ {
3285
+ "id": "scale_ai",
3286
+ "name": "Scale AI (SEAL)",
3287
+ "country": "United States",
3288
+ "api_endpoint": "https://api.scale.com/v1",
3289
+ "api_style": "openai",
3290
+ "supports_model_listing": true,
3291
+ "supports_streaming": true,
3292
+ "supports_vision": true,
3293
+ "supports_embeddings": true,
3294
+ "supports_audio": false,
3295
+ "supports_reasoning": true,
3296
+ "openai_compatible": true,
3297
+ "free_models_available": false,
3298
+ "paid_models_available": true,
3299
+ "needs_key": true,
3300
+ "description": "Enterprise generative AI platform, frontier model evaluations and SEAL inference",
3301
+ "last_updated": "2026-09-01T00:00:00+00:00",
3302
+ "fallback_models": [
3303
+ "seal-reasoner-v1",
3304
+ "seal-code-70b",
3305
+ "scale-llama-3-70b"
3306
+ ],
3307
+ "default_model": "seal-reasoner-v1"
3308
+ },
3309
+ {
3310
+ "id": "scaleway",
3311
+ "name": "Scaleway Generative APIs",
3312
+ "country": "France",
3313
+ "api_endpoint": "https://api.scaleway.ai/v1",
3314
+ "api_style": "openai",
3315
+ "supports_model_listing": true,
3316
+ "supports_streaming": true,
3317
+ "supports_vision": true,
3318
+ "supports_embeddings": true,
3319
+ "supports_audio": false,
3320
+ "supports_reasoning": true,
3321
+ "openai_compatible": true,
3322
+ "free_models_available": false,
3323
+ "paid_models_available": true,
3324
+ "needs_key": true,
3325
+ "description": "European sovereign cloud infrastructure and hosted open-weight models",
3326
+ "last_updated": "2026-09-01T00:00:00+00:00",
3327
+ "fallback_models": [
3328
+ "llama-3.1-70b-instruct",
3329
+ "mistral-large-2407",
3330
+ "qwen2.5-coder-32b-instruct"
3331
+ ],
3332
+ "default_model": "llama-3.1-70b-instruct"
3333
+ },
3334
+ {
3335
+ "id": "sensetime",
3336
+ "name": "SenseTime (SenseNova)",
3337
+ "country": "China",
3338
+ "api_endpoint": "https://api.sensenova.cn/v1",
3339
+ "api_style": "openai",
3340
+ "supports_model_listing": true,
3341
+ "supports_streaming": true,
3342
+ "supports_vision": true,
3343
+ "supports_embeddings": true,
3344
+ "supports_audio": true,
3345
+ "supports_reasoning": true,
3346
+ "openai_compatible": true,
3347
+ "free_models_available": true,
3348
+ "paid_models_available": true,
3349
+ "needs_key": true,
3350
+ "description": "SenseNova multimodal foundation model family by SenseTime",
3351
+ "last_updated": "2026-09-01T00:00:00+00:00",
3352
+ "fallback_models": [
3353
+ "SenseChat-5",
3354
+ "SenseNova-Omni",
3355
+ "SenseChat-Vision",
3356
+ "nova-5.5-chat",
3357
+ "nova-5-chat",
3358
+ "nova-4-chat"
3359
+ ],
3360
+ "default_model": "SenseChat-5"
3361
+ },
3362
+ {
3363
+ "id": "seven",
3364
+ "name": "Seven AI",
3365
+ "country": "Germany",
3366
+ "api_endpoint": "https://api.seven.ai/v1",
3367
+ "api_style": "openai",
3368
+ "supports_model_listing": true,
3369
+ "supports_streaming": true,
3370
+ "supports_vision": false,
3371
+ "supports_embeddings": false,
3372
+ "supports_audio": false,
3373
+ "supports_reasoning": false,
3374
+ "openai_compatible": true,
3375
+ "free_models_available": true,
3376
+ "paid_models_available": true,
3377
+ "needs_key": true,
3378
+ "description": "German AI platform",
3379
+ "last_updated": "2026-07-31T05:17:14+00:00",
3380
+ "fallback_models": [
3381
+ "default"
3382
+ ],
3383
+ "default_model": "default"
3384
+ },
3385
+ {
3386
+ "id": "sglang_hosted",
3387
+ "name": "SGLang Server (Local / Remote)",
3388
+ "country": "International",
3389
+ "api_endpoint": "http://localhost:30000/v1",
3390
+ "api_style": "openai",
3391
+ "supports_model_listing": true,
3392
+ "supports_streaming": true,
3393
+ "supports_vision": true,
3394
+ "supports_embeddings": false,
3395
+ "supports_audio": false,
3396
+ "supports_reasoning": true,
3397
+ "openai_compatible": true,
3398
+ "free_models_available": true,
3399
+ "paid_models_available": false,
3400
+ "needs_key": false,
3401
+ "description": "Fast serving framework for complex reasoning, multi-turn chat, and structured generation",
3402
+ "last_updated": "2026-09-01T00:00:00+00:00",
3403
+ "fallback_models": [
3404
+ "default",
3405
+ "deepseek-ai/DeepSeek-R1-Distill-Qwen-32B",
3406
+ "meta-llama/Llama-3.3-70B-Instruct"
3407
+ ],
3408
+ "default_model": "default"
3409
+ },
3410
+ {
3411
+ "id": "siliconflow",
3412
+ "name": "SiliconFlow",
3413
+ "country": "China",
3414
+ "api_endpoint": "https://api.siliconflow.cn/v1",
3415
+ "api_style": "openai",
3416
+ "supports_model_listing": true,
3417
+ "supports_streaming": true,
3418
+ "supports_vision": false,
3419
+ "supports_embeddings": false,
3420
+ "supports_audio": false,
3421
+ "supports_reasoning": false,
3422
+ "openai_compatible": true,
3423
+ "free_models_available": true,
3424
+ "paid_models_available": true,
3425
+ "needs_key": true,
3426
+ "description": "Chinese open-model inference platform",
3427
+ "last_updated": "2026-07-31T05:17:14+00:00",
3428
+ "fallback_models": [
3429
+ "deepseek-v3",
3430
+ "glm-4",
3431
+ "qwen2.5-72b"
3432
+ ],
3433
+ "default_model": "deepseek-v3"
3434
+ },
3435
+ {
3436
+ "id": "snowflake",
3437
+ "name": "Snowflake Cortex",
3438
+ "country": "United States",
3439
+ "api_endpoint": "https://<account>.snowflakecomputing.com/api/v2/cortex/chat/completions",
3440
+ "api_style": "openai",
3441
+ "supports_model_listing": true,
3442
+ "supports_streaming": true,
3443
+ "supports_vision": false,
3444
+ "supports_embeddings": false,
3445
+ "supports_audio": false,
3446
+ "supports_reasoning": false,
3447
+ "openai_compatible": true,
3448
+ "free_models_available": false,
3449
+ "paid_models_available": true,
3450
+ "needs_key": true,
3451
+ "description": "Snowflake Cortex — Enterprise data AI",
3452
+ "last_updated": "2026-09-03T00:00:00+00:00",
3453
+ "fallback_models": [
3454
+ "snowflake-arctic",
3455
+ "llama-3.1-70b",
3456
+ "llama-3.1-8b"
3457
+ ],
3458
+ "default_model": "snowflake-arctic"
3459
+ },
3460
+ {
3461
+ "id": "sourcegraph-cody",
3462
+ "name": "Sourcegraph Cody",
3463
+ "country": "United States",
3464
+ "api_endpoint": "https://sourcegraph.com/.api/completions/stream",
3465
+ "api_style": "openai",
3466
+ "supports_model_listing": true,
3467
+ "supports_streaming": true,
3468
+ "supports_vision": false,
3469
+ "supports_embeddings": true,
3470
+ "supports_audio": false,
3471
+ "supports_reasoning": true,
3472
+ "openai_compatible": true,
3473
+ "free_models_available": true,
3474
+ "paid_models_available": true,
3475
+ "needs_key": true,
3476
+ "description": "Code intelligence and multi-repo context assistant",
3477
+ "last_updated": "2026-09-01T00:00:00+00:00",
3478
+ "fallback_models": [
3479
+ "cody-chat",
3480
+ "claude-3.5-sonnet",
3481
+ "gemini-1.5-pro"
3482
+ ],
3483
+ "default_model": "cody-chat"
3484
+ },
3485
+ {
3486
+ "id": "spark",
3487
+ "name": "Spark AI (iFlytek)",
3488
+ "country": "China",
3489
+ "api_endpoint": "https://spark-api.xf-yun.com/v1",
3490
+ "api_style": "openai",
3491
+ "supports_model_listing": true,
3492
+ "supports_streaming": true,
3493
+ "supports_vision": false,
3494
+ "supports_embeddings": false,
3495
+ "supports_audio": true,
3496
+ "supports_reasoning": false,
3497
+ "openai_compatible": true,
3498
+ "free_models_available": true,
3499
+ "paid_models_available": true,
3500
+ "needs_key": true,
3501
+ "description": "iFlytek Spark — Chinese LLM",
3502
+ "last_updated": "2026-07-31T05:17:14+00:00",
3503
+ "fallback_models": [
3504
+ "spark-3.5",
3505
+ "spark-4.0"
3506
+ ],
3507
+ "default_model": "spark-3.5"
3508
+ },
3509
+ {
3510
+ "id": "stability",
3511
+ "name": "Stability AI",
3512
+ "country": "United Kingdom",
3513
+ "api_endpoint": "https://api.stability.ai/v1",
3514
+ "api_style": "openai",
3515
+ "supports_model_listing": true,
3516
+ "supports_streaming": false,
3517
+ "supports_vision": true,
3518
+ "supports_embeddings": false,
3519
+ "supports_audio": false,
3520
+ "supports_reasoning": false,
3521
+ "openai_compatible": false,
3522
+ "free_models_available": true,
3523
+ "paid_models_available": true,
3524
+ "needs_key": true,
3525
+ "description": "Stable Diffusion — image generation",
3526
+ "last_updated": "2026-07-31T05:17:14+00:00",
3527
+ "fallback_models": [
3528
+ "stable-diffusion-3.5",
3529
+ "stable-image-ultra"
3530
+ ],
3531
+ "default_model": "stable-diffusion-3.5"
3532
+ },
3533
+ {
3534
+ "id": "crfm_helm",
3535
+ "name": "Stanford CRFM (HELM)",
3536
+ "country": "United States",
3537
+ "api_endpoint": "https://crfm.stanford.edu/helm/api/v1",
3538
+ "api_style": "openai",
3539
+ "supports_model_listing": true,
3540
+ "supports_streaming": false,
3541
+ "supports_vision": true,
3542
+ "supports_embeddings": true,
3543
+ "supports_audio": false,
3544
+ "supports_reasoning": true,
3545
+ "openai_compatible": true,
3546
+ "free_models_available": true,
3547
+ "paid_models_available": false,
3548
+ "needs_key": true,
3549
+ "description": "Academic evaluation benchmark and model exploration API from Stanford University",
3550
+ "last_updated": "2026-09-01T00:00:00+00:00",
3551
+ "fallback_models": [
3552
+ "helm-eval-suite",
3553
+ "stanford-crfm/llama-2-7b"
3554
+ ],
3555
+ "default_model": "helm-eval-suite"
3556
+ },
3557
+ {
3558
+ "id": "stepfun",
3559
+ "name": "StepFun (Step)",
3560
+ "country": "China",
3561
+ "api_endpoint": "https://api.stepfun.com/v1",
3562
+ "api_style": "openai",
3563
+ "supports_model_listing": true,
3564
+ "supports_streaming": true,
3565
+ "supports_vision": true,
3566
+ "supports_embeddings": false,
3567
+ "supports_audio": false,
3568
+ "supports_reasoning": false,
3569
+ "openai_compatible": true,
3570
+ "free_models_available": true,
3571
+ "paid_models_available": true,
3572
+ "needs_key": true,
3573
+ "description": "Step 2 — StepFun's long-context LLM",
3574
+ "last_updated": "2026-09-03T00:00:00+00:00",
3575
+ "fallback_models": [
3576
+ "step-2-16k",
3577
+ "step-2-32k",
3578
+ "step-1-flash"
3579
+ ],
3580
+ "default_model": "step-2-16k"
3581
+ },
3582
+ {
3583
+ "id": "stochastic",
3584
+ "name": "Stochastic AI",
3585
+ "country": "United States",
3586
+ "api_endpoint": "https://api.stochastic.ai/v1",
3587
+ "api_style": "openai",
3588
+ "supports_model_listing": true,
3589
+ "supports_streaming": true,
3590
+ "supports_vision": false,
3591
+ "supports_embeddings": false,
3592
+ "supports_audio": false,
3593
+ "supports_reasoning": false,
3594
+ "openai_compatible": true,
3595
+ "free_models_available": true,
3596
+ "paid_models_available": true,
3597
+ "needs_key": true,
3598
+ "description": "Model inference API",
3599
+ "last_updated": "2026-07-31T05:17:14+00:00",
3600
+ "fallback_models": [
3601
+ "llama3-70b"
3602
+ ],
3603
+ "default_model": "llama3-70b"
3604
+ },
3605
+ {
3606
+ "id": "suno",
3607
+ "name": "Suno AI Music",
3608
+ "country": "United States",
3609
+ "api_endpoint": "https://api.suno.ai/v1",
3610
+ "api_style": "openai",
3611
+ "supports_model_listing": true,
3612
+ "supports_streaming": true,
3613
+ "supports_vision": false,
3614
+ "supports_embeddings": false,
3615
+ "supports_audio": true,
3616
+ "supports_reasoning": false,
3617
+ "openai_compatible": false,
3618
+ "free_models_available": false,
3619
+ "paid_models_available": true,
3620
+ "needs_key": true,
3621
+ "description": "Generative musical composition, arrangement, and vocal synthesis",
3622
+ "last_updated": "2026-09-01T00:00:00+00:00",
3623
+ "fallback_models": [
3624
+ "chirp-v3-5",
3625
+ "chirp-v3"
3626
+ ],
3627
+ "default_model": "chirp-v3-5"
3628
+ },
3629
+ {
3630
+ "id": "tabby_api",
3631
+ "name": "TabbyAPI (ExLlamaV2)",
3632
+ "country": "International",
3633
+ "api_endpoint": "http://localhost:5000/v1",
3634
+ "api_style": "openai",
3635
+ "supports_model_listing": true,
3636
+ "supports_streaming": true,
3637
+ "supports_vision": false,
3638
+ "supports_embeddings": false,
3639
+ "supports_audio": false,
3640
+ "supports_reasoning": true,
3641
+ "openai_compatible": true,
3642
+ "free_models_available": true,
3643
+ "paid_models_available": false,
3644
+ "needs_key": true,
3645
+ "description": "Fast local inference server hosting EXL2 quantized models with OpenAI-compatible API",
3646
+ "last_updated": "2026-09-01T00:00:00+00:00",
3647
+ "fallback_models": [
3648
+ "current",
3649
+ "llama-3.3-70b-instruct-exl2",
3650
+ "mistral-large-2407-exl2"
3651
+ ],
3652
+ "default_model": "current"
3653
+ },
3654
+ {
3655
+ "id": "tabnine",
3656
+ "name": "Tabnine AI",
3657
+ "country": "Israel",
3658
+ "api_endpoint": "https://api.tabnine.com/v1",
3659
+ "api_style": "openai",
3660
+ "supports_model_listing": true,
3661
+ "supports_streaming": true,
3662
+ "supports_vision": false,
3663
+ "supports_embeddings": false,
3664
+ "supports_audio": false,
3665
+ "supports_reasoning": false,
3666
+ "openai_compatible": true,
3667
+ "free_models_available": true,
3668
+ "paid_models_available": true,
3669
+ "needs_key": true,
3670
+ "description": "Private, secure code completion and generative chat",
3671
+ "last_updated": "2026-09-01T00:00:00+00:00",
3672
+ "fallback_models": [
3673
+ "tabnine-chat",
3674
+ "tabnine-protected"
3675
+ ],
3676
+ "default_model": "tabnine-chat"
3677
+ },
3678
+ {
3679
+ "id": "tencent",
3680
+ "name": "Tencent (Hunyuan)",
3681
+ "country": "China",
3682
+ "api_endpoint": "https://api.hunyuan.cloud.tencent.com/v1",
3683
+ "api_style": "openai",
3684
+ "supports_model_listing": true,
3685
+ "supports_streaming": true,
3686
+ "supports_vision": false,
3687
+ "supports_embeddings": false,
3688
+ "supports_audio": false,
3689
+ "supports_reasoning": false,
3690
+ "openai_compatible": true,
3691
+ "free_models_available": true,
3692
+ "paid_models_available": true,
3693
+ "needs_key": true,
3694
+ "description": "Hunyuan — Tencent's LLM",
3695
+ "last_updated": "2026-07-31T05:17:14+00:00",
3696
+ "fallback_models": [
3697
+ "hunyuan-pro",
3698
+ "hunyuan-standard"
3699
+ ],
3700
+ "default_model": "hunyuan-pro"
3701
+ },
3702
+ {
3703
+ "id": "tencent-hunyuan",
3704
+ "name": "Tencent Hunyuan",
3705
+ "country": "China",
3706
+ "api_endpoint": "https://hunyuan.tencentcloudapi.com",
3707
+ "api_style": "openai",
3708
+ "supports_model_listing": true,
3709
+ "supports_streaming": true,
3710
+ "supports_vision": true,
3711
+ "supports_embeddings": false,
3712
+ "supports_audio": false,
3713
+ "supports_reasoning": false,
3714
+ "openai_compatible": true,
3715
+ "free_models_available": true,
3716
+ "paid_models_available": true,
3717
+ "needs_key": true,
3718
+ "description": "Hunyuan — Tencent's flagship LLM",
3719
+ "last_updated": "2026-09-03T00:00:00+00:00",
3720
+ "fallback_models": [
3721
+ "hunyuan-pro",
3722
+ "hunyuan-standard",
3723
+ "hunyuan-lite"
3724
+ ],
3725
+ "default_model": "hunyuan-pro"
3726
+ },
3727
+ {
3728
+ "id": "text-generation-webui",
3729
+ "name": "text-generation-webui",
3730
+ "country": "International",
3731
+ "api_endpoint": "http://localhost:5000/v1",
3732
+ "api_style": "openai",
3733
+ "supports_model_listing": true,
3734
+ "supports_streaming": true,
3735
+ "supports_vision": false,
3736
+ "supports_embeddings": false,
3737
+ "supports_audio": false,
3738
+ "supports_reasoning": false,
3739
+ "openai_compatible": true,
3740
+ "free_models_available": true,
3741
+ "paid_models_available": false,
3742
+ "needs_key": false,
3743
+ "description": "text-generation-webui — Gradio UI for LLMs",
3744
+ "last_updated": "2026-09-03T00:00:00+00:00",
3745
+ "fallback_models": [
3746
+ "default"
3747
+ ],
3748
+ "default_model": "default"
3749
+ },
3750
+ {
3751
+ "id": "titan",
3752
+ "name": "Titan AI",
3753
+ "country": "United States",
3754
+ "api_endpoint": "https://api.titanai.io/v1",
3755
+ "api_style": "openai",
3756
+ "supports_model_listing": true,
3757
+ "supports_streaming": true,
3758
+ "supports_vision": false,
3759
+ "supports_embeddings": false,
3760
+ "supports_audio": false,
3761
+ "supports_reasoning": false,
3762
+ "openai_compatible": true,
3763
+ "free_models_available": true,
3764
+ "paid_models_available": true,
3765
+ "needs_key": true,
3766
+ "description": "AI platform",
3767
+ "last_updated": "2026-07-31T05:17:14+00:00",
3768
+ "fallback_models": [
3769
+ "default"
3770
+ ],
3771
+ "default_model": "default"
3772
+ },
3773
+ {
3774
+ "id": "together",
3775
+ "name": "Together AI",
3776
+ "country": "United States",
3777
+ "api_endpoint": "https://api.together.xyz/v1",
3778
+ "api_style": "openai",
3779
+ "supports_model_listing": true,
3780
+ "supports_streaming": true,
3781
+ "supports_vision": true,
3782
+ "supports_embeddings": false,
3783
+ "supports_audio": false,
3784
+ "supports_reasoning": false,
3785
+ "openai_compatible": true,
3786
+ "free_models_available": true,
3787
+ "paid_models_available": true,
3788
+ "needs_key": true,
3789
+ "description": "Hosted open-source model inference",
3790
+ "last_updated": "2026-07-31T05:17:14+00:00",
3791
+ "fallback_models": [
3792
+ "deepseek-ai/DeepSeek-V3",
3793
+ "google/gemma-2-27b-it",
3794
+ "meta-llama/Llama-3.1-8B-Instruct-Turbo",
3795
+ "meta-llama/Llama-3.3-70B-Instruct-Turbo",
3796
+ "mistralai/Mixtral-8x22B-Instruct-v0.1",
3797
+ "meta-llama/Llama-3.1-70B-Instruct-Turbo",
3798
+ "deepseek-ai/DeepSeek-R1"
3799
+ ],
3800
+ "default_model": "meta-llama/Llama-3.3-70B-Instruct-Turbo"
3801
+ },
3802
+ {
3803
+ "id": "trulens",
3804
+ "name": "TruLens AI",
3805
+ "country": "United States",
3806
+ "api_endpoint": "https://api.trulens.ai/v1",
3807
+ "api_style": "openai",
3808
+ "supports_model_listing": true,
3809
+ "supports_streaming": true,
3810
+ "supports_vision": false,
3811
+ "supports_embeddings": false,
3812
+ "supports_audio": false,
3813
+ "supports_reasoning": false,
3814
+ "openai_compatible": true,
3815
+ "free_models_available": true,
3816
+ "paid_models_available": true,
3817
+ "needs_key": true,
3818
+ "description": "AI evaluation & monitoring",
3819
+ "last_updated": "2026-07-31T05:17:14+00:00",
3820
+ "fallback_models": [
3821
+ "default"
3822
+ ],
3823
+ "default_model": "default"
3824
+ },
3825
+ {
3826
+ "id": "upstage",
3827
+ "name": "Upstage (Solar)",
3828
+ "country": "South Korea",
3829
+ "api_endpoint": "https://api.upstage.ai/v1/solar",
3830
+ "api_style": "openai",
3831
+ "supports_model_listing": true,
3832
+ "supports_streaming": true,
3833
+ "supports_vision": false,
3834
+ "supports_embeddings": true,
3835
+ "supports_audio": false,
3836
+ "supports_reasoning": false,
3837
+ "openai_compatible": true,
3838
+ "free_models_available": true,
3839
+ "paid_models_available": true,
3840
+ "needs_key": true,
3841
+ "description": "Solar 10.7B — Upstage's efficient LLM",
3842
+ "last_updated": "2026-09-03T00:00:00+00:00",
3843
+ "fallback_models": [
3844
+ "solar-10.7b-instruct",
3845
+ "solar-1-mini",
3846
+ "embedding-passage"
3847
+ ],
3848
+ "default_model": "solar-10.7b-instruct"
3849
+ },
3850
+ {
3851
+ "id": "vllm",
3852
+ "name": "vLLM",
3853
+ "country": "International",
3854
+ "api_endpoint": "http://localhost:8000/v1",
3855
+ "api_style": "openai",
3856
+ "supports_model_listing": true,
3857
+ "supports_streaming": true,
3858
+ "supports_vision": false,
3859
+ "supports_embeddings": false,
3860
+ "supports_audio": false,
3861
+ "supports_reasoning": false,
3862
+ "openai_compatible": true,
3863
+ "free_models_available": true,
3864
+ "paid_models_available": false,
3865
+ "needs_key": false,
3866
+ "description": "Self-hosted high-throughput inference",
3867
+ "last_updated": "2026-07-31T05:17:14+00:00",
3868
+ "fallback_models": [
3869
+ "llama3",
3870
+ "qwen2",
3871
+ "default",
3872
+ "meta-llama/Llama-3.1-8B-Instruct",
3873
+ "deepseek-ai/DeepSeek-R1-Distill-Qwen-14B"
3874
+ ],
3875
+ "default_model": "llama3"
3876
+ },
3877
+ {
3878
+ "id": "voiceflow",
3879
+ "name": "Voiceflow",
3880
+ "country": "United States",
3881
+ "api_endpoint": "https://api.voiceflow.com/v1",
3882
+ "api_style": "openai",
3883
+ "supports_model_listing": true,
3884
+ "supports_streaming": true,
3885
+ "supports_vision": false,
3886
+ "supports_embeddings": false,
3887
+ "supports_audio": false,
3888
+ "supports_reasoning": false,
3889
+ "openai_compatible": true,
3890
+ "free_models_available": true,
3891
+ "paid_models_available": true,
3892
+ "needs_key": true,
3893
+ "description": "Conversational AI platform",
3894
+ "last_updated": "2026-07-31T05:17:14+00:00",
3895
+ "fallback_models": [
3896
+ "default"
3897
+ ],
3898
+ "default_model": "default"
3899
+ },
3900
+ {
3901
+ "id": "voyage",
3902
+ "name": "Voyage AI",
3903
+ "country": "United States",
3904
+ "api_endpoint": "https://api.voyageai.com/v1",
3905
+ "api_style": "openai",
3906
+ "supports_model_listing": true,
3907
+ "supports_streaming": false,
3908
+ "supports_vision": false,
3909
+ "supports_embeddings": true,
3910
+ "supports_audio": false,
3911
+ "supports_reasoning": false,
3912
+ "openai_compatible": true,
3913
+ "free_models_available": true,
3914
+ "paid_models_available": true,
3915
+ "needs_key": true,
3916
+ "description": "State-of-the-art embedding models",
3917
+ "last_updated": "2026-07-31T05:17:14+00:00",
3918
+ "fallback_models": [
3919
+ "voyage-3",
3920
+ "voyage-code-3"
3921
+ ],
3922
+ "default_model": "voyage-3"
3923
+ },
3924
+ {
3925
+ "id": "weaviate",
3926
+ "name": "Weaviate Cloud AI",
3927
+ "country": "Netherlands",
3928
+ "api_endpoint": "https://api.weaviate.io/v1",
3929
+ "api_style": "openai",
3930
+ "supports_model_listing": true,
3931
+ "supports_streaming": false,
3932
+ "supports_vision": true,
3933
+ "supports_embeddings": true,
3934
+ "supports_audio": false,
3935
+ "supports_reasoning": false,
3936
+ "openai_compatible": false,
3937
+ "free_models_available": true,
3938
+ "paid_models_available": true,
3939
+ "needs_key": true,
3940
+ "description": "AI-native vector search engine with generative search modules",
3941
+ "last_updated": "2026-09-01T00:00:00+00:00",
3942
+ "fallback_models": [
3943
+ "text2vec-openai",
3944
+ "generative-cohere"
3945
+ ],
3946
+ "default_model": "text2vec-openai"
3947
+ },
3948
+ {
3949
+ "id": "writer",
3950
+ "name": "Writer AI",
3951
+ "country": "United States",
3952
+ "api_endpoint": "https://api.writer.com/v1",
3953
+ "api_style": "openai",
3954
+ "supports_model_listing": true,
3955
+ "supports_streaming": true,
3956
+ "supports_vision": false,
3957
+ "supports_embeddings": false,
3958
+ "supports_audio": false,
3959
+ "supports_reasoning": false,
3960
+ "openai_compatible": true,
3961
+ "free_models_available": true,
3962
+ "paid_models_available": true,
3963
+ "needs_key": true,
3964
+ "description": "Palmyra — enterprise LLM",
3965
+ "last_updated": "2026-07-31T05:17:14+00:00",
3966
+ "fallback_models": [
3967
+ "palmyra-large",
3968
+ "palmyra-x-004"
3969
+ ],
3970
+ "default_model": "palmyra-large"
3971
+ },
3972
+ {
3973
+ "id": "xai",
3974
+ "name": "xAI (Grok)",
3975
+ "country": "United States",
3976
+ "api_endpoint": "https://api.x.ai/v1",
3977
+ "api_style": "openai",
3978
+ "supports_model_listing": true,
3979
+ "supports_streaming": true,
3980
+ "supports_vision": true,
3981
+ "supports_embeddings": false,
3982
+ "supports_audio": false,
3983
+ "supports_reasoning": true,
3984
+ "openai_compatible": true,
3985
+ "free_models_available": true,
3986
+ "paid_models_available": true,
3987
+ "needs_key": true,
3988
+ "description": "Grok-3 & GLM — xAI's flagship models",
3989
+ "last_updated": "2026-09-03T00:00:00+00:00",
3990
+ "fallback_models": [
3991
+ "grok-2-1212",
3992
+ "grok-2-latest",
3993
+ "grok-2-vision-1212",
3994
+ "grok-3",
3995
+ "grok-3-mini",
3996
+ "grok-beta",
3997
+ "glm-5.2",
3998
+ "glm-5.3"
3999
+ ],
4000
+ "default_model": "grok-3"
4001
+ },
4002
+ {
4003
+ "id": "xenova",
4004
+ "name": "Xenova (Transformers.js)",
4005
+ "country": "Canada",
4006
+ "api_endpoint": "https://api.xenova.ai/v1",
4007
+ "api_style": "openai",
4008
+ "supports_model_listing": true,
4009
+ "supports_streaming": true,
4010
+ "supports_vision": false,
4011
+ "supports_embeddings": false,
4012
+ "supports_audio": false,
4013
+ "supports_reasoning": false,
4014
+ "openai_compatible": true,
4015
+ "free_models_available": true,
4016
+ "paid_models_available": false,
4017
+ "needs_key": false,
4018
+ "description": "Browser/Edge inference via Transformers.js",
4019
+ "last_updated": "2026-07-31T05:17:14+00:00",
4020
+ "fallback_models": [
4021
+ "llama-3.2-1b",
4022
+ "phi-3-mini"
4023
+ ],
4024
+ "default_model": "llama-3.2-1b"
4025
+ },
4026
+ {
4027
+ "id": "zhipu",
4028
+ "name": "Zhipu AI (GLM)",
4029
+ "country": "China",
4030
+ "api_endpoint": "https://open.bigmodel.cn/api/paas/v4",
4031
+ "api_style": "openai",
4032
+ "supports_model_listing": true,
4033
+ "supports_streaming": true,
4034
+ "supports_vision": true,
4035
+ "supports_embeddings": false,
4036
+ "supports_audio": false,
4037
+ "supports_reasoning": false,
4038
+ "openai_compatible": true,
4039
+ "free_models_available": true,
4040
+ "paid_models_available": true,
4041
+ "needs_key": true,
4042
+ "description": "GLM-4 — Zhipu's multilingual LLM",
4043
+ "last_updated": "2026-07-31T05:17:14+00:00",
4044
+ "fallback_models": [
4045
+ "glm-4-air",
4046
+ "glm-4-plus",
4047
+ "glm-4v"
4048
+ ],
4049
+ "default_model": "glm-4-air"
4050
+ },
4051
+ {
4052
+ "id": "zilliz",
4053
+ "name": "Zilliz Cloud",
4054
+ "country": "United States",
4055
+ "api_endpoint": "https://api.zilliz.com/v1",
4056
+ "api_style": "openai",
4057
+ "supports_model_listing": true,
4058
+ "supports_streaming": false,
4059
+ "supports_vision": false,
4060
+ "supports_embeddings": true,
4061
+ "supports_audio": false,
4062
+ "supports_reasoning": false,
4063
+ "openai_compatible": true,
4064
+ "free_models_available": true,
4065
+ "paid_models_available": true,
4066
+ "needs_key": true,
4067
+ "description": "Vector database & embedding API",
4068
+ "last_updated": "2026-09-03T00:00:00+00:00",
4069
+ "fallback_models": [
4070
+ "default"
4071
+ ],
4072
+ "default_model": "default"
4073
+ }
4074
+ ]
4075
+ }