@mastra/mcp-docs-server 1.2.18-alpha.1 → 1.2.18-alpha.3
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/.docs/integrations/agentic-ui/ai-sdk-ui.md +2 -1
- package/.docs/integrations/frameworks/next-js.md +3 -2
- package/.docs/models/gateways/netlify.md +227 -70
- package/.docs/models/gateways/openrouter.md +3 -4
- package/.docs/models/gateways/vercel.md +5 -3
- package/.docs/models/index.md +1 -1
- package/.docs/models/providers/crossmodel.md +3 -2
- package/.docs/models/providers/edenai.md +5 -7
- package/.docs/models/providers/empiriolabs.md +2 -1
- package/.docs/models/providers/huggingface.md +2 -1
- package/.docs/models/providers/hyper.md +9 -9
- package/.docs/models/providers/inceptron.md +4 -5
- package/.docs/models/providers/kilo.md +11 -12
- package/.docs/models/providers/llmgateway.md +2 -2
- package/.docs/models/providers/ofox.md +4 -4
- package/.docs/models/providers/opencode-go.md +1 -1
- package/.docs/models/providers/opencode.md +1 -1
- package/.docs/models/providers/pioneer.md +12 -12
- package/.docs/models/providers/runinfra.md +2 -1
- package/.docs/models/providers/umans-ai-coding-plan.md +1 -2
- package/.docs/models/providers/umans-ai.md +1 -2
- package/.docs/reference/agent-controller/agent-controller-class.md +1 -1
- package/.docs/reference/ai-sdk/handle-chat-stream.md +2 -1
- package/CHANGELOG.md +7 -0
- package/package.json +5 -5
|
@@ -135,7 +135,7 @@ If you don't want to run Mastra's server and instead use frameworks like Next.js
|
|
|
135
135
|
|
|
136
136
|
They return a `ReadableStream` that you can wrap with [`createUIMessageStreamResponse()`](https://ai-sdk.dev/docs/reference/ai-sdk-ui/create-ui-message-stream-response).
|
|
137
137
|
|
|
138
|
-
> **AI SDK
|
|
138
|
+
> **AI SDK version compatibility:** The framework-agnostic handlers default to AI SDK v5 for backward compatibility. Pass the version that matches your installed AI SDK, such as `version: 'v7'` for AI SDK v7. For best TypeScript inference with `handleChatStream()` and `handleNetworkStream()`, pass `messages` as `UIMessage[]` from your installed `ai` version.
|
|
139
139
|
|
|
140
140
|
The examples below show you how to use them with Next.js App Router.
|
|
141
141
|
|
|
@@ -153,6 +153,7 @@ export async function POST(req: Request) {
|
|
|
153
153
|
const stream = await handleChatStream({
|
|
154
154
|
mastra,
|
|
155
155
|
agentId: 'weatherAgent',
|
|
156
|
+
version: 'v7',
|
|
156
157
|
params,
|
|
157
158
|
})
|
|
158
159
|
return createUIMessageStreamResponse({ stream })
|
|
@@ -159,7 +159,7 @@ Create `src/app/api/chat/route.ts`:
|
|
|
159
159
|
|
|
160
160
|
```ts
|
|
161
161
|
import { handleChatStream } from '@mastra/ai-sdk'
|
|
162
|
-
import {
|
|
162
|
+
import { toAISdkMessages } from '@mastra/ai-sdk/ui'
|
|
163
163
|
import { createUIMessageStreamResponse } from 'ai'
|
|
164
164
|
import { mastra } from '@/mastra'
|
|
165
165
|
import { NextResponse } from 'next/server'
|
|
@@ -172,6 +172,7 @@ export async function POST(req: Request) {
|
|
|
172
172
|
const stream = await handleChatStream({
|
|
173
173
|
mastra,
|
|
174
174
|
agentId: 'weather-agent',
|
|
175
|
+
version: 'v7',
|
|
175
176
|
params: {
|
|
176
177
|
...params,
|
|
177
178
|
memory: {
|
|
@@ -197,7 +198,7 @@ export async function GET() {
|
|
|
197
198
|
console.log('No previous messages found.')
|
|
198
199
|
}
|
|
199
200
|
|
|
200
|
-
const uiMessages =
|
|
201
|
+
const uiMessages = toAISdkMessages(response?.messages || [], { version: 'v7' })
|
|
201
202
|
|
|
202
203
|
return NextResponse.json(uiMessages)
|
|
203
204
|
}
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Netlify
|
|
4
4
|
|
|
5
|
-
Netlify AI Gateway provides unified access to multiple providers with built-in caching and observability. Access
|
|
5
|
+
Netlify AI Gateway provides unified access to multiple providers with built-in caching and observability. Access 224 models through Mastra's model router.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Netlify documentation](https://docs.netlify.com/build/ai-gateway/overview/).
|
|
8
8
|
|
|
@@ -35,72 +35,229 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
35
35
|
|
|
36
36
|
## Available models
|
|
37
37
|
|
|
38
|
-
| Model
|
|
39
|
-
|
|
|
40
|
-
| `anthropic/claude-fable-5`
|
|
41
|
-
| `anthropic/claude-haiku-4-5`
|
|
42
|
-
| `anthropic/claude-haiku-4-5-20251001`
|
|
43
|
-
| `anthropic/claude-opus-4-5`
|
|
44
|
-
| `anthropic/claude-opus-4-5-20251101`
|
|
45
|
-
| `anthropic/claude-opus-4-6`
|
|
46
|
-
| `anthropic/claude-opus-4-7`
|
|
47
|
-
| `anthropic/claude-opus-4-8`
|
|
48
|
-
| `anthropic/claude-opus-5`
|
|
49
|
-
| `anthropic/claude-sonnet-4-5`
|
|
50
|
-
| `anthropic/claude-sonnet-4-5-20250929`
|
|
51
|
-
| `anthropic/claude-sonnet-4-6`
|
|
52
|
-
| `anthropic/claude-sonnet-5`
|
|
53
|
-
| `gemini/gemini-2.5-flash`
|
|
54
|
-
| `gemini/gemini-2.5-flash-image`
|
|
55
|
-
| `gemini/gemini-2.5-flash-lite`
|
|
56
|
-
| `gemini/gemini-2.5-pro`
|
|
57
|
-
| `gemini/gemini-3-flash-preview`
|
|
58
|
-
| `gemini/gemini-3-pro-image`
|
|
59
|
-
| `gemini/gemini-3.1-flash-image`
|
|
60
|
-
| `gemini/gemini-3.1-flash-lite`
|
|
61
|
-
| `gemini/gemini-3.1-flash-lite-image`
|
|
62
|
-
| `gemini/gemini-3.1-pro-preview`
|
|
63
|
-
| `gemini/gemini-3.1-pro-preview-customtools`
|
|
64
|
-
| `gemini/gemini-3.5-flash`
|
|
65
|
-
| `gemini/gemini-3.5-flash-lite`
|
|
66
|
-
| `gemini/gemini-3.6-flash`
|
|
67
|
-
| `gemini/gemini-3.7-flash`
|
|
68
|
-
| `gemini/gemini-flash-latest`
|
|
69
|
-
| `gemini/gemini-flash-lite-latest`
|
|
70
|
-
| `openai/chat-latest`
|
|
71
|
-
| `openai/gpt-4.1`
|
|
72
|
-
| `openai/gpt-4.1-mini`
|
|
73
|
-
| `openai/gpt-4.1-nano`
|
|
74
|
-
| `openai/gpt-4o`
|
|
75
|
-
| `openai/gpt-4o-mini`
|
|
76
|
-
| `openai/gpt-5`
|
|
77
|
-
| `openai/gpt-5-2025-08-07`
|
|
78
|
-
| `openai/gpt-5-mini`
|
|
79
|
-
| `openai/gpt-5-mini-2025-08-07`
|
|
80
|
-
| `openai/gpt-5-nano`
|
|
81
|
-
| `openai/gpt-5-pro`
|
|
82
|
-
| `openai/gpt-5.1`
|
|
83
|
-
| `openai/gpt-5.1-2025-11-13`
|
|
84
|
-
| `openai/gpt-5.2`
|
|
85
|
-
| `openai/gpt-5.2-2025-12-11`
|
|
86
|
-
| `openai/gpt-5.2-pro`
|
|
87
|
-
| `openai/gpt-5.2-pro-2025-12-11`
|
|
88
|
-
| `openai/gpt-5.3-codex`
|
|
89
|
-
| `openai/gpt-5.4`
|
|
90
|
-
| `openai/gpt-5.4-2026-03-05`
|
|
91
|
-
| `openai/gpt-5.4-mini`
|
|
92
|
-
| `openai/gpt-5.4-mini-2026-03-17`
|
|
93
|
-
| `openai/gpt-5.4-nano`
|
|
94
|
-
| `openai/gpt-5.4-nano-2026-03-17`
|
|
95
|
-
| `openai/gpt-5.4-pro`
|
|
96
|
-
| `openai/gpt-5.4-pro-2026-03-05`
|
|
97
|
-
| `openai/gpt-5.5`
|
|
98
|
-
| `openai/gpt-5.5-2026-04-23`
|
|
99
|
-
| `openai/gpt-5.5-pro`
|
|
100
|
-
| `openai/gpt-5.5-pro-2026-04-23`
|
|
101
|
-
| `openai/gpt-5.6-luna`
|
|
102
|
-
| `openai/gpt-5.6-sol`
|
|
103
|
-
| `openai/gpt-5.6-terra`
|
|
104
|
-
| `openai/o3`
|
|
105
|
-
| `openai/o3-mini`
|
|
106
|
-
| `openai/o4-mini`
|
|
38
|
+
| Model |
|
|
39
|
+
| --------------------------------------------------------------------- |
|
|
40
|
+
| `anthropic/claude-fable-5` |
|
|
41
|
+
| `anthropic/claude-haiku-4-5` |
|
|
42
|
+
| `anthropic/claude-haiku-4-5-20251001` |
|
|
43
|
+
| `anthropic/claude-opus-4-5` |
|
|
44
|
+
| `anthropic/claude-opus-4-5-20251101` |
|
|
45
|
+
| `anthropic/claude-opus-4-6` |
|
|
46
|
+
| `anthropic/claude-opus-4-7` |
|
|
47
|
+
| `anthropic/claude-opus-4-8` |
|
|
48
|
+
| `anthropic/claude-opus-5` |
|
|
49
|
+
| `anthropic/claude-sonnet-4-5` |
|
|
50
|
+
| `anthropic/claude-sonnet-4-5-20250929` |
|
|
51
|
+
| `anthropic/claude-sonnet-4-6` |
|
|
52
|
+
| `anthropic/claude-sonnet-5` |
|
|
53
|
+
| `gemini/gemini-2.5-flash` |
|
|
54
|
+
| `gemini/gemini-2.5-flash-image` |
|
|
55
|
+
| `gemini/gemini-2.5-flash-lite` |
|
|
56
|
+
| `gemini/gemini-2.5-pro` |
|
|
57
|
+
| `gemini/gemini-3-flash-preview` |
|
|
58
|
+
| `gemini/gemini-3-pro-image` |
|
|
59
|
+
| `gemini/gemini-3.1-flash-image` |
|
|
60
|
+
| `gemini/gemini-3.1-flash-lite` |
|
|
61
|
+
| `gemini/gemini-3.1-flash-lite-image` |
|
|
62
|
+
| `gemini/gemini-3.1-pro-preview` |
|
|
63
|
+
| `gemini/gemini-3.1-pro-preview-customtools` |
|
|
64
|
+
| `gemini/gemini-3.5-flash` |
|
|
65
|
+
| `gemini/gemini-3.5-flash-lite` |
|
|
66
|
+
| `gemini/gemini-3.6-flash` |
|
|
67
|
+
| `gemini/gemini-3.7-flash` |
|
|
68
|
+
| `gemini/gemini-flash-latest` |
|
|
69
|
+
| `gemini/gemini-flash-lite-latest` |
|
|
70
|
+
| `openai/chat-latest` |
|
|
71
|
+
| `openai/gpt-4.1` |
|
|
72
|
+
| `openai/gpt-4.1-mini` |
|
|
73
|
+
| `openai/gpt-4.1-nano` |
|
|
74
|
+
| `openai/gpt-4o` |
|
|
75
|
+
| `openai/gpt-4o-mini` |
|
|
76
|
+
| `openai/gpt-5` |
|
|
77
|
+
| `openai/gpt-5-2025-08-07` |
|
|
78
|
+
| `openai/gpt-5-mini` |
|
|
79
|
+
| `openai/gpt-5-mini-2025-08-07` |
|
|
80
|
+
| `openai/gpt-5-nano` |
|
|
81
|
+
| `openai/gpt-5-pro` |
|
|
82
|
+
| `openai/gpt-5.1` |
|
|
83
|
+
| `openai/gpt-5.1-2025-11-13` |
|
|
84
|
+
| `openai/gpt-5.2` |
|
|
85
|
+
| `openai/gpt-5.2-2025-12-11` |
|
|
86
|
+
| `openai/gpt-5.2-pro` |
|
|
87
|
+
| `openai/gpt-5.2-pro-2025-12-11` |
|
|
88
|
+
| `openai/gpt-5.3-codex` |
|
|
89
|
+
| `openai/gpt-5.4` |
|
|
90
|
+
| `openai/gpt-5.4-2026-03-05` |
|
|
91
|
+
| `openai/gpt-5.4-mini` |
|
|
92
|
+
| `openai/gpt-5.4-mini-2026-03-17` |
|
|
93
|
+
| `openai/gpt-5.4-nano` |
|
|
94
|
+
| `openai/gpt-5.4-nano-2026-03-17` |
|
|
95
|
+
| `openai/gpt-5.4-pro` |
|
|
96
|
+
| `openai/gpt-5.4-pro-2026-03-05` |
|
|
97
|
+
| `openai/gpt-5.5` |
|
|
98
|
+
| `openai/gpt-5.5-2026-04-23` |
|
|
99
|
+
| `openai/gpt-5.5-pro` |
|
|
100
|
+
| `openai/gpt-5.5-pro-2026-04-23` |
|
|
101
|
+
| `openai/gpt-5.6-luna` |
|
|
102
|
+
| `openai/gpt-5.6-sol` |
|
|
103
|
+
| `openai/gpt-5.6-terra` |
|
|
104
|
+
| `openai/o3` |
|
|
105
|
+
| `openai/o3-mini` |
|
|
106
|
+
| `openai/o4-mini` |
|
|
107
|
+
| `openrouter/~deepseek/deepseek-v4-flash-latest` |
|
|
108
|
+
| `openrouter/~moonshotai/kimi-latest` |
|
|
109
|
+
| `openrouter/~x-ai/grok-latest` |
|
|
110
|
+
| `openrouter/allenai/olmo-3-32b-think` |
|
|
111
|
+
| `openrouter/anthracite-org/magnum-v4-72b` |
|
|
112
|
+
| `openrouter/arcee-ai/trinity-large-thinking` |
|
|
113
|
+
| `openrouter/arcee-ai/virtuoso-large` |
|
|
114
|
+
| `openrouter/baidu/ernie-4.5-vl-424b-a47b` |
|
|
115
|
+
| `openrouter/bytedance-seed/seed-1.6` |
|
|
116
|
+
| `openrouter/bytedance-seed/seed-1.6-flash` |
|
|
117
|
+
| `openrouter/bytedance-seed/seed-2-1-turbo` |
|
|
118
|
+
| `openrouter/bytedance-seed/seed-2.0-code` |
|
|
119
|
+
| `openrouter/bytedance-seed/seed-2.0-lite` |
|
|
120
|
+
| `openrouter/bytedance-seed/seed-2.0-mini` |
|
|
121
|
+
| `openrouter/bytedance/ui-tars-1.5-7b` |
|
|
122
|
+
| `openrouter/cognitivecomputations/dolphin-mistral-24b-venice-edition` |
|
|
123
|
+
| `openrouter/deepcogito/cogito-v2.1-671b` |
|
|
124
|
+
| `openrouter/deepseek/deepseek-chat` |
|
|
125
|
+
| `openrouter/deepseek/deepseek-chat-v3-0324` |
|
|
126
|
+
| `openrouter/deepseek/deepseek-chat-v3.1` |
|
|
127
|
+
| `openrouter/deepseek/deepseek-r1` |
|
|
128
|
+
| `openrouter/deepseek/deepseek-r1-0528` |
|
|
129
|
+
| `openrouter/deepseek/deepseek-r1-distill-llama-70b` |
|
|
130
|
+
| `openrouter/deepseek/deepseek-v3.1-terminus` |
|
|
131
|
+
| `openrouter/deepseek/deepseek-v3.2` |
|
|
132
|
+
| `openrouter/deepseek/deepseek-v3.2-exp` |
|
|
133
|
+
| `openrouter/deepseek/deepseek-v4-flash` |
|
|
134
|
+
| `openrouter/deepseek/deepseek-v4-flash-0731` |
|
|
135
|
+
| `openrouter/deepseek/deepseek-v4-pro` |
|
|
136
|
+
| `openrouter/deepseek/deepseek-v4-pro-0813` |
|
|
137
|
+
| `openrouter/google/gemma-2-27b-it` |
|
|
138
|
+
| `openrouter/google/gemma-3-12b-it` |
|
|
139
|
+
| `openrouter/google/gemma-3-27b-it` |
|
|
140
|
+
| `openrouter/google/gemma-3-4b-it` |
|
|
141
|
+
| `openrouter/google/gemma-3n-e4b-it` |
|
|
142
|
+
| `openrouter/google/gemma-4-26b-a4b-it` |
|
|
143
|
+
| `openrouter/google/gemma-4-31b-it` |
|
|
144
|
+
| `openrouter/gryphe/mythomax-l2-13b` |
|
|
145
|
+
| `openrouter/ibm-granite/granite-4.1-8b` |
|
|
146
|
+
| `openrouter/inception/mercury-2` |
|
|
147
|
+
| `openrouter/inclusionai/ling-2.6-1t` |
|
|
148
|
+
| `openrouter/inclusionai/ling-2.6-flash` |
|
|
149
|
+
| `openrouter/inclusionai/ling-3.0-flash` |
|
|
150
|
+
| `openrouter/inclusionai/ring-2.6-1t` |
|
|
151
|
+
| `openrouter/meta-llama/llama-3.1-70b-instruct` |
|
|
152
|
+
| `openrouter/meta-llama/llama-3.1-8b-instruct` |
|
|
153
|
+
| `openrouter/meta-llama/llama-3.2-3b-instruct` |
|
|
154
|
+
| `openrouter/meta-llama/llama-3.3-70b-instruct` |
|
|
155
|
+
| `openrouter/meta-llama/llama-4-maverick` |
|
|
156
|
+
| `openrouter/meta-llama/llama-4-scout` |
|
|
157
|
+
| `openrouter/meta-llama/llama-guard-4-12b` |
|
|
158
|
+
| `openrouter/meta/muse-glimmer-30b` |
|
|
159
|
+
| `openrouter/microsoft/phi-4` |
|
|
160
|
+
| `openrouter/microsoft/wizardlm-2-8x22b` |
|
|
161
|
+
| `openrouter/minimax/minimax-m1` |
|
|
162
|
+
| `openrouter/minimax/minimax-m2` |
|
|
163
|
+
| `openrouter/minimax/minimax-m2.1` |
|
|
164
|
+
| `openrouter/minimax/minimax-m2.5` |
|
|
165
|
+
| `openrouter/minimax/minimax-m2.7` |
|
|
166
|
+
| `openrouter/minimax/minimax-m3` |
|
|
167
|
+
| `openrouter/mistralai/mistral-nemo` |
|
|
168
|
+
| `openrouter/mistralai/mistral-small-24b-instruct-2501` |
|
|
169
|
+
| `openrouter/mistralai/mistral-small-2603` |
|
|
170
|
+
| `openrouter/mistralai/mistral-small-3.2-24b-instruct` |
|
|
171
|
+
| `openrouter/moonshotai/kimi-k2` |
|
|
172
|
+
| `openrouter/moonshotai/kimi-k2-0905` |
|
|
173
|
+
| `openrouter/moonshotai/kimi-k2-thinking` |
|
|
174
|
+
| `openrouter/moonshotai/kimi-k2.5` |
|
|
175
|
+
| `openrouter/moonshotai/kimi-k2.6` |
|
|
176
|
+
| `openrouter/moonshotai/kimi-k2.7-code` |
|
|
177
|
+
| `openrouter/moonshotai/kimi-k3` |
|
|
178
|
+
| `openrouter/morph/morph-v3-fast` |
|
|
179
|
+
| `openrouter/morph/morph-v3-large` |
|
|
180
|
+
| `openrouter/nousresearch/hermes-3-llama-3.1-405b` |
|
|
181
|
+
| `openrouter/nousresearch/hermes-3-llama-3.1-70b` |
|
|
182
|
+
| `openrouter/nousresearch/hermes-4-405b` |
|
|
183
|
+
| `openrouter/nousresearch/hermes-4-70b` |
|
|
184
|
+
| `openrouter/nvidia/nemotron-3-nano-30b-a3b` |
|
|
185
|
+
| `openrouter/nvidia/nemotron-3-super-120b-a12b` |
|
|
186
|
+
| `openrouter/nvidia/nemotron-3-ultra-550b-a55b` |
|
|
187
|
+
| `openrouter/nvidia/nemotron-3.5-lightning` |
|
|
188
|
+
| `openrouter/openai/gpt-oss-120b` |
|
|
189
|
+
| `openrouter/openai/gpt-oss-20b` |
|
|
190
|
+
| `openrouter/openai/gpt-oss-safeguard-20b` |
|
|
191
|
+
| `openrouter/openrouter/auto` |
|
|
192
|
+
| `openrouter/openrouter/auto-beta` |
|
|
193
|
+
| `openrouter/openrouter/bodybuilder` |
|
|
194
|
+
| `openrouter/openrouter/free` |
|
|
195
|
+
| `openrouter/openrouter/fusion` |
|
|
196
|
+
| `openrouter/openrouter/pareto-code` |
|
|
197
|
+
| `openrouter/perceptron/perceptron-mk1` |
|
|
198
|
+
| `openrouter/perplexity/sonar` |
|
|
199
|
+
| `openrouter/perplexity/sonar-deep-research` |
|
|
200
|
+
| `openrouter/perplexity/sonar-pro` |
|
|
201
|
+
| `openrouter/perplexity/sonar-pro-search` |
|
|
202
|
+
| `openrouter/perplexity/sonar-reasoning-pro` |
|
|
203
|
+
| `openrouter/qwen/qwen-2.5-72b-instruct` |
|
|
204
|
+
| `openrouter/qwen/qwen-2.5-7b-instruct` |
|
|
205
|
+
| `openrouter/qwen/qwen2.5-vl-72b-instruct` |
|
|
206
|
+
| `openrouter/qwen/qwen3-14b` |
|
|
207
|
+
| `openrouter/qwen/qwen3-235b-a22b-2507` |
|
|
208
|
+
| `openrouter/qwen/qwen3-235b-a22b-thinking-2507` |
|
|
209
|
+
| `openrouter/qwen/qwen3-30b-a3b` |
|
|
210
|
+
| `openrouter/qwen/qwen3-30b-a3b-instruct-2507` |
|
|
211
|
+
| `openrouter/qwen/qwen3-32b` |
|
|
212
|
+
| `openrouter/qwen/qwen3-coder` |
|
|
213
|
+
| `openrouter/qwen/qwen3-coder-30b-a3b-instruct` |
|
|
214
|
+
| `openrouter/qwen/qwen3-coder-next` |
|
|
215
|
+
| `openrouter/qwen/qwen3-next-80b-a3b-instruct` |
|
|
216
|
+
| `openrouter/qwen/qwen3-next-80b-a3b-thinking` |
|
|
217
|
+
| `openrouter/qwen/qwen3-vl-235b-a22b-instruct` |
|
|
218
|
+
| `openrouter/qwen/qwen3-vl-235b-a22b-thinking` |
|
|
219
|
+
| `openrouter/qwen/qwen3-vl-30b-a3b-instruct` |
|
|
220
|
+
| `openrouter/qwen/qwen3-vl-8b-instruct` |
|
|
221
|
+
| `openrouter/qwen/qwen3.5-122b-a10b` |
|
|
222
|
+
| `openrouter/qwen/qwen3.5-27b` |
|
|
223
|
+
| `openrouter/qwen/qwen3.5-35b-a3b` |
|
|
224
|
+
| `openrouter/qwen/qwen3.5-397b-a17b` |
|
|
225
|
+
| `openrouter/qwen/qwen3.5-9b` |
|
|
226
|
+
| `openrouter/qwen/qwen3.6-27b` |
|
|
227
|
+
| `openrouter/qwen/qwen3.6-35b-a3b` |
|
|
228
|
+
| `openrouter/qwen/qwen3.8-2.4t-a95b` |
|
|
229
|
+
| `openrouter/qwen/qwen3.8-27b` |
|
|
230
|
+
| `openrouter/rekaai/reka-edge` |
|
|
231
|
+
| `openrouter/rekaai/reka-flash-3` |
|
|
232
|
+
| `openrouter/relace/relace-apply-3` |
|
|
233
|
+
| `openrouter/relace/relace-search` |
|
|
234
|
+
| `openrouter/sao10k/l3-lunaris-8b` |
|
|
235
|
+
| `openrouter/sao10k/l3.1-euryale-70b` |
|
|
236
|
+
| `openrouter/sao10k/l3.3-euryale-70b` |
|
|
237
|
+
| `openrouter/stepfun/step-3.7-flash` |
|
|
238
|
+
| `openrouter/tencent/hy3` |
|
|
239
|
+
| `openrouter/thedrummer/cydonia-24b-v4.1` |
|
|
240
|
+
| `openrouter/thedrummer/rocinante-12b` |
|
|
241
|
+
| `openrouter/thedrummer/skyfall-36b-v2` |
|
|
242
|
+
| `openrouter/thedrummer/unslopnemo-12b` |
|
|
243
|
+
| `openrouter/thinkingmachines/inkling` |
|
|
244
|
+
| `openrouter/thinkingmachines/inkling-small` |
|
|
245
|
+
| `openrouter/undi95/remm-slerp-l2-13b` |
|
|
246
|
+
| `openrouter/x-ai/grok-4.20` |
|
|
247
|
+
| `openrouter/x-ai/grok-4.20-multi-agent` |
|
|
248
|
+
| `openrouter/x-ai/grok-4.3` |
|
|
249
|
+
| `openrouter/x-ai/grok-4.5` |
|
|
250
|
+
| `openrouter/x-ai/grok-4.6` |
|
|
251
|
+
| `openrouter/x-ai/grok-build-0.1` |
|
|
252
|
+
| `openrouter/xiaomi/mimo-v2.5` |
|
|
253
|
+
| `openrouter/xiaomi/mimo-v2.5-pro` |
|
|
254
|
+
| `openrouter/z-ai/glm-4.5-air` |
|
|
255
|
+
| `openrouter/z-ai/glm-4.5v` |
|
|
256
|
+
| `openrouter/z-ai/glm-4.6` |
|
|
257
|
+
| `openrouter/z-ai/glm-4.6v` |
|
|
258
|
+
| `openrouter/z-ai/glm-4.7` |
|
|
259
|
+
| `openrouter/z-ai/glm-4.7-flash` |
|
|
260
|
+
| `openrouter/z-ai/glm-5` |
|
|
261
|
+
| `openrouter/z-ai/glm-5.1` |
|
|
262
|
+
| `openrouter/z-ai/glm-5.2` |
|
|
263
|
+
| `openrouter/z-ai/glm-5.2:free` |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# OpenRouter
|
|
4
4
|
|
|
5
|
-
OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access
|
|
5
|
+
OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 352 models through Mastra's model router.
|
|
6
6
|
|
|
7
7
|
Learn more in the [OpenRouter documentation](https://openrouter.ai/models).
|
|
8
8
|
|
|
@@ -15,7 +15,7 @@ const agent = new Agent({
|
|
|
15
15
|
id: "my-agent",
|
|
16
16
|
name: "My Agent",
|
|
17
17
|
instructions: "You are a helpful assistant",
|
|
18
|
-
model: "openrouter/
|
|
18
|
+
model: "openrouter/aion-labs/aion-2.0"
|
|
19
19
|
});
|
|
20
20
|
```
|
|
21
21
|
|
|
@@ -47,7 +47,7 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
47
47
|
| `~openai/gpt-latest` |
|
|
48
48
|
| `~openai/gpt-mini-latest` |
|
|
49
49
|
| `~x-ai/grok-latest` |
|
|
50
|
-
|
|
|
50
|
+
| `~z-ai/glm-latest` |
|
|
51
51
|
| `aion-labs/aion-2.0` |
|
|
52
52
|
| `aion-labs/aion-3.0` |
|
|
53
53
|
| `aion-labs/aion-3.0-mini` |
|
|
@@ -150,7 +150,6 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
150
150
|
| `kwaipilot/kat-coder-pro-v2` |
|
|
151
151
|
| `kwaipilot/kat-coder-pro-v2.5` |
|
|
152
152
|
| `liquid/lfm-2.5-2.6b:free` |
|
|
153
|
-
| `mancer/weaver` |
|
|
154
153
|
| `meituan/longcat-2.0` |
|
|
155
154
|
| `meta-llama/llama-3.1-70b-instruct` |
|
|
156
155
|
| `meta-llama/llama-3.1-8b-instruct` |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Vercel
|
|
4
4
|
|
|
5
|
-
Vercel aggregates models from multiple providers with enhanced features like rate limiting and failover. Access
|
|
5
|
+
Vercel aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 350 models through Mastra's model router.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Vercel documentation](https://ai-sdk.dev/providers/ai-sdk-providers).
|
|
8
8
|
|
|
@@ -136,9 +136,13 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
136
136
|
| `deepseek/deepseek-v4-pro` |
|
|
137
137
|
| `deepseek/deepseek-v4-pro-0813` |
|
|
138
138
|
| `fish-audio/s1` |
|
|
139
|
+
| `fish-audio/s1-free` |
|
|
139
140
|
| `fish-audio/s2-pro` |
|
|
141
|
+
| `fish-audio/s2-pro-free` |
|
|
140
142
|
| `fish-audio/s2.1-pro` |
|
|
143
|
+
| `fish-audio/s2.1-pro-free` |
|
|
141
144
|
| `fish-audio/transcribe-1` |
|
|
145
|
+
| `fish-audio/transcribe-1-free` |
|
|
142
146
|
| `google/gemini-2.5-flash` |
|
|
143
147
|
| `google/gemini-2.5-flash-image` |
|
|
144
148
|
| `google/gemini-2.5-flash-lite` |
|
|
@@ -372,8 +376,6 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
372
376
|
| `zai/glm-4.5-air` |
|
|
373
377
|
| `zai/glm-4.5v` |
|
|
374
378
|
| `zai/glm-4.6` |
|
|
375
|
-
| `zai/glm-4.6v` |
|
|
376
|
-
| `zai/glm-4.6v-flash` |
|
|
377
379
|
| `zai/glm-4.7` |
|
|
378
380
|
| `zai/glm-4.7-flash` |
|
|
379
381
|
| `zai/glm-4.7-flashx` |
|
package/.docs/models/index.md
CHANGED
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Model Providers
|
|
4
4
|
|
|
5
|
-
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to
|
|
5
|
+
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to 6358 models from 179 providers through a single API.
|
|
6
6
|
|
|
7
7
|
## Features
|
|
8
8
|
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# CrossModel
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 52 CrossModel models through Mastra's model router. Authentication is handled automatically using the `CROSSMODEL_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [CrossModel documentation](https://www.crossmodel.ai/docs).
|
|
8
8
|
|
|
@@ -87,6 +87,7 @@ for await (const chunk of stream) {
|
|
|
87
87
|
| `crossmodel/z-ai/glm-5-turbo` | 200K | | | | | | $0.90 | $4 |
|
|
88
88
|
| `crossmodel/z-ai/glm-5.1` | 200K | | | | | | $1 | $4 |
|
|
89
89
|
| `crossmodel/z-ai/glm-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
90
|
+
| `crossmodel/z-ai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
90
91
|
|
|
91
92
|
## Advanced configuration
|
|
92
93
|
|
|
@@ -116,7 +117,7 @@ const agent = new Agent({
|
|
|
116
117
|
model: ({ requestContext }) => {
|
|
117
118
|
const useAdvanced = requestContext.task === "complex";
|
|
118
119
|
return useAdvanced
|
|
119
|
-
? "crossmodel/z-ai/glm-5.
|
|
120
|
+
? "crossmodel/z-ai/glm-5.3"
|
|
120
121
|
: "crossmodel/anthropic/claude-fable-5";
|
|
121
122
|
}
|
|
122
123
|
});
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Eden AI
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 233 Eden AI models through Mastra's model router. Authentication is handled automatically using the `EDENAI_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Eden AI documentation](https://docs.edenai.co).
|
|
8
8
|
|
|
@@ -66,7 +66,6 @@ for await (const chunk of stream) {
|
|
|
66
66
|
| `edenai/cloudflare/@cf/openai/gpt-oss-120b` | 128K | | | | | | $0.35 | $0.75 |
|
|
67
67
|
| `edenai/cloudflare/@cf/openai/gpt-oss-20b` | 128K | | | | | | $0.20 | $0.30 |
|
|
68
68
|
| `edenai/cloudflare/@cf/qwen/qwen2.5-coder-32b-instruct` | 33K | | | | | | $0.66 | $1 |
|
|
69
|
-
| `edenai/cloudflare/@cf/qwen/qwen3.8-27b` | 262K | | | | | | $0.45 | $3 |
|
|
70
69
|
| `edenai/cloudflare/@cf/zai-org/glm-4.7-flash` | 131K | | | | | | $0.06 | $0.40 |
|
|
71
70
|
| `edenai/cohere/command-a-03-2025` | 288K | | | | | | $3 | $10 |
|
|
72
71
|
| `edenai/cohere/command-r-08-2024` | 128K | | | | | | $0.15 | $0.60 |
|
|
@@ -91,7 +90,6 @@ for await (const chunk of stream) {
|
|
|
91
90
|
| `edenai/deepinfra/nvidia/Nemotron-3-Nano-30B-A3B` | 262K | | | | | | $0.05 | $0.20 |
|
|
92
91
|
| `edenai/deepinfra/openai/gpt-oss-120b` | 131K | | | | | | $0.04 | $0.17 |
|
|
93
92
|
| `edenai/deepinfra/openai/gpt-oss-20b` | 131K | | | | | | $0.03 | $0.14 |
|
|
94
|
-
| `edenai/deepinfra/Qwen/Qwen3.8-27B` | 262K | | | | | | $0.40 | $3 |
|
|
95
93
|
| `edenai/deepinfra/stepfun-ai/Step-3.5-Flash` | 262K | | | | | | $0.09 | $0.30 |
|
|
96
94
|
| `edenai/deepinfra/stepfun-ai/Step-3.7-Flash` | 262K | | | | | | $0.20 | $1 |
|
|
97
95
|
| `edenai/deepinfra/tencent/Hy3` | 262K | | | | | | $0.14 | $0.58 |
|
|
@@ -145,7 +143,7 @@ for await (const chunk of stream) {
|
|
|
145
143
|
| `edenai/mistral/codestral-latest` | 256K | | | | | | $1 | $3 |
|
|
146
144
|
| `edenai/mistral/devstral-2512` | 262K | | | | | | $0.40 | $2 |
|
|
147
145
|
| `edenai/mistral/devstral-medium-latest` | 262K | | | | | | $0.40 | $2 |
|
|
148
|
-
| `edenai/mistral/magistral-medium-latest` |
|
|
146
|
+
| `edenai/mistral/magistral-medium-latest` | 262K | | | | | | $2 | $5 |
|
|
149
147
|
| `edenai/mistral/mistral-large-2512` | 262K | | | | | | $0.50 | $2 |
|
|
150
148
|
| `edenai/mistral/mistral-large-latest` | 262K | | | | | | $2 | $6 |
|
|
151
149
|
| `edenai/mistral/mistral-medium-2505` | 131K | | | | | | $0.40 | $2 |
|
|
@@ -203,7 +201,7 @@ for await (const chunk of stream) {
|
|
|
203
201
|
| `edenai/perplexityai/sonar-pro` | 200K | | | | | | $3 | $15 |
|
|
204
202
|
| `edenai/perplexityai/sonar-reasoning-pro` | 128K | | | | | | $2 | $8 |
|
|
205
203
|
| `edenai/qwen/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.20 | $0.40 |
|
|
206
|
-
| `edenai/qwen/deepseek-v4-pro-0813` | 1.0M | | | | | | $
|
|
204
|
+
| `edenai/qwen/deepseek-v4-pro-0813` | 1.0M | | | | | | $0.73 | $2 |
|
|
207
205
|
| `edenai/qwen/qwen-max` | 33K | | | | | | $2 | $6 |
|
|
208
206
|
| `edenai/qwen/qwen-vl-max` | 131K | | | | | | $0.80 | $3 |
|
|
209
207
|
| `edenai/qwen/qwen-vl-plus` | 131K | | | | | | $0.21 | $0.63 |
|
|
@@ -219,14 +217,14 @@ for await (const chunk of stream) {
|
|
|
219
217
|
| `edenai/qwen/qwen3-vl-235b-a22b-instruct` | 131K | | | | | | $0.40 | $2 |
|
|
220
218
|
| `edenai/qwen/qwen3-vl-235b-a22b-thinking` | 131K | | | | | | $0.40 | $4 |
|
|
221
219
|
| `edenai/qwen/qwen3.8-2.4t-a95b` | 1.0M | | | | | | $2 | $6 |
|
|
220
|
+
| `edenai/qwen/qwen3.8-27b` | 1.0M | | | | | | $0.50 | $3 |
|
|
222
221
|
| `edenai/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
|
|
223
222
|
| `edenai/qwen/qwq-plus` | 131K | | | | | | $0.80 | $2 |
|
|
224
223
|
| `edenai/scaleway/deepseek-v4-flash-0731` | 256K | | | | | | $0.46 | $0.93 |
|
|
225
|
-
| `edenai/scaleway/gpt-oss-120b` | 128K | | | | | | $0.17 | $0.
|
|
224
|
+
| `edenai/scaleway/gpt-oss-120b` | 128K | | | | | | $0.17 | $0.70 |
|
|
226
225
|
| `edenai/scaleway/llama-3.3-70b-instruct` | 128K | | | | | | $1 | $1 |
|
|
227
226
|
| `edenai/tensorx/deepseek/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.25 | $0.30 |
|
|
228
227
|
| `edenai/tensorx/moonshotai/kimi-k2.5` | 262K | | | | | | $0.50 | $3 |
|
|
229
|
-
| `edenai/tensorx/qwen/qwen3.8-27b` | 262K | | | | | | $0.40 | $2 |
|
|
230
228
|
| `edenai/together_ai/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
231
229
|
| `edenai/together_ai/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
232
230
|
| `edenai/together_ai/meta-models/Muse-Glimmer-30B` | 131K | | | | | | $0.35 | $2 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# EmpirioLabs AI
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 55 EmpirioLabs AI models through Mastra's model router. Authentication is handled automatically using the `EMPIRIOLABS_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [EmpirioLabs AI documentation](https://docs.empiriolabs.ai).
|
|
8
8
|
|
|
@@ -45,6 +45,7 @@ for await (const chunk of stream) {
|
|
|
45
45
|
| `empiriolabs/fugu-ultra-v1-1` | 1.0M | | | | | | $5 | $30 |
|
|
46
46
|
| `empiriolabs/gemma-4-26b-a4b` | 262K | | | | | | $0.05 | $0.29 |
|
|
47
47
|
| `empiriolabs/glm-4-5-flash` | 200K | | | | | | — | — |
|
|
48
|
+
| `empiriolabs/glm-4-6v-flash` | 128K | | | | | | — | — |
|
|
48
49
|
| `empiriolabs/glm-4-7-flash` | 200K | | | | | | — | — |
|
|
49
50
|
| `empiriolabs/glm-5-1` | 202K | | | | | | $0.82 | $3 |
|
|
50
51
|
| `empiriolabs/glm-5-2` | 1.0M | | | | | | $1 | $4 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Hugging Face
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 69 Hugging Face models through Mastra's model router. Authentication is handled automatically using the `HF_TOKEN` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Hugging Face documentation](https://huggingface.co).
|
|
8
8
|
|
|
@@ -99,6 +99,7 @@ for await (const chunk of stream) {
|
|
|
99
99
|
| `huggingface/zai-org/GLM-4.5-Air` | 131K | | | | | | $0.13 | $0.85 |
|
|
100
100
|
| `huggingface/zai-org/GLM-4.5V` | 66K | | | | | | $0.60 | $2 |
|
|
101
101
|
| `huggingface/zai-org/GLM-4.6` | 205K | | | | | | $0.55 | $2 |
|
|
102
|
+
| `huggingface/zai-org/GLM-4.6V-Flash` | 131K | | | | | | $0.30 | $0.90 |
|
|
102
103
|
| `huggingface/zai-org/GLM-4.7` | 205K | | | | | | $0.60 | $2 |
|
|
103
104
|
| `huggingface/zai-org/GLM-4.7-Flash` | 200K | | | | | | — | — |
|
|
104
105
|
| `huggingface/zai-org/GLM-5` | 203K | | | | | | $1 | $3 |
|
|
@@ -37,21 +37,21 @@ for await (const chunk of stream) {
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ---------------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
39
|
| `hyper/deepseek-v4-flash` | 1.0M | | | | | | $0.20 | $0.40 |
|
|
40
|
-
| `hyper/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.
|
|
40
|
+
| `hyper/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.44 | $1 |
|
|
41
41
|
| `hyper/deepseek-v4-pro` | 1.0M | | | | | | $2 | $5 |
|
|
42
42
|
| `hyper/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
43
43
|
| `hyper/gemma-4-26b-a4b-it` | 256K | | | | | | $0.12 | $0.42 |
|
|
44
|
-
| `hyper/glm-5` | 203K | | | | | | $0.
|
|
44
|
+
| `hyper/glm-5` | 203K | | | | | | $0.86 | $3 |
|
|
45
45
|
| `hyper/glm-5.1` | 203K | | | | | | $1 | $4 |
|
|
46
|
-
| `hyper/glm-5.2` | 1.0M | | | | | | $
|
|
47
|
-
| `hyper/gpt-oss-120b` | 131K | | | | | | $0.
|
|
48
|
-
| `hyper/kimi-k2.5` | 262K | | | | | | $0.
|
|
49
|
-
| `hyper/kimi-k2.6` | 262K | | | | | | $
|
|
50
|
-
| `hyper/kimi-k2.7-code` | 256K | | | | | | $
|
|
46
|
+
| `hyper/glm-5.2` | 1.0M | | | | | | $2 | $5 |
|
|
47
|
+
| `hyper/gpt-oss-120b` | 131K | | | | | | $0.19 | $0.70 |
|
|
48
|
+
| `hyper/kimi-k2.5` | 262K | | | | | | $0.52 | $3 |
|
|
49
|
+
| `hyper/kimi-k2.6` | 262K | | | | | | $1 | $4 |
|
|
50
|
+
| `hyper/kimi-k2.7-code` | 256K | | | | | | $1 | $4 |
|
|
51
51
|
| `hyper/kimi-k3` | 1.0M | | | | | | $3 | $16 |
|
|
52
|
-
| `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.
|
|
52
|
+
| `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.60 | $0.74 |
|
|
53
53
|
| `hyper/llama-4-maverick-17b-128e-instruct-fp8` | 430K | | | | | | $0.27 | $0.90 |
|
|
54
|
-
| `hyper/minimax-m2.7` | 262K | | | | | | $0.
|
|
54
|
+
| `hyper/minimax-m2.7` | 262K | | | | | | $0.41 | $2 |
|
|
55
55
|
| `hyper/minimax-m3` | 512K | | | | | | $0.33 | $1 |
|
|
56
56
|
| `hyper/qwen3-coder-480b-a35b-instruct-int4-mixed-ar` | 106K | | | | | | $0.45 | $2 |
|
|
57
57
|
| `hyper/qwen3-next-80b-a3b-instruct` | 262K | | | | | | $0.12 | $1 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Inceptron
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 4 Inceptron models through Mastra's model router. Authentication is handled automatically using the `INCEPTRON_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Inceptron documentation](https://docs.inceptron.io).
|
|
8
8
|
|
|
@@ -17,7 +17,7 @@ const agent = new Agent({
|
|
|
17
17
|
id: "my-agent",
|
|
18
18
|
name: "My Agent",
|
|
19
19
|
instructions: "You are a helpful assistant",
|
|
20
|
-
model: "inceptron/
|
|
20
|
+
model: "inceptron/deepseek-ai/DeepSeek-V4-Flash-0731"
|
|
21
21
|
});
|
|
22
22
|
|
|
23
23
|
// Generate a response
|
|
@@ -37,7 +37,6 @@ for await (const chunk of stream) {
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ---------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
39
|
| `inceptron/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.13 | $0.28 |
|
|
40
|
-
| `inceptron/MiniMaxAI/MiniMax-M2.5` | 197K | | | | | | $0.22 | $0.90 |
|
|
41
40
|
| `inceptron/moonshotai/Kimi-K2.6` | 262K | | | | | | $0.60 | $3 |
|
|
42
41
|
| `inceptron/moonshotai/Kimi-K2.7-Code` | 262K | | | | | | $0.67 | $3 |
|
|
43
42
|
| `inceptron/zai-org/GLM-5.2` | 1.0M | | | | | | $0.75 | $3 |
|
|
@@ -52,7 +51,7 @@ const agent = new Agent({
|
|
|
52
51
|
name: "custom-agent",
|
|
53
52
|
model: {
|
|
54
53
|
url: "https://api.inceptron.io/v1",
|
|
55
|
-
id: "inceptron/
|
|
54
|
+
id: "inceptron/deepseek-ai/DeepSeek-V4-Flash-0731",
|
|
56
55
|
apiKey: process.env.INCEPTRON_API_KEY,
|
|
57
56
|
headers: {
|
|
58
57
|
"X-Custom-Header": "value"
|
|
@@ -71,7 +70,7 @@ const agent = new Agent({
|
|
|
71
70
|
const useAdvanced = requestContext.task === "complex";
|
|
72
71
|
return useAdvanced
|
|
73
72
|
? "inceptron/zai-org/GLM-5.2"
|
|
74
|
-
: "inceptron/
|
|
73
|
+
: "inceptron/deepseek-ai/DeepSeek-V4-Flash-0731";
|
|
75
74
|
}
|
|
76
75
|
});
|
|
77
76
|
```
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Kilo Gateway
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 359 Kilo Gateway models through Mastra's model router. Authentication is handled automatically using the `KILO_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Kilo Gateway documentation](https://kilo.ai).
|
|
8
8
|
|
|
@@ -17,7 +17,7 @@ const agent = new Agent({
|
|
|
17
17
|
id: "my-agent",
|
|
18
18
|
name: "My Agent",
|
|
19
19
|
instructions: "You are a helpful assistant",
|
|
20
|
-
model: "kilo/
|
|
20
|
+
model: "kilo/aion-labs/aion-2.0"
|
|
21
21
|
});
|
|
22
22
|
|
|
23
23
|
// Generate a response
|
|
@@ -40,14 +40,14 @@ for await (const chunk of stream) {
|
|
|
40
40
|
| `kilo/~anthropic/claude-haiku-latest` | 200K | | | | | | $1 | $5 |
|
|
41
41
|
| `kilo/~anthropic/claude-opus-latest` | 1.0M | | | | | | $5 | $25 |
|
|
42
42
|
| `kilo/~anthropic/claude-sonnet-latest` | 1.0M | | | | | | $2 | $10 |
|
|
43
|
-
| `kilo/~deepseek/deepseek-v4-flash-latest` |
|
|
43
|
+
| `kilo/~deepseek/deepseek-v4-flash-latest` | 1.0M | | | | | | $0.07 | $0.14 |
|
|
44
44
|
| `kilo/~google/gemini-flash-latest` | 1.0M | | | | | | $0.38 | $2 |
|
|
45
45
|
| `kilo/~google/gemini-pro-latest` | 1.0M | | | | | | $2 | $12 |
|
|
46
46
|
| `kilo/~moonshotai/kimi-latest` | 975K | | | | | | $3 | $13 |
|
|
47
47
|
| `kilo/~openai/gpt-latest` | 1.1M | | | | | | $3 | $15 |
|
|
48
48
|
| `kilo/~openai/gpt-mini-latest` | 400K | | | | | | $0.75 | $5 |
|
|
49
49
|
| `kilo/~x-ai/grok-latest` | 500K | | | | | | $2 | $6 |
|
|
50
|
-
| `kilo/
|
|
50
|
+
| `kilo/~z-ai/glm-latest` | 1.0M | | | | | | $1 | $4 |
|
|
51
51
|
| `kilo/aion-labs/aion-2.0` | 131K | | | | | | $0.80 | $2 |
|
|
52
52
|
| `kilo/aion-labs/aion-3.0` | 131K | | | | | | $3 | $6 |
|
|
53
53
|
| `kilo/aion-labs/aion-3.0-mini` | 131K | | | | | | $0.70 | $1 |
|
|
@@ -155,7 +155,6 @@ for await (const chunk of stream) {
|
|
|
155
155
|
| `kilo/kwaipilot/kat-coder-pro-v2` | 256K | | | | | | $0.30 | $1 |
|
|
156
156
|
| `kilo/kwaipilot/kat-coder-pro-v2.5` | 256K | | | | | | $0.74 | $3 |
|
|
157
157
|
| `kilo/liquid/lfm-2.5-2.6b:free` | 128K | | | | | | — | — |
|
|
158
|
-
| `kilo/mancer/weaver` | 8K | | | | | | $0.50 | $0.75 |
|
|
159
158
|
| `kilo/meituan/longcat-2.0` | 1.0M | | | | | | $0.75 | $3 |
|
|
160
159
|
| `kilo/meta-llama/llama-3.1-70b-instruct` | 131K | | | | | | $0.40 | $0.40 |
|
|
161
160
|
| `kilo/meta-llama/llama-3.1-8b-instruct` | 131K | | | | | | $0.02 | $0.04 |
|
|
@@ -175,7 +174,7 @@ for await (const chunk of stream) {
|
|
|
175
174
|
| `kilo/minimax/minimax-m2` | 205K | | | | | | $0.30 | $1 |
|
|
176
175
|
| `kilo/minimax/minimax-m2-her` | 66K | | | | | | $0.30 | $1 |
|
|
177
176
|
| `kilo/minimax/minimax-m2.1` | 205K | | | | | | $0.30 | $1 |
|
|
178
|
-
| `kilo/minimax/minimax-m2.5` |
|
|
177
|
+
| `kilo/minimax/minimax-m2.5` | 66K | | | | | | $0.30 | $1 |
|
|
179
178
|
| `kilo/minimax/minimax-m2.7` | 205K | | | | | | $0.30 | $1 |
|
|
180
179
|
| `kilo/minimax/minimax-m3` | 524K | | | | | | $0.30 | $1 |
|
|
181
180
|
| `kilo/mistralai/codestral-2508` | 256K | | | | | | $0.30 | $0.90 |
|
|
@@ -261,6 +260,7 @@ for await (const chunk of stream) {
|
|
|
261
260
|
| `kilo/openai/gpt-5.6-luna` | 1.1M | | | | | | $0.20 | $1 |
|
|
262
261
|
| `kilo/openai/gpt-5.6-luna-pro` | 1.1M | | | | | | $0.20 | $1 |
|
|
263
262
|
| `kilo/openai/gpt-5.6-sol` | 1.1M | | | | | | $5 | $30 |
|
|
263
|
+
| `kilo/openai/gpt-5.6-sol-discounted` | 1.1M | | | | | | $3 | $15 |
|
|
264
264
|
| `kilo/openai/gpt-5.6-sol-pro` | 1.1M | | | | | | $5 | $30 |
|
|
265
265
|
| `kilo/openai/gpt-5.6-terra` | 1.1M | | | | | | $2 | $12 |
|
|
266
266
|
| `kilo/openai/gpt-5.6-terra-pro` | 1.1M | | | | | | $2 | $12 |
|
|
@@ -341,7 +341,7 @@ for await (const chunk of stream) {
|
|
|
341
341
|
| `kilo/qwen/qwen3.7-max` | 1.0M | | | | | | $1 | $4 |
|
|
342
342
|
| `kilo/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.32 | $1 |
|
|
343
343
|
| `kilo/qwen/qwen3.8-2.4t-a95b` | 1.0M | | | | | | $2 | $6 |
|
|
344
|
-
| `kilo/qwen/qwen3.8-27b` | 262K | | | | | | $0.
|
|
344
|
+
| `kilo/qwen/qwen3.8-27b` | 262K | | | | | | $0.55 | $3 |
|
|
345
345
|
| `kilo/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
|
|
346
346
|
| `kilo/rekaai/reka-edge` | 16K | | | | | | $0.10 | $0.10 |
|
|
347
347
|
| `kilo/rekaai/reka-flash-3` | 66K | | | | | | $0.10 | $0.20 |
|
|
@@ -356,13 +356,12 @@ for await (const chunk of stream) {
|
|
|
356
356
|
| `kilo/stealth/claude-opus-4.7` | 1.0M | | | | | | $4 | $20 |
|
|
357
357
|
| `kilo/stealth/claude-opus-4.8` | 1.0M | | | | | | $4 | $20 |
|
|
358
358
|
| `kilo/stealth/claude-sonnet-4.6` | 1.0M | | | | | | $2 | $12 |
|
|
359
|
-
| `kilo/stealth/gpt-5.6-sol` | 1.1M | | | | | | $4 | $24 |
|
|
360
359
|
| `kilo/stealth/qwen3.6-plus` | 1.0M | | | | | | $0.25 | $2 |
|
|
361
360
|
| `kilo/stepfun/step-3.5-flash` | 262K | | | | | | $0.10 | $0.30 |
|
|
362
361
|
| `kilo/stepfun/step-3.7-flash` | 256K | | | | | | $0.20 | $1 |
|
|
363
362
|
| `kilo/stepfun/step-3.7-flash:free` | 262K | | | | | | — | — |
|
|
364
363
|
| `kilo/tencent/hunyuan-a13b-instruct` | 131K | | | | | | $0.14 | $0.57 |
|
|
365
|
-
| `kilo/tencent/hy3` | 262K | | | | | | $0.
|
|
364
|
+
| `kilo/tencent/hy3` | 262K | | | | | | $0.13 | $0.53 |
|
|
366
365
|
| `kilo/tencent/hy3-preview` | 262K | | | | | | $0.18 | $0.60 |
|
|
367
366
|
| `kilo/tencent/hy3:free` | 262K | | | | | | — | — |
|
|
368
367
|
| `kilo/thedrummer/cydonia-24b-v4.1` | 131K | | | | | | $0.30 | $0.50 |
|
|
@@ -407,7 +406,7 @@ const agent = new Agent({
|
|
|
407
406
|
name: "custom-agent",
|
|
408
407
|
model: {
|
|
409
408
|
url: "https://api.kilo.ai/api/gateway",
|
|
410
|
-
id: "kilo/
|
|
409
|
+
id: "kilo/aion-labs/aion-2.0",
|
|
411
410
|
apiKey: process.env.KILO_API_KEY,
|
|
412
411
|
headers: {
|
|
413
412
|
"X-Custom-Header": "value"
|
|
@@ -425,8 +424,8 @@ const agent = new Agent({
|
|
|
425
424
|
model: ({ requestContext }) => {
|
|
426
425
|
const useAdvanced = requestContext.task === "complex";
|
|
427
426
|
return useAdvanced
|
|
428
|
-
? "kilo/~
|
|
429
|
-
: "kilo/
|
|
427
|
+
? "kilo/~z-ai/glm-latest"
|
|
428
|
+
: "kilo/aion-labs/aion-2.0";
|
|
430
429
|
}
|
|
431
430
|
});
|
|
432
431
|
```
|
|
@@ -37,7 +37,6 @@ for await (const chunk of stream) {
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ---------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
39
|
| `llmgateway/auto` | 128K | | | | | | — | — |
|
|
40
|
-
| `llmgateway/claude-3-opus` | 200K | | | | | | $15 | $75 |
|
|
41
40
|
| `llmgateway/claude-fable-5` | 1.0M | | | | | | $10 | $50 |
|
|
42
41
|
| `llmgateway/claude-haiku-4-5` | 200K | | | | | | $1 | $5 |
|
|
43
42
|
| `llmgateway/claude-haiku-4-5-20251001` | 200K | | | | | | $1 | $5 |
|
|
@@ -55,7 +54,7 @@ for await (const chunk of stream) {
|
|
|
55
54
|
| `llmgateway/cosmos3-super-reasoner` | 262K | | | | | | $0.10 | $0.30 |
|
|
56
55
|
| `llmgateway/custom` | 128K | | | | | | — | — |
|
|
57
56
|
| `llmgateway/deepseek-v3.2` | 164K | | | | | | $0.26 | $0.38 |
|
|
58
|
-
| `llmgateway/deepseek-v4-flash` | 1.1M | | | | | | $0.
|
|
57
|
+
| `llmgateway/deepseek-v4-flash` | 1.1M | | | | | | $0.08 | $0.15 |
|
|
59
58
|
| `llmgateway/deepseek-v4-pro` | 1.1M | | | | | | $0.43 | $0.87 |
|
|
60
59
|
| `llmgateway/ernie-4.5-vl-424b-a47b` | 123K | | | | | | $0.42 | $1 |
|
|
61
60
|
| `llmgateway/fugu-ultra` | 1.0M | | | | | | $5 | $30 |
|
|
@@ -88,6 +87,7 @@ for await (const chunk of stream) {
|
|
|
88
87
|
| `llmgateway/glm-5` | 203K | | | | | | $0.72 | $2 |
|
|
89
88
|
| `llmgateway/glm-5.1` | 205K | | | | | | $0.93 | $3 |
|
|
90
89
|
| `llmgateway/glm-5.2` | 1.0M | | | | | | $0.55 | $2 |
|
|
90
|
+
| `llmgateway/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
91
91
|
| `llmgateway/gpt-3.5-turbo` | 16K | | | | | | $0.50 | $2 |
|
|
92
92
|
| `llmgateway/gpt-4` | 8K | | | | | | $30 | $60 |
|
|
93
93
|
| `llmgateway/gpt-4-turbo` | 128K | | | | | | $10 | $30 |
|
|
@@ -63,7 +63,7 @@ for await (const chunk of stream) {
|
|
|
63
63
|
| `ofox/bailian/qwen3.5-plus` | 1.0M | | | | | | $0.40 | $2 |
|
|
64
64
|
| `ofox/bailian/qwen3.6-27b` | 256K | | | | | | $0.60 | $4 |
|
|
65
65
|
| `ofox/bailian/qwen3.6-flash` | 1.0M | | | | | | $0.25 | $2 |
|
|
66
|
-
| `ofox/bailian/qwen3.6-max-preview` | 262K | | | | | | $2 | $
|
|
66
|
+
| `ofox/bailian/qwen3.6-max-preview` | 262K | | | | | | $2 | $13 |
|
|
67
67
|
| `ofox/bailian/qwen3.6-plus` | 1.0M | | | | | | $0.50 | $3 |
|
|
68
68
|
| `ofox/bailian/qwen3.7-max` | 1.0M | | | | | | $3 | $8 |
|
|
69
69
|
| `ofox/bailian/qwen3.7-plus` | 1.0M | | | | | | $0.40 | $2 |
|
|
@@ -125,8 +125,8 @@ for await (const chunk of stream) {
|
|
|
125
125
|
| `ofox/volcengine/doubao-seed-2.0-lite` | 256K | | | | | | $0.13 | $0.76 |
|
|
126
126
|
| `ofox/volcengine/doubao-seed-2.0-mini` | 256K | | | | | | $0.06 | $0.56 |
|
|
127
127
|
| `ofox/volcengine/doubao-seed-2.0-pro` | 256K | | | | | | $0.67 | $3 |
|
|
128
|
-
| `ofox/volcengine/doubao-seed-2.1-pro` | 256K | | | | | | $0.
|
|
129
|
-
| `ofox/volcengine/doubao-seed-2.1-turbo` | 256K | | | | | | $0.
|
|
128
|
+
| `ofox/volcengine/doubao-seed-2.1-pro` | 256K | | | | | | $0.71 | $4 |
|
|
129
|
+
| `ofox/volcengine/doubao-seed-2.1-turbo` | 256K | | | | | | $0.35 | $2 |
|
|
130
130
|
| `ofox/volcengine/doubao-seed-character` | 256K | | | | | | $0.18 | $0.88 |
|
|
131
131
|
| `ofox/volcengine/doubao-seed-evolving` | 256K | | | | | | $0.88 | $4 |
|
|
132
132
|
| `ofox/x-ai/grok-4.1-fast` | 2.0M | | | | | | $0.20 | $0.50 |
|
|
@@ -140,7 +140,7 @@ for await (const chunk of stream) {
|
|
|
140
140
|
| `ofox/z-ai/glm-5` | 205K | | | | | | $1 | $3 |
|
|
141
141
|
| `ofox/z-ai/glm-5-turbo` | 200K | | | | | | $1 | $4 |
|
|
142
142
|
| `ofox/z-ai/glm-5.1` | 200K | | | | | | $1 | $4 |
|
|
143
|
-
| `ofox/z-ai/glm-5.2` | 1.0M | | | | | | $
|
|
143
|
+
| `ofox/z-ai/glm-5.2` | 1.0M | | | | | | $0.98 | $3 |
|
|
144
144
|
| `ofox/z-ai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
145
145
|
| `ofox/z-ai/glm-5v-turbo` | 200K | | | | | | $1 | $4 |
|
|
146
146
|
|
|
@@ -43,7 +43,7 @@ for await (const chunk of stream) {
|
|
|
43
43
|
| `opencode-go/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
44
44
|
| `opencode-go/gpt-5.6-luna` | 1.1M | | | | | | $0.20 | $1 |
|
|
45
45
|
| `opencode-go/grok-4.5` | 500K | | | | | | $2 | $6 |
|
|
46
|
-
| `opencode-go/hy3` | 256K | | | | | | $0.
|
|
46
|
+
| `opencode-go/hy3` | 256K | | | | | | $0.02 | $0.07 |
|
|
47
47
|
| `opencode-go/kimi-k2.6` | 262K | | | | | | $0.95 | $4 |
|
|
48
48
|
| `opencode-go/kimi-k2.7-code` | 262K | | | | | | $0.95 | $4 |
|
|
49
49
|
| `opencode-go/kimi-k3` | 1.0M | | | | | | $3 | $15 |
|
|
@@ -78,7 +78,7 @@ for await (const chunk of stream) {
|
|
|
78
78
|
| `opencode/gpt-5.5` | 1.1M | | | | | | $5 | $30 |
|
|
79
79
|
| `opencode/gpt-5.5-pro` | 1.1M | | | | | | $30 | $180 |
|
|
80
80
|
| `opencode/gpt-5.6-luna` | 1.1M | | | | | | $0.20 | $1 |
|
|
81
|
-
| `opencode/gpt-5.6-sol` | 1.1M | | | | | | $
|
|
81
|
+
| `opencode/gpt-5.6-sol` | 1.1M | | | | | | $3 | $15 |
|
|
82
82
|
| `opencode/gpt-5.6-terra` | 1.1M | | | | | | $3 | $15 |
|
|
83
83
|
| `opencode/grok-4.5` | 500K | | | | | | $2 | $6 |
|
|
84
84
|
| `opencode/grok-4.6` | 500K | | | | | | $2 | $6 |
|
|
@@ -90,29 +90,29 @@ for await (const chunk of stream) {
|
|
|
90
90
|
| `pioneer/grok-4.5` | 500K | | | | | | $2 | $6 |
|
|
91
91
|
| `pioneer/HuggingFaceTB/SmolLM3-3B-Base` | 33K | | | | | | $0.15 | $0.15 |
|
|
92
92
|
| `pioneer/LiquidAI/LFM2-24B-A2B` | 33K | | | | | | $0.03 | $0.12 |
|
|
93
|
-
| `pioneer/meta-llama/Llama-3.1-8B-Instruct` |
|
|
93
|
+
| `pioneer/meta-llama/Llama-3.1-8B-Instruct` | 128K | | | | | | $0.20 | $0.20 |
|
|
94
94
|
| `pioneer/meta-llama/Llama-3.2-1B` | 131K | | | | | | $0.10 | $0.10 |
|
|
95
95
|
| `pioneer/meta-llama/Llama-3.2-1B-Instruct` | 131K | | | | | | $0.10 | $0.20 |
|
|
96
96
|
| `pioneer/meta-llama/Llama-3.2-3B` | 131K | | | | | | $0.10 | $0.10 |
|
|
97
97
|
| `pioneer/meta-llama/Llama-3.2-3B-Instruct` | 131K | | | | | | $0.10 | $0.34 |
|
|
98
|
-
| `pioneer/meta-llama/Llama-3.3-70B-Instruct` |
|
|
98
|
+
| `pioneer/meta-llama/Llama-3.3-70B-Instruct` | 16K | | | | | | $0.90 | $0.90 |
|
|
99
99
|
| `pioneer/meta/muse-spark-1.1` | 1.0M | | | | | | $1 | $4 |
|
|
100
100
|
| `pioneer/MiniMaxAI/MiniMax-M2.7` | 205K | | | | | | $0.28 | $1 |
|
|
101
101
|
| `pioneer/MiniMaxAI/MiniMax-M3` | 1.0M | | | | | | $0.30 | $1 |
|
|
102
102
|
| `pioneer/mistral-large-3` | 256K | | | | | | $0.50 | $2 |
|
|
103
|
-
| `pioneer/mistral-medium-3.5` |
|
|
103
|
+
| `pioneer/mistral-medium-3.5` | 256K | | | | | | $2 | $8 |
|
|
104
104
|
| `pioneer/mistralai/Codestral-22B-v0.1` | 128K | | | | | | $0.30 | $0.90 |
|
|
105
105
|
| `pioneer/mistralai/Magistral-Small-2506` | 128K | | | | | | $0.50 | $2 |
|
|
106
106
|
| `pioneer/mistralai/Ministral-8B-Instruct-2410` | 128K | | | | | | $0.15 | $0.15 |
|
|
107
107
|
| `pioneer/mistralai/Mistral-7B-Instruct-v0.3` | 33K | | | | | | $0.20 | $0.20 |
|
|
108
|
-
| `pioneer/mistralai/Mistral-Nemo-Instruct-2407` |
|
|
109
|
-
| `pioneer/mistralai/Mistral-Small-4-119B-2603` |
|
|
108
|
+
| `pioneer/mistralai/Mistral-Nemo-Instruct-2407` | 128K | | | | | | $0.02 | $0.03 |
|
|
109
|
+
| `pioneer/mistralai/Mistral-Small-4-119B-2603` | 32K | | | | | | $0.15 | $0.60 |
|
|
110
110
|
| `pioneer/mistralai/Pixtral-12B-2409` | 128K | | | | | | $0.15 | $0.15 |
|
|
111
111
|
| `pioneer/moonshotai/Kimi-K2.6` | 262K | | | | | | $0.95 | $4 |
|
|
112
|
-
| `pioneer/moonshotai/Kimi-K2.7-Code` |
|
|
112
|
+
| `pioneer/moonshotai/Kimi-K2.7-Code` | 256K | | | | | | $0.95 | $4 |
|
|
113
113
|
| `pioneer/moonshotai/Kimi-K3` | 1.0M | | | | | | $3 | $15 |
|
|
114
114
|
| `pioneer/nvidia/NVIDIA-Nemotron-3-Nano-30B-A3B-BF16` | 262K | | | | | | $0.05 | $0.20 |
|
|
115
|
-
| `pioneer/nvidia/NVIDIA-Nemotron-3-Super-120B-A12B-FP8` |
|
|
115
|
+
| `pioneer/nvidia/NVIDIA-Nemotron-3-Super-120B-A12B-FP8` | 256K | | | | | | $0.09 | $0.45 |
|
|
116
116
|
| `pioneer/nvidia/NVIDIA-Nemotron-3-Ultra-550B-A55B-BF16` | 1.0M | | | | | | $0.50 | $3 |
|
|
117
117
|
| `pioneer/openai/gpt-oss-120b` | 131K | | | | | | $0.15 | $0.60 |
|
|
118
118
|
| `pioneer/openai/gpt-oss-20b` | 131K | | | | | | $0.07 | $0.30 |
|
|
@@ -123,20 +123,20 @@ for await (const chunk of stream) {
|
|
|
123
123
|
| `pioneer/Qwen/Qwen3-235B-A22B-Instruct-2507` | 262K | | | | | | $1 | $1 |
|
|
124
124
|
| `pioneer/Qwen/Qwen3-32B` | 131K | | | | | | $0.90 | $0.90 |
|
|
125
125
|
| `pioneer/Qwen/Qwen3-4B-Base` | 33K | | | | | | $0.15 | $0.15 |
|
|
126
|
-
| `pioneer/Qwen/Qwen3-4B-Instruct-2507` |
|
|
127
|
-
| `pioneer/Qwen/Qwen3-8B` |
|
|
126
|
+
| `pioneer/Qwen/Qwen3-4B-Instruct-2507` | 33K | | | | | | $0.20 | $0.20 |
|
|
127
|
+
| `pioneer/Qwen/Qwen3-8B` | 41K | | | | | | $0.20 | $0.20 |
|
|
128
128
|
| `pioneer/Qwen/Qwen3.5-9B` | 33K | | | | | | $0.30 | $0.30 |
|
|
129
129
|
| `pioneer/Qwen/Qwen3.6-27B` | 33K | | | | | | $0.60 | $0.60 |
|
|
130
130
|
| `pioneer/Qwen/Qwen3.6-35B-A3B` | 262K | | | | | | $0.14 | $1 |
|
|
131
131
|
| `pioneer/qwen3.6-flash` | 1.0M | | | | | | $0.19 | $1 |
|
|
132
|
-
| `pioneer/qwen3.6-max-preview` |
|
|
132
|
+
| `pioneer/qwen3.6-max-preview` | 240K | | | | | | $1 | $6 |
|
|
133
133
|
| `pioneer/qwen3.6-plus` | 1.0M | | | | | | $0.33 | $2 |
|
|
134
|
-
| `pioneer/qwen3.7-max` |
|
|
134
|
+
| `pioneer/qwen3.7-max` | 991K | | | | | | $1 | $4 |
|
|
135
135
|
| `pioneer/qwen3.7-plus` | 1.0M | | | | | | $0.32 | $1 |
|
|
136
136
|
| `pioneer/sakana/fugu-ultra` | 1.0M | | | | | | $5 | $30 |
|
|
137
137
|
| `pioneer/XiaomiMiMo/MiMo-V2.5` | 1.1M | | | | | | $0.14 | $0.28 |
|
|
138
138
|
| `pioneer/XiaomiMiMo/MiMo-V2.5-Pro` | 1.1M | | | | | | $0.43 | $0.87 |
|
|
139
|
-
| `pioneer/zai-org/GLM-5.1` |
|
|
139
|
+
| `pioneer/zai-org/GLM-5.1` | 202K | | | | | | $0.98 | $3 |
|
|
140
140
|
| `pioneer/zai-org/GLM-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
141
141
|
|
|
142
142
|
## Advanced configuration
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# RunInfra
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 5 RunInfra models through Mastra's model router. Authentication is handled automatically using the `RUNINFRA_GATEWAY_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [RunInfra documentation](https://runinfra.ai/docs).
|
|
8
8
|
|
|
@@ -37,6 +37,7 @@ for await (const chunk of stream) {
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ------------------------------------------------------------ | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
39
|
| `runinfra/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.13 | $0.27 |
|
|
40
|
+
| `runinfra/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $0.60 | $2 |
|
|
40
41
|
| `runinfra/Inferact/Qwen3.8-2.4T-A95B-NVFP4` | 262K | | | | | | $2 | $6 |
|
|
41
42
|
| `runinfra/nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-BF16` | 262K | | | | | | $0.05 | $0.15 |
|
|
42
43
|
| `runinfra/Qwen/Qwen3.8-27B` | 262K | | | | | | $0.10 | $0.40 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Umans AI Coding Plan
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 8 Umans AI Coding Plan models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_CODING_PLAN_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Umans AI Coding Plan documentation](https://app.umans.ai/offers/code/docs).
|
|
8
8
|
|
|
@@ -40,7 +40,6 @@ for await (const chunk of stream) {
|
|
|
40
40
|
| `umans-ai-coding-plan/umans-deepseek-v4-flash-0731` | 1.0M | | | | | | — | — |
|
|
41
41
|
| `umans-ai-coding-plan/umans-deepseek-v4-pro-0813` | 1.0M | | | | | | — | — |
|
|
42
42
|
| `umans-ai-coding-plan/umans-flash` | 262K | | | | | | — | — |
|
|
43
|
-
| `umans-ai-coding-plan/umans-glm-5.1` | 205K | | | | | | — | — |
|
|
44
43
|
| `umans-ai-coding-plan/umans-glm-5.2` | 406K | | | | | | — | — |
|
|
45
44
|
| `umans-ai-coding-plan/umans-kimi-k2.7` | 262K | | | | | | — | — |
|
|
46
45
|
| `umans-ai-coding-plan/umans-kimi-k3` | 1.0M | | | | | | — | — |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Umans AI
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 7 Umans AI models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Umans AI documentation](https://app.umans.ai/offers/code/docs/orgs).
|
|
8
8
|
|
|
@@ -40,7 +40,6 @@ for await (const chunk of stream) {
|
|
|
40
40
|
| `umans-ai/umans-deepseek-v4-flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
41
41
|
| `umans-ai/umans-deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
42
42
|
| `umans-ai/umans-flash` | 262K | | | | | | $0.15 | $1 |
|
|
43
|
-
| `umans-ai/umans-glm-5.1` | 205K | | | | | | $1 | $4 |
|
|
44
43
|
| `umans-ai/umans-glm-5.2` | 406K | | | | | | $1 | $4 |
|
|
45
44
|
| `umans-ai/umans-kimi-k2.7` | 262K | | | | | | $0.95 | $4 |
|
|
46
45
|
| `umans-ai/umans-kimi-k3` | 1.0M | | | | | | $3 | $15 |
|
|
@@ -292,7 +292,7 @@ interface ActiveThreadRun {
|
|
|
292
292
|
|
|
293
293
|
#### `init()`
|
|
294
294
|
|
|
295
|
-
Initialize shared storage,
|
|
295
|
+
Initialize shared storage, propagate runtime services to agents, and start configured interval handlers. Configured workspaces initialize lazily when used. Repeated calls reuse the same initialization promise.
|
|
296
296
|
|
|
297
297
|
```typescript
|
|
298
298
|
await controller.init()
|
|
@@ -6,7 +6,7 @@ Framework-agnostic handler for streaming agent chat in AI SDK-compatible format.
|
|
|
6
6
|
|
|
7
7
|
`handleChatStream()` returns a `ReadableStream` that you can wrap with [`createUIMessageStreamResponse()`](https://ai-sdk.dev/docs/reference/ai-sdk-ui/create-ui-message-stream-response).
|
|
8
8
|
|
|
9
|
-
`handleChatStream()`
|
|
9
|
+
`handleChatStream()` defaults to AI SDK v5 for backward compatibility. Pass the version that matches your installed AI SDK, such as `version: 'v7'` for AI SDK v7.
|
|
10
10
|
|
|
11
11
|
Use [`chatRoute()`](https://mastra.ai/reference/ai-sdk/chat-route) if you want to create a chat route inside a Mastra server.
|
|
12
12
|
|
|
@@ -41,6 +41,7 @@ export async function POST(req: Request) {
|
|
|
41
41
|
const stream = await handleChatStream({
|
|
42
42
|
mastra,
|
|
43
43
|
agentId: 'weatherAgent',
|
|
44
|
+
version: 'v7',
|
|
44
45
|
params,
|
|
45
46
|
messageMetadata: () => ({ createdAt: new Date().toISOString() }),
|
|
46
47
|
})
|
package/CHANGELOG.md
CHANGED
|
@@ -1,5 +1,12 @@
|
|
|
1
1
|
# @mastra/mcp-docs-server
|
|
2
2
|
|
|
3
|
+
## 1.2.18-alpha.2
|
|
4
|
+
|
|
5
|
+
### Patch Changes
|
|
6
|
+
|
|
7
|
+
- Updated dependencies [[`d23e75d`](https://github.com/mastra-ai/mastra/commit/d23e75d57cc7cf5b9bfdbee896bf5a6a2484fed7), [`c8faa4e`](https://github.com/mastra-ai/mastra/commit/c8faa4e1cfebaec56b65e754e90b9fe46d153359), [`f2031a4`](https://github.com/mastra-ai/mastra/commit/f2031a47445e8f67a89ba1309036816f97ab7a65), [`8e529d4`](https://github.com/mastra-ai/mastra/commit/8e529d4ac754efef04b225841349e0da9edf89a6)]:
|
|
8
|
+
- @mastra/core@1.61.0-alpha.1
|
|
9
|
+
|
|
3
10
|
## 1.2.18-alpha.0
|
|
4
11
|
|
|
5
12
|
### Patch Changes
|
package/package.json
CHANGED
|
@@ -1,6 +1,6 @@
|
|
|
1
1
|
{
|
|
2
2
|
"name": "@mastra/mcp-docs-server",
|
|
3
|
-
"version": "1.2.18-alpha.
|
|
3
|
+
"version": "1.2.18-alpha.3",
|
|
4
4
|
"description": "MCP server for accessing Mastra.ai documentation, changelogs, and news.",
|
|
5
5
|
"type": "module",
|
|
6
6
|
"main": "dist/index.js",
|
|
@@ -28,8 +28,8 @@
|
|
|
28
28
|
"jsdom": "^26.1.0",
|
|
29
29
|
"local-pkg": "^1.1.2",
|
|
30
30
|
"zod": "^4.4.3",
|
|
31
|
-
"@mastra/
|
|
32
|
-
"@mastra/
|
|
31
|
+
"@mastra/core": "1.61.0-alpha.1",
|
|
32
|
+
"@mastra/mcp": "^1.17.1-alpha.0"
|
|
33
33
|
},
|
|
34
34
|
"devDependencies": {
|
|
35
35
|
"@hono/node-server": "^2.0.0",
|
|
@@ -46,8 +46,8 @@
|
|
|
46
46
|
"typescript": "^6.0.3",
|
|
47
47
|
"vitest": "4.1.10",
|
|
48
48
|
"@internal/lint": "0.0.124",
|
|
49
|
-
"@
|
|
50
|
-
"@
|
|
49
|
+
"@internal/types-builder": "0.0.99",
|
|
50
|
+
"@mastra/core": "1.61.0-alpha.1"
|
|
51
51
|
},
|
|
52
52
|
"homepage": "https://mastra.ai",
|
|
53
53
|
"repository": {
|