@mastra/mcp-docs-server 1.2.27-alpha.15 → 1.2.27-alpha.17
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/.docs/models/gateways/openrouter.md +2 -1
- package/.docs/models/index.md +1 -1
- package/.docs/models/providers/kilo.md +9 -9
- package/.docs/models/providers/nano-gpt.md +1 -1
- package/.docs/models/providers/opencode.md +5 -1
- package/.docs/models/providers/zai.md +3 -2
- package/.docs/models/providers/zhipuai.md +3 -2
- package/.docs/reference/agents/generate.md +1 -1
- package/.docs/reference/agents/network.md +1 -1
- package/.docs/reference/streaming/agents/stream.md +1 -1
- package/package.json +4 -4
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# OpenRouter
|
|
6
6
|
|
|
7
|
-
OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access
|
|
7
|
+
OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 372 models through Mastra's model router.
|
|
8
8
|
|
|
9
9
|
Learn more in the [OpenRouter documentation](https://openrouter.ai/models).
|
|
10
10
|
|
|
@@ -408,4 +408,5 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
408
408
|
| `z-ai/glm-5.2:free` |
|
|
409
409
|
| `z-ai/glm-5.3` |
|
|
410
410
|
| `z-ai/glm-5.3-flash` |
|
|
411
|
+
| `z-ai/glm-5.3-flashx` |
|
|
411
412
|
| `z-ai/glm-5v-turbo` |
|
package/.docs/models/index.md
CHANGED
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Model Providers
|
|
6
6
|
|
|
7
|
-
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to
|
|
7
|
+
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to 7359 models from 209 providers through a single API.
|
|
8
8
|
|
|
9
9
|
## Features
|
|
10
10
|
|
|
@@ -42,12 +42,12 @@ for await (const chunk of stream) {
|
|
|
42
42
|
| `kilo/~anthropic/claude-haiku-latest` | 200K | | | | | | $1 | $5 |
|
|
43
43
|
| `kilo/~anthropic/claude-opus-latest` | 1.0M | | | | | | $5 | $25 |
|
|
44
44
|
| `kilo/~anthropic/claude-sonnet-latest` | 1.0M | | | | | | $2 | $10 |
|
|
45
|
-
| `kilo/~deepseek/deepseek-flash-latest` | 1.0M | | | | | | $0.
|
|
45
|
+
| `kilo/~deepseek/deepseek-flash-latest` | 1.0M | | | | | | $0.13 | $0.52 |
|
|
46
46
|
| `kilo/~deepseek/deepseek-pro-latest` | 1.0M | | | | | | $0.58 | $2 |
|
|
47
|
-
| `kilo/~deepseek/deepseek-v4-flash-latest` | 1.0M | | | | | | $0.
|
|
47
|
+
| `kilo/~deepseek/deepseek-v4-flash-latest` | 1.0M | | | | | | $0.04 | $0.08 |
|
|
48
48
|
| `kilo/~google/gemini-flash-latest` | 1.0M | | | | | | $0.75 | $4 |
|
|
49
49
|
| `kilo/~google/gemini-pro-latest` | 1.0M | | | | | | $2 | $12 |
|
|
50
|
-
| `kilo/~moonshotai/kimi-latest` | 1.0M | | | | | | $2 | $
|
|
50
|
+
| `kilo/~moonshotai/kimi-latest` | 1.0M | | | | | | $2 | $9 |
|
|
51
51
|
| `kilo/~openai/gpt-astra-latest` | 1.1M | | | | | | $10 | $50 |
|
|
52
52
|
| `kilo/~openai/gpt-luna-latest` | 1.1M | | | | | | $0.20 | $1 |
|
|
53
53
|
| `kilo/~openai/gpt-mini-latest` | 400K | | | | | | $0.75 | $5 |
|
|
@@ -212,7 +212,7 @@ for await (const chunk of stream) {
|
|
|
212
212
|
| `kilo/moonshotai/kimi-k2.5` | 262K | | | | | | $0.60 | $3 |
|
|
213
213
|
| `kilo/moonshotai/kimi-k2.6` | 262K | | | | | | $0.80 | $3 |
|
|
214
214
|
| `kilo/moonshotai/kimi-k2.7-code` | 262K | | | | | | $0.95 | $4 |
|
|
215
|
-
| `kilo/moonshotai/kimi-k3` | 1.0M | | | | | | $2 | $
|
|
215
|
+
| `kilo/moonshotai/kimi-k3` | 1.0M | | | | | | $2 | $9 |
|
|
216
216
|
| `kilo/morph/morph-v3-fast` | 82K | | | | | | $0.80 | $1 |
|
|
217
217
|
| `kilo/morph/morph-v3-large` | 262K | | | | | | $0.90 | $2 |
|
|
218
218
|
| `kilo/nex-agi/nex-n2.5-mini:free` | 262K | | | | | | — | — |
|
|
@@ -224,11 +224,11 @@ for await (const chunk of stream) {
|
|
|
224
224
|
| `kilo/nvidia/nemotron-3-nano-omni-30b-a3b-reasoning:free` | 256K | | | | | | — | — |
|
|
225
225
|
| `kilo/nvidia/nemotron-3-super-120b-a12b` | 262K | | | | | | $0.08 | $0.45 |
|
|
226
226
|
| `kilo/nvidia/nemotron-3-super-120b-a12b:free` | 262K | | | | | | — | — |
|
|
227
|
-
| `kilo/nvidia/nemotron-3-ultra-550b-a55b` |
|
|
227
|
+
| `kilo/nvidia/nemotron-3-ultra-550b-a55b` | 203K | | | | | | $0.50 | $2 |
|
|
228
228
|
| `kilo/nvidia/nemotron-3-ultra-550b-a55b:free` | 1.0M | | | | | | — | — |
|
|
229
229
|
| `kilo/nvidia/nemotron-3.5-content-safety` | 131K | | | | | | $0.20 | $0.20 |
|
|
230
230
|
| `kilo/nvidia/nemotron-3.5-content-safety:free` | 128K | | | | | | — | — |
|
|
231
|
-
| `kilo/nvidia/nemotron-3.5-lightning` | 262K | | | | | | $0.
|
|
231
|
+
| `kilo/nvidia/nemotron-3.5-lightning` | 262K | | | | | | $0.04 | $0.18 |
|
|
232
232
|
| `kilo/nvidia/nemotron-3.5-lightning:free` | 1.0M | | | | | | — | — |
|
|
233
233
|
| `kilo/openai/gpt-3.5-turbo` | 16K | | | | | | $0.50 | $2 |
|
|
234
234
|
| `kilo/openai/gpt-3.5-turbo-0613` | 4K | | | | | | $1 | $2 |
|
|
@@ -270,7 +270,6 @@ for await (const chunk of stream) {
|
|
|
270
270
|
| `kilo/openai/gpt-5.6-luna` | 1.1M | | | | | | $0.20 | $1 |
|
|
271
271
|
| `kilo/openai/gpt-5.6-luna-pro` | 1.1M | | | | | | $0.20 | $1 |
|
|
272
272
|
| `kilo/openai/gpt-5.6-sol` | 1.1M | | | | | | $4 | $20 |
|
|
273
|
-
| `kilo/openai/gpt-5.6-sol-discounted` | 1.1M | | | | | | $2 | $10 |
|
|
274
273
|
| `kilo/openai/gpt-5.6-sol-pro` | 1.1M | | | | | | $4 | $20 |
|
|
275
274
|
| `kilo/openai/gpt-5.6-terra` | 1.1M | | | | | | $2 | $12 |
|
|
276
275
|
| `kilo/openai/gpt-5.6-terra-pro` | 1.1M | | | | | | $2 | $12 |
|
|
@@ -280,7 +279,7 @@ for await (const chunk of stream) {
|
|
|
280
279
|
| `kilo/openai/gpt-audio-mini` | 128K | | | | | | $0.60 | $2 |
|
|
281
280
|
| `kilo/openai/gpt-chat-latest` | 400K | | | | | | $5 | $30 |
|
|
282
281
|
| `kilo/openai/gpt-oss-120b` | 131K | | | | | | $0.03 | $0.17 |
|
|
283
|
-
| `kilo/openai/gpt-oss-20b` | 131K | | | | | | $0.02 | $0.
|
|
282
|
+
| `kilo/openai/gpt-oss-20b` | 131K | | | | | | $0.02 | $0.09 |
|
|
284
283
|
| `kilo/openai/gpt-oss-safeguard-20b` | 131K | | | | | | $0.07 | $0.30 |
|
|
285
284
|
| `kilo/openai/o1` | 200K | | | | | | $15 | $60 |
|
|
286
285
|
| `kilo/openai/o1-pro` | 200K | | | | | | $150 | $600 |
|
|
@@ -380,7 +379,7 @@ for await (const chunk of stream) {
|
|
|
380
379
|
| `kilo/tencent/hy-mt2-1.8b` | 8K | | | | | | $0.04 | $0.18 |
|
|
381
380
|
| `kilo/tencent/hy-mt2-30b-a3b` | 8K | | | | | | $0.07 | $0.29 |
|
|
382
381
|
| `kilo/tencent/hy-mt2-7b` | 8K | | | | | | $0.07 | $0.29 |
|
|
383
|
-
| `kilo/tencent/hy3` | 262K | | | | | | $0.
|
|
382
|
+
| `kilo/tencent/hy3` | 262K | | | | | | $0.13 | $0.53 |
|
|
384
383
|
| `kilo/tencent/hy3-preview` | 262K | | | | | | $0.18 | $0.60 |
|
|
385
384
|
| `kilo/tencent/hy4-preview` | 1.0M | | | | | | $0.83 | $3 |
|
|
386
385
|
| `kilo/thedrummer/cydonia-24b-v4.1` | 131K | | | | | | $0.30 | $0.50 |
|
|
@@ -416,6 +415,7 @@ for await (const chunk of stream) {
|
|
|
416
415
|
| `kilo/z-ai/glm-5.2:free` | 33K | | | | | | — | — |
|
|
417
416
|
| `kilo/z-ai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
418
417
|
| `kilo/z-ai/glm-5.3-flash` | 1.0M | | | | | | $0.15 | $0.50 |
|
|
418
|
+
| `kilo/z-ai/glm-5.3-flashx` | 1.0M | | | | | | $0.37 | $1 |
|
|
419
419
|
| `kilo/z-ai/glm-5v-turbo` | 203K | | | | | | $1 | $4 |
|
|
420
420
|
|
|
421
421
|
Model availability, capabilities, context windows, and pricing are sourced from [models.dev](https://models.dev) and may change.
|
|
@@ -241,6 +241,7 @@ for await (const chunk of stream) {
|
|
|
241
241
|
| `nano-gpt/google/gemma-4-26b-a4b-it:thinking` | 262K | | | | | | $0.13 | $0.40 |
|
|
242
242
|
| `nano-gpt/google/gemma-4-31b-it` | 262K | | | | | | $0.10 | $0.45 |
|
|
243
243
|
| `nano-gpt/google/gemma-4-31b-it:thinking` | 262K | | | | | | $0.10 | $0.35 |
|
|
244
|
+
| `nano-gpt/google/gemma4-31b-splituntied` | 262K | | | | | | $0.10 | $0.30 |
|
|
244
245
|
| `nano-gpt/Gryphe/MythoMax-L2-13b` | 4K | | | | | | $0.10 | $0.10 |
|
|
245
246
|
| `nano-gpt/hermes-high` | 1.0M | | | | | | $1 | $3 |
|
|
246
247
|
| `nano-gpt/hermes-low` | 1.0M | | | | | | $1 | $3 |
|
|
@@ -494,7 +495,6 @@ for await (const chunk of stream) {
|
|
|
494
495
|
| `nano-gpt/sarvam-105b` | 131K | | | | | | $0.05 | $0.21 |
|
|
495
496
|
| `nano-gpt/shisa-ai/shisa-v2-llama3.3-70b` | 128K | | | | | | $0.50 | $0.50 |
|
|
496
497
|
| `nano-gpt/shisa-ai/shisa-v2.1-llama3.3-70b` | 33K | | | | | | $0.50 | $0.50 |
|
|
497
|
-
| `nano-gpt/slowburn/gemma4-31b-splituntied` | 262K | | | | | | $0.10 | $0.30 |
|
|
498
498
|
| `nano-gpt/soob3123/amoral-gemma3-27B-v2` | 33K | | | | | | $0.30 | $0.30 |
|
|
499
499
|
| `nano-gpt/soob3123/GrayLine-Qwen3-8B` | 33K | | | | | | $0.30 | $0.30 |
|
|
500
500
|
| `nano-gpt/soob3123/Veiled-Calla-12B` | 33K | | | | | | $0.30 | $0.30 |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# OpenCode Zen
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 107 OpenCode Zen models through Mastra's model router. Authentication is handled automatically using the `OPENCODE_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [OpenCode Zen documentation](https://opencode.ai/docs/zen).
|
|
10
10
|
|
|
@@ -54,6 +54,7 @@ for await (const chunk of stream) {
|
|
|
54
54
|
| `opencode/deepseek-v4-flash` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
55
55
|
| `opencode/deepseek-v4-flash-vision-exp` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
56
56
|
| `opencode/deepseek-v4-pro` | 1.0M | | | | | | $2 | $4 |
|
|
57
|
+
| `opencode/deepseek-v4.1-flash` | 1.0M | | | | | | $0.30 | $1 |
|
|
57
58
|
| `opencode/gemini-3-flash` | 1.0M | | | | | | $0.50 | $3 |
|
|
58
59
|
| `opencode/gemini-3.1-pro` | 1.0M | | | | | | $2 | $12 |
|
|
59
60
|
| `opencode/gemini-3.5-flash` | 1.0M | | | | | | $2 | $9 |
|
|
@@ -90,6 +91,8 @@ for await (const chunk of stream) {
|
|
|
90
91
|
| `opencode/grok-4.5` | 500K | | | | | | $2 | $6 |
|
|
91
92
|
| `opencode/grok-4.6` | 500K | | | | | | $2 | $6 |
|
|
92
93
|
| `opencode/grok-build-0.1` | 256K | | | | | | $1 | $2 |
|
|
94
|
+
| `opencode/jev-1.13` | 64K | | | | | | $0.04 | — |
|
|
95
|
+
| `opencode/jev-1.13-free` | 64K | | | | | | — | — |
|
|
93
96
|
| `opencode/jev-latest` | 64K | | | | | | $0.04 | — |
|
|
94
97
|
| `opencode/kimi-k2.5` | 262K | | | | | | $0.60 | $3 |
|
|
95
98
|
| `opencode/kimi-k2.6` | 262K | | | | | | $0.95 | $4 |
|
|
@@ -108,6 +111,7 @@ for await (const chunk of stream) {
|
|
|
108
111
|
| `opencode/nemotron-3.5-lightning-free` | 262K | | | | | | — | — |
|
|
109
112
|
| `opencode/qwen3.5-plus` | 262K | | | | | | $0.20 | $1 |
|
|
110
113
|
| `opencode/qwen3.6-plus` | 262K | | | | | | $0.50 | $3 |
|
|
114
|
+
| `opencode/qwen3.8-flash` | 1.0M | | | | | | $0.15 | $0.47 |
|
|
111
115
|
|
|
112
116
|
Model availability, capabilities, context windows, and pricing are sourced from [models.dev](https://models.dev) and may change.
|
|
113
117
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Z.AI
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 17 Z.AI models through Mastra's model router. Authentication is handled automatically using the `ZHIPU_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Z.AI documentation](https://docs.z.ai/guides/overview/pricing).
|
|
10
10
|
|
|
@@ -52,7 +52,8 @@ for await (const chunk of stream) {
|
|
|
52
52
|
| `zai/glm-5.1` | 200K | | | | | | $1 | $4 |
|
|
53
53
|
| `zai/glm-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
54
54
|
| `zai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
55
|
-
| `zai/glm-5.3-flash` | 1.0M | | | | | | $0.
|
|
55
|
+
| `zai/glm-5.3-flash` | 1.0M | | | | | | $0.15 | $0.50 |
|
|
56
|
+
| `zai/glm-5.3-flashx` | 1.0M | | | | | | $0.37 | $1 |
|
|
56
57
|
| `zai/glm-5v-turbo` | 200K | | | | | | $1 | $4 |
|
|
57
58
|
|
|
58
59
|
Model availability, capabilities, context windows, and pricing are sourced from [models.dev](https://models.dev) and may change.
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Zhipu AI
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 16 Zhipu AI models through Mastra's model router. Authentication is handled automatically using the `ZHIPU_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Zhipu AI documentation](https://docs.z.ai/guides/overview/pricing).
|
|
10
10
|
|
|
@@ -51,7 +51,8 @@ for await (const chunk of stream) {
|
|
|
51
51
|
| `zhipuai/glm-5.1` | 200K | | | | | | $1 | $4 |
|
|
52
52
|
| `zhipuai/glm-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
53
53
|
| `zhipuai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
54
|
-
| `zhipuai/glm-5.3-flash` | 1.0M | | | | | | $0.
|
|
54
|
+
| `zhipuai/glm-5.3-flash` | 1.0M | | | | | | $0.15 | $0.50 |
|
|
55
|
+
| `zhipuai/glm-5.3-flashx` | 1.0M | | | | | | $0.37 | $1 |
|
|
55
56
|
| `zhipuai/glm-5v-turbo` | 200K | | | | | | $5 | $22 |
|
|
56
57
|
|
|
57
58
|
Model availability, capabilities, context windows, and pricing are sourced from [models.dev](https://models.dev) and may change.
|
|
@@ -168,7 +168,7 @@ const result = await agent.generate('message for agent')
|
|
|
168
168
|
|
|
169
169
|
**options.modelSettings.frequencyPenalty** (`number`): Penalty for token frequency (-2 to 2). Reduces repetition of frequent tokens.
|
|
170
170
|
|
|
171
|
-
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
171
|
+
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. It must be an object whose configured values are positive, finite numbers of milliseconds. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
172
172
|
|
|
173
173
|
**options.modelSettings.stopSequences** (`string[]`): Stop sequences. If set, the model will stop generating text when one of the stop sequences is generated.
|
|
174
174
|
|
|
@@ -102,7 +102,7 @@ await agent.network(`
|
|
|
102
102
|
|
|
103
103
|
**options.modelSettings.frequencyPenalty** (`number`): Penalty for token frequency (-2 to 2). Reduces repetition of frequent tokens.
|
|
104
104
|
|
|
105
|
-
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
105
|
+
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. It must be an object whose configured values are positive, finite numbers of milliseconds. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
106
106
|
|
|
107
107
|
**options.modelSettings.stopSequences** (`string[]`): Stop sequences. If set, the model will stop generating text when one of the stop sequences is generated.
|
|
108
108
|
|
|
@@ -162,7 +162,7 @@ const stream = await agent.stream('message for agent')
|
|
|
162
162
|
|
|
163
163
|
**options.modelSettings.frequencyPenalty** (`number`): Penalty for token frequency (-2 to 2). Reduces repetition of frequent tokens.
|
|
164
164
|
|
|
165
|
-
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
165
|
+
**options.modelSettings.timeout** (`object`): Time-based execution budget for the run. It must be an object whose configured values are positive, finite numbers of milliseconds. Accepts totalMs, the maximum duration of the entire agent run across every loop iteration, tool call and retry, and stepMs, the maximum duration of a single model call including the time spent consuming its stream. Exceeding either budget fails with a MastraTimeoutError. A totalMs timeout ends the run and does not try fallback models, because it is a hard deadline for the whole run. A stepMs timeout is not retried against the same model but does advance to the next entry in models when fallback models are configured. Also accepts firstChunkMs, which only applies to streaming calls and is the maximum time the model may take to emit its first content-bearing chunk (text, reasoning, tool call, file or source; stream-start and metadata chunks do not count). A firstChunkMs timeout fails with timeoutType: 'firstChunk', behaves like stepMs for fallback, and is reset for each provider retry attempt. Nested timeout keys are merged across call-time and per-model settings.
|
|
166
166
|
|
|
167
167
|
**options.modelSettings.stopSequences** (`string[]`): Stop sequences. If set, the model will stop generating text when one of the stop sequences is generated.
|
|
168
168
|
|
package/package.json
CHANGED
|
@@ -1,6 +1,6 @@
|
|
|
1
1
|
{
|
|
2
2
|
"name": "@mastra/mcp-docs-server",
|
|
3
|
-
"version": "1.2.27-alpha.
|
|
3
|
+
"version": "1.2.27-alpha.17",
|
|
4
4
|
"description": "MCP server for accessing Mastra.ai documentation, changelogs, and news.",
|
|
5
5
|
"type": "module",
|
|
6
6
|
"main": "dist/index.js",
|
|
@@ -27,7 +27,7 @@
|
|
|
27
27
|
"@modelcontextprotocol/sdk": "^1.27.1",
|
|
28
28
|
"local-pkg": "^1.1.2",
|
|
29
29
|
"zod": "^4.6.4",
|
|
30
|
-
"@mastra/core": "1.68.0-alpha.
|
|
30
|
+
"@mastra/core": "1.68.0-alpha.8"
|
|
31
31
|
},
|
|
32
32
|
"devDependencies": {
|
|
33
33
|
"@hono/node-server": "^2.0.0",
|
|
@@ -42,9 +42,9 @@
|
|
|
42
42
|
"tsx": "^4.23.1",
|
|
43
43
|
"typescript": "^7.0.2",
|
|
44
44
|
"vitest": "4.1.11",
|
|
45
|
-
"@internal/lint": "0.0.133",
|
|
46
45
|
"@internal/types-builder": "0.0.108",
|
|
47
|
-
"@mastra/core": "1.68.0-alpha.
|
|
46
|
+
"@mastra/core": "1.68.0-alpha.8",
|
|
47
|
+
"@internal/lint": "0.0.133"
|
|
48
48
|
},
|
|
49
49
|
"homepage": "https://mastra.ai",
|
|
50
50
|
"repository": {
|