@mastra/mcp-docs-server 1.2.17-alpha.11 → 1.2.17-alpha.13

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -19,6 +19,7 @@ List of required environment variables for each model provider and gateway suppo
19
19
  | [Alibaba Token Plan](https://mastra.ai/models/providers/alibaba-token-plan) | `alibaba-token-plan/*` | `ALIBABA_TOKEN_PLAN_API_KEY` |
20
20
  | [Alibaba Token Plan (China)](https://mastra.ai/models/providers/alibaba-token-plan-cn) | `alibaba-token-plan-cn/*` | `ALIBABA_TOKEN_PLAN_API_KEY` |
21
21
  | [Ambient](https://mastra.ai/models/providers/ambient) | `ambient/*` | `AMBIENT_API_KEY` |
22
+ | [AMD](https://mastra.ai/models/providers/amd) | `amd/*` | `AMD_API_KEY` |
22
23
  | [Anthropic](https://mastra.ai/models/providers/anthropic) | `anthropic/*` | `ANTHROPIC_API_KEY` |
23
24
  | [AnyAPI](https://mastra.ai/models/providers/anyapi) | `anyapi/*` | `ANYAPI_API_KEY` |
24
25
  | [Atomic Chat](https://mastra.ai/models/providers/atomic-chat) | `atomic-chat/*` | `ATOMIC_CHAT_API_KEY` |
@@ -127,9 +128,11 @@ List of required environment variables for each model provider and gateway suppo
127
128
  | [Regolo AI](https://mastra.ai/models/providers/regolo-ai) | `regolo-ai/*` | `REGOLO_API_KEY` |
128
129
  | [Requesty](https://mastra.ai/models/providers/requesty) | `requesty/*` | `REQUESTY_API_KEY` |
129
130
  | [routing.run](https://mastra.ai/models/providers/routing-run) | `routing-run/*` | `ROUTING_RUN_API_KEY` |
131
+ | [RunInfra](https://mastra.ai/models/providers/runinfra) | `runinfra/*` | `RUNINFRA_GATEWAY_KEY` |
130
132
  | [Sakana AI](https://mastra.ai/models/providers/sakana) | `sakana/*` | `SAKANA_API_KEY` |
131
133
  | [Sarvam AI](https://mastra.ai/models/providers/sarvam) | `sarvam/*` | `SARVAM_API_KEY` |
132
134
  | [Scaleway](https://mastra.ai/models/providers/scaleway) | `scaleway/*` | `SCALEWAY_API_KEY` |
135
+ | [SCNet Token Plan](https://mastra.ai/models/providers/scnet-token-plan) | `scnet-token-plan/*` | `SCNET_API_KEY` |
133
136
  | [SCX.ai](https://mastra.ai/models/providers/scx) | `scx/*` | `SCX_API_KEY` |
134
137
  | [SiliconFlow](https://mastra.ai/models/providers/siliconflow) | `siliconflow/*` | `SILICONFLOW_API_KEY` |
135
138
  | [SiliconFlow (China)](https://mastra.ai/models/providers/siliconflow-cn) | `siliconflow-cn/*` | `SILICONFLOW_CN_API_KEY` |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![OpenRouter logo](https://models.dev/logos/openrouter.svg)OpenRouter
4
4
 
5
- OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 351 models through Mastra's model router.
5
+ OpenRouter aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 352 models through Mastra's model router.
6
6
 
7
7
  Learn more in the [OpenRouter documentation](https://openrouter.ai/models).
8
8
 
@@ -386,4 +386,5 @@ ANTHROPIC_API_KEY=ant-...
386
386
  | `z-ai/glm-5-turbo` |
387
387
  | `z-ai/glm-5.1` |
388
388
  | `z-ai/glm-5.2` |
389
+ | `z-ai/glm-5.2:free` |
389
390
  | `z-ai/glm-5v-turbo` |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # Model Providers
4
4
 
5
- Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to 5962 models from 172 providers through a single API.
5
+ Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to 6043 models from 175 providers through a single API.
6
6
 
7
7
  ## Features
8
8
 
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Alibaba Token Plan (China) logo](https://models.dev/logos/alibaba-token-plan-cn.svg)Alibaba Token Plan (China)
4
4
 
5
- Access 24 Alibaba Token Plan (China) models through Mastra's model router. Authentication is handled automatically using the `ALIBABA_TOKEN_PLAN_API_KEY` environment variable.
5
+ Access 25 Alibaba Token Plan (China) models through Mastra's model router. Authentication is handled automatically using the `ALIBABA_TOKEN_PLAN_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Alibaba Token Plan (China) documentation](https://www.alibabacloud.com/help/zh/model-studio/token-plan-overview).
8
8
 
@@ -40,6 +40,7 @@ for await (const chunk of stream) {
40
40
  | `alibaba-token-plan-cn/deepseek-v4-flash` | 1.0M | | | | | | — | — |
41
41
  | `alibaba-token-plan-cn/deepseek-v4-flash-0731` | 1.0M | | | | | | — | — |
42
42
  | `alibaba-token-plan-cn/deepseek-v4-pro` | 1.0M | | | | | | — | — |
43
+ | `alibaba-token-plan-cn/deepseek-v4-pro-0813` | 1.0M | | | | | | — | — |
43
44
  | `alibaba-token-plan-cn/glm-5` | 203K | | | | | | — | — |
44
45
  | `alibaba-token-plan-cn/glm-5.1` | 203K | | | | | | — | — |
45
46
  | `alibaba-token-plan-cn/glm-5.2` | 1.0M | | | | | | — | — |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Alibaba Token Plan logo](https://models.dev/logos/alibaba-token-plan.svg)Alibaba Token Plan
4
4
 
5
- Access 24 Alibaba Token Plan models through Mastra's model router. Authentication is handled automatically using the `ALIBABA_TOKEN_PLAN_API_KEY` environment variable.
5
+ Access 25 Alibaba Token Plan models through Mastra's model router. Authentication is handled automatically using the `ALIBABA_TOKEN_PLAN_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Alibaba Token Plan documentation](https://www.alibabacloud.com/help/en/model-studio/token-plan-overview).
8
8
 
@@ -40,6 +40,7 @@ for await (const chunk of stream) {
40
40
  | `alibaba-token-plan/deepseek-v4-flash` | 1.0M | | | | | | — | — |
41
41
  | `alibaba-token-plan/deepseek-v4-flash-0731` | 1.0M | | | | | | — | — |
42
42
  | `alibaba-token-plan/deepseek-v4-pro` | 1.0M | | | | | | — | — |
43
+ | `alibaba-token-plan/deepseek-v4-pro-0813` | 1.0M | | | | | | — | — |
43
44
  | `alibaba-token-plan/glm-5` | 203K | | | | | | — | — |
44
45
  | `alibaba-token-plan/glm-5.1` | 203K | | | | | | — | — |
45
46
  | `alibaba-token-plan/glm-5.2` | 1.0M | | | | | | — | — |
@@ -0,0 +1,73 @@
1
+ > Discover all available pages from the documentation index: https://mastra.ai/llms.txt
2
+
3
+ # ![AMD logo](https://models.dev/logos/amd.svg)AMD
4
+
5
+ Access 1 AMD model through Mastra's model router. Authentication is handled automatically using the `AMD_API_KEY` environment variable.
6
+
7
+ Learn more in the [AMD documentation](https://developer.amd.com.cn/radeon/tokenfactory).
8
+
9
+ ```bash
10
+ AMD_API_KEY=your-api-key
11
+ ```
12
+
13
+ ```typescript
14
+ import { Agent } from "@mastra/core/agent";
15
+
16
+ const agent = new Agent({
17
+ id: "my-agent",
18
+ name: "My Agent",
19
+ instructions: "You are a helpful assistant",
20
+ model: "amd/DeepSeek-V4-Flash"
21
+ });
22
+
23
+ // Generate a response
24
+ const response = await agent.generate("Hello!");
25
+
26
+ // Stream a response
27
+ const stream = await agent.stream("Tell me a story");
28
+ for await (const chunk of stream) {
29
+ console.log(chunk);
30
+ }
31
+ ```
32
+
33
+ > **Note:** Mastra uses the OpenAI-compatible `/chat/completions` endpoint. Some provider-specific features may not be available. Check the [AMD documentation](https://developer.amd.com.cn/radeon/tokenfactory) for details.
34
+
35
+ ## Models
36
+
37
+ | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
+ | ----------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
+ | `amd/DeepSeek-V4-Flash` | 1.0M | | | | | | $0.14 | $0.28 |
40
+
41
+ ## Advanced configuration
42
+
43
+ ### Custom headers
44
+
45
+ ```typescript
46
+ const agent = new Agent({
47
+ id: "custom-agent",
48
+ name: "custom-agent",
49
+ model: {
50
+ url: "https://developer.amd.com.cn/radeon/api/v1",
51
+ id: "amd/DeepSeek-V4-Flash",
52
+ apiKey: process.env.AMD_API_KEY,
53
+ headers: {
54
+ "X-Custom-Header": "value"
55
+ }
56
+ }
57
+ });
58
+ ```
59
+
60
+ ### Dynamic model selection
61
+
62
+ ```typescript
63
+ const agent = new Agent({
64
+ id: "dynamic-agent",
65
+ name: "Dynamic Agent",
66
+ model: ({ requestContext }) => {
67
+ const useAdvanced = requestContext.task === "complex";
68
+ return useAdvanced
69
+ ? "amd/DeepSeek-V4-Flash"
70
+ : "amd/DeepSeek-V4-Flash";
71
+ }
72
+ });
73
+ ```
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Chutes logo](https://models.dev/logos/chutes.svg)Chutes
4
4
 
5
- Access 13 Chutes models through Mastra's model router. Authentication is handled automatically using the `CHUTES_API_KEY` environment variable.
5
+ Access 14 Chutes models through Mastra's model router. Authentication is handled automatically using the `CHUTES_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Chutes documentation](https://llm.chutes.ai).
8
8
 
@@ -46,6 +46,7 @@ for await (const chunk of stream) {
46
46
  | `chutes/Qwen/Qwen3-32B-TEE` | 41K | | | | | | $0.10 | $0.42 |
47
47
  | `chutes/Qwen/Qwen3.5-397B-A17B-TEE` | 262K | | | | | | $0.45 | $3 |
48
48
  | `chutes/Qwen/Qwen3.6-27B-TEE` | 262K | | | | | | $0.30 | $2 |
49
+ | `chutes/Qwen/Qwen3.8-27B-TEE` | 262K | | | | | | $0.40 | $3 |
49
50
  | `chutes/unsloth/Mistral-Nemo-Instruct-2407-TEE` | 131K | | | | | | $0.02 | $0.10 |
50
51
  | `chutes/zai-org/GLM-5.1-TEE` | 203K | | | | | | $0.98 | $3 |
51
52
  | `chutes/zai-org/GLM-5.2-TEE` | 1.0M | | | | | | $1 | $4 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![CrossModel logo](https://models.dev/logos/crossmodel.svg)CrossModel
4
4
 
5
- Access 50 CrossModel models through Mastra's model router. Authentication is handled automatically using the `CROSSMODEL_API_KEY` environment variable.
5
+ Access 51 CrossModel models through Mastra's model router. Authentication is handled automatically using the `CROSSMODEL_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [CrossModel documentation](https://www.crossmodel.ai/docs).
8
8
 
@@ -53,6 +53,7 @@ for await (const chunk of stream) {
53
53
  | `crossmodel/gemini/gemini-3.5-flash` | 1.0M | | | | | | $2 | $9 |
54
54
  | `crossmodel/gemini/gemini-3.5-flash-lite` | 1.0M | | | | | | $0.30 | $3 |
55
55
  | `crossmodel/gemini/gemini-3.6-flash` | 1.0M | | | | | | $2 | $8 |
56
+ | `crossmodel/gemini/gemini-3.7-flash` | 1.0M | | | | | | $0.75 | $4 |
56
57
  | `crossmodel/minimax/minimax-m2.7` | 205K | | | | | | $0.33 | $1 |
57
58
  | `crossmodel/minimax/minimax-m3` | 1.0M | | | | | | $0.33 | $1 |
58
59
  | `crossmodel/moonshot/kimi-k2.5` | 262K | | | | | | $0.62 | $3 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Deep Infra logo](https://models.dev/logos/deepinfra.svg)Deep Infra
4
4
 
5
- Access 56 Deep Infra models through Mastra's model router. Authentication is handled automatically using the `DEEPINFRA_API_KEY` environment variable.
5
+ Access 58 Deep Infra models through Mastra's model router. Authentication is handled automatically using the `DEEPINFRA_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Deep Infra documentation](https://deepinfra.com/models).
8
8
 
@@ -45,6 +45,7 @@ for await (const chunk of stream) {
45
45
  | `deepinfra/deepseek-ai/DeepSeek-V4-Flash` | 1.0M | | | | | | $0.09 | $0.18 |
46
46
  | `deepinfra/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.08 | $0.18 |
47
47
  | `deepinfra/deepseek-ai/DeepSeek-V4-Pro` | 1.0M | | | | | | $1 | $3 |
48
+ | `deepinfra/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $1 | $3 |
48
49
  | `deepinfra/google/gemma-4-26B-A4B-it` | 262K | | | | | | $0.07 | $0.34 |
49
50
  | `deepinfra/google/gemma-4-31B-it` | 262K | | | | | | $0.13 | $0.38 |
50
51
  | `deepinfra/google/gemma-4-E4B-it` | 131K | | | | | | $0.02 | $0.10 |
@@ -74,6 +75,7 @@ for await (const chunk of stream) {
74
75
  | `deepinfra/Qwen/Qwen3.6-27B` | 262K | | | | | | $0.32 | $3 |
75
76
  | `deepinfra/Qwen/Qwen3.6-35B-A3B` | 262K | | | | | | $0.10 | $0.95 |
76
77
  | `deepinfra/Qwen/Qwen3.7-Max` | 256K | | | | | | $3 | $8 |
78
+ | `deepinfra/Qwen/Qwen3.8-2.4T-A95B` | 262K | | | | | | $2 | $6 |
77
79
  | `deepinfra/Qwen/Qwen3.8-Max` | 256K | | | | | | $2 | $5 |
78
80
  | `deepinfra/stepfun-ai/Step-3.7-Flash` | 262K | | | | | | $0.20 | $1 |
79
81
  | `deepinfra/tencent/Hy3` | 262K | | | | | | $0.14 | $0.58 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![DigitalOcean logo](https://models.dev/logos/digitalocean.svg)DigitalOcean
4
4
 
5
- Access 91 DigitalOcean models through Mastra's model router. Authentication is handled automatically using the `DIGITALOCEAN_ACCESS_TOKEN` environment variable.
5
+ Access 92 DigitalOcean models through Mastra's model router. Authentication is handled automatically using the `DIGITALOCEAN_ACCESS_TOKEN` environment variable.
6
6
 
7
7
  Learn more in the [DigitalOcean documentation](https://docs.digitalocean.com/products/gradient-ai-platform/details/models/).
8
8
 
@@ -61,6 +61,7 @@ for await (const chunk of stream) {
61
61
  | `digitalocean/deepseek-v3` | 164K | | | | | | — | — |
62
62
  | `digitalocean/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.08 | $0.25 |
63
63
  | `digitalocean/deepseek-v4-pro` | 1.0M | | | | | | $0.87 | $2 |
64
+ | `digitalocean/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
64
65
  | `digitalocean/e5-large-v2` | 512 | | | | | | $0.02 | — |
65
66
  | `digitalocean/fal-ai/elevenlabs/tts/multilingual-v2` | — | | | | | | — | — |
66
67
  | `digitalocean/fal-ai/fast-sdxl` | — | | | | | | — | — |
@@ -69,7 +70,7 @@ for await (const chunk of stream) {
69
70
  | `digitalocean/gemma-4-31B-it` | 256K | | | | | | $0.18 | $0.50 |
70
71
  | `digitalocean/glm-5` | 64K | | | | | | $0.75 | $2 |
71
72
  | `digitalocean/glm-5.1` | 164K | | | | | | $0.97 | $4 |
72
- | `digitalocean/glm-5.2` | 262K | | | | | | $0.63 | $2 |
73
+ | `digitalocean/glm-5.2` | 262K | | | | | | $0.70 | $2 |
73
74
  | `digitalocean/gte-large-en-v1.5` | 8K | | | | | | $0.09 | — |
74
75
  | `digitalocean/kimi-k2.5` | 262K | | | | | | $0.38 | $2 |
75
76
  | `digitalocean/kimi-k2.6` | 262K | | | | | | $0.76 | $3 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Eden AI logo](https://models.dev/logos/edenai.svg)Eden AI
4
4
 
5
- Access 220 Eden AI models through Mastra's model router. Authentication is handled automatically using the `EDENAI_API_KEY` environment variable.
5
+ Access 226 Eden AI models through Mastra's model router. Authentication is handled automatically using the `EDENAI_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Eden AI documentation](https://docs.edenai.co).
8
8
 
@@ -60,6 +60,8 @@ for await (const chunk of stream) {
60
60
  | `edenai/azure/gpt-5.2-codex` | 272K | | | | | | $2 | $14 |
61
61
  | `edenai/cerebras/gpt-oss-120b` | 131K | | | | | | $0.35 | $0.75 |
62
62
  | `edenai/cloudflare/@cf/aisingapore/gemma-sea-lion-v4-27b-it` | 128K | | | | | | $0.35 | $0.56 |
63
+ | `edenai/cloudflare/@cf/deepseek-ai/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.44 | $1 |
64
+ | `edenai/cloudflare/@cf/deepseek-ai/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
63
65
  | `edenai/cloudflare/@cf/meta/llama-guard-3-8b` | 131K | | | | | | $0.48 | $0.03 |
64
66
  | `edenai/cloudflare/@cf/openai/gpt-oss-120b` | 128K | | | | | | $0.35 | $0.75 |
65
67
  | `edenai/cloudflare/@cf/openai/gpt-oss-20b` | 128K | | | | | | $0.20 | $0.30 |
@@ -77,6 +79,7 @@ for await (const chunk of stream) {
77
79
  | `edenai/deepinfra/deepseek-ai/DeepSeek-V3` | 164K | | | | | | $0.32 | $0.89 |
78
80
  | `edenai/deepinfra/deepseek-ai/DeepSeek-V3-0324` | 164K | | | | | | $0.24 | $0.90 |
79
81
  | `edenai/deepinfra/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.08 | $0.18 |
82
+ | `edenai/deepinfra/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $1 | $3 |
80
83
  | `edenai/deepinfra/meta-llama/Llama-3.2-11B-Vision-Instruct` | 131K | | | | | | $0.34 | $0.34 |
81
84
  | `edenai/deepinfra/meta-llama/Llama-3.3-70B-Instruct` | 131K | | | | | | $0.10 | $0.32 |
82
85
  | `edenai/deepinfra/meta-llama/Llama-Guard-3-8B` | 131K | | | | | | $0.06 | $0.06 |
@@ -194,6 +197,7 @@ for await (const chunk of stream) {
194
197
  | `edenai/ovhcloud/gpt-oss-120b` | 131K | | | | | | $0.09 | $0.47 |
195
198
  | `edenai/ovhcloud/gpt-oss-20b` | 131K | | | | | | $0.05 | $0.18 |
196
199
  | `edenai/perplexityai/sonar` | 127K | | | | | | $1 | $1 |
200
+ | `edenai/perplexityai/sonar-deep-research` | 128K | | | | | | $2 | $8 |
197
201
  | `edenai/perplexityai/sonar-pro` | 200K | | | | | | $3 | $15 |
198
202
  | `edenai/perplexityai/sonar-reasoning-pro` | 128K | | | | | | $2 | $8 |
199
203
  | `edenai/qwen/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.20 | $0.40 |
@@ -209,12 +213,14 @@ for await (const chunk of stream) {
209
213
  | `edenai/qwen/qwen3-max` | 262K | | | | | | $1 | $6 |
210
214
  | `edenai/qwen/qwen3-next-80b-a3b-instruct` | 131K | | | | | | $0.15 | $1 |
211
215
  | `edenai/qwen/qwen3-next-80b-a3b-thinking` | 131K | | | | | | $0.15 | $1 |
216
+ | `edenai/qwen/qwen3.8-2.4t-a95b` | 1.0M | | | | | | $2 | $6 |
212
217
  | `edenai/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
213
218
  | `edenai/qwen/qwq-plus` | 131K | | | | | | $0.80 | $2 |
214
- | `edenai/scaleway/deepseek-v4-flash-0731` | 256K | | | | | | $0.46 | $0.92 |
219
+ | `edenai/scaleway/deepseek-v4-flash-0731` | 256K | | | | | | $0.46 | $0.93 |
215
220
  | `edenai/scaleway/gpt-oss-120b` | 128K | | | | | | $0.17 | $0.69 |
216
221
  | `edenai/scaleway/llama-3.3-70b-instruct` | 128K | | | | | | $1 | $1 |
217
222
  | `edenai/together_ai/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
223
+ | `edenai/together_ai/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $1 | $4 |
218
224
  | `edenai/together_ai/meta-models/Muse-Glimmer-30B` | 131K | | | | | | $0.35 | $2 |
219
225
  | `edenai/together_ai/nvidia/nemotron-3-ultra-550b-a55b` | 512K | | | | | | $0.60 | $4 |
220
226
  | `edenai/together_ai/openai/gpt-oss-120b` | 131K | | | | | | $0.15 | $0.60 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![EmpirioLabs AI logo](https://models.dev/logos/empiriolabs.svg)EmpirioLabs AI
4
4
 
5
- Access 44 EmpirioLabs AI models through Mastra's model router. Authentication is handled automatically using the `EMPIRIOLABS_API_KEY` environment variable.
5
+ Access 53 EmpirioLabs AI models through Mastra's model router. Authentication is handled automatically using the `EMPIRIOLABS_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [EmpirioLabs AI documentation](https://docs.empiriolabs.ai).
8
8
 
@@ -41,6 +41,8 @@ for await (const chunk of stream) {
41
41
  | `empiriolabs/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
42
42
  | `empiriolabs/deepseek-v4-pro` | 1.0M | | | | | | $2 | $3 |
43
43
  | `empiriolabs/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
44
+ | `empiriolabs/fugu-ultra-v1-0` | 1.0M | | | | | | $8 | $45 |
45
+ | `empiriolabs/fugu-ultra-v1-1` | 1.0M | | | | | | $5 | $30 |
44
46
  | `empiriolabs/gemma-4-26b-a4b` | 262K | | | | | | $0.05 | $0.29 |
45
47
  | `empiriolabs/glm-4-5-flash` | 200K | | | | | | — | — |
46
48
  | `empiriolabs/glm-4-7-flash` | 200K | | | | | | — | — |
@@ -57,7 +59,9 @@ for await (const chunk of stream) {
57
59
  | `empiriolabs/minimax-m3` | 1.0M | | | | | | $0.23 | $0.90 |
58
60
  | `empiriolabs/mistral-medium-3` | 130K | | | | | | — | — |
59
61
  | `empiriolabs/mistral-small-4` | 256K | | | | | | $0.15 | $0.60 |
62
+ | `empiriolabs/muse-glimmer-30b` | 131K | | | | | | $0.20 | $0.80 |
60
63
  | `empiriolabs/muse-spark-1-1` | 1.0M | | | | | | $1 | $4 |
64
+ | `empiriolabs/muse-spark-1-2` | 1.0M | | | | | | $1 | $4 |
61
65
  | `empiriolabs/qwen3-5-122b-a10b` | 256K | | | | | | $0.12 | $0.92 |
62
66
  | `empiriolabs/qwen3-5-27b` | 256K | | | | | | $0.09 | $0.69 |
63
67
  | `empiriolabs/qwen3-5-35b-a3b` | 256K | | | | | | $0.06 | $0.46 |
@@ -74,9 +78,14 @@ for await (const chunk of stream) {
74
78
  | `empiriolabs/qwen3-7-flash` | 1.0M | | | | | | $0.03 | $0.13 |
75
79
  | `empiriolabs/qwen3-7-max` | 1.0M | | | | | | $3 | $8 |
76
80
  | `empiriolabs/qwen3-7-plus` | 1.0M | | | | | | $0.40 | $2 |
81
+ | `empiriolabs/qwen3-8-27b` | 262K | | | | | | $0.17 | $0.50 |
77
82
  | `empiriolabs/qwen3-8-max` | 1.0M | | | | | | $2 | $6 |
78
83
  | `empiriolabs/qwen3-max` | 256K | | | | | | $1 | $6 |
79
84
  | `empiriolabs/seed-2-0-code` | 256K | | | | | | $0.40 | $2 |
85
+ | `empiriolabs/seed-2-0-lite` | 256K | | | | | | $0.31 | $3 |
86
+ | `empiriolabs/seed-2-0-mini` | 256K | | | | | | $0.12 | $0.50 |
87
+ | `empiriolabs/seed-2-0-pro` | 256K | | | | | | $0.63 | $4 |
88
+ | `empiriolabs/seed-2-1-turbo` | 256K | | | | | | $0.63 | $3 |
80
89
  | `empiriolabs/step-3-5-flash` | 256K | | | | | | $0.10 | $0.30 |
81
90
  | `empiriolabs/step-3-5-flash-2603` | 256K | | | | | | $0.10 | $0.30 |
82
91
  | `empiriolabs/step-3-7-flash` | 256K | | | | | | $0.20 | $1 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Hugging Face logo](https://models.dev/logos/huggingface.svg)Hugging Face
4
4
 
5
- Access 65 Hugging Face models through Mastra's model router. Authentication is handled automatically using the `HF_TOKEN` environment variable.
5
+ Access 66 Hugging Face models through Mastra's model router. Authentication is handled automatically using the `HF_TOKEN` environment variable.
6
6
 
7
7
  Learn more in the [Hugging Face documentation](https://huggingface.co).
8
8
 
@@ -84,6 +84,7 @@ for await (const chunk of stream) {
84
84
  | `huggingface/Qwen/Qwen3.5-9B` | 262K | | | | | | $0.17 | $0.25 |
85
85
  | `huggingface/Qwen/Qwen3.6-27B` | 262K | | | | | | $0.47 | $3 |
86
86
  | `huggingface/Qwen/Qwen3.6-35B-A3B` | 262K | | | | | | $0.15 | $0.95 |
87
+ | `huggingface/Qwen/Qwen3.8-2.4T-A95B` | 262K | | | | | | $3 | $6 |
87
88
  | `huggingface/stepfun-ai/Step-3.5-Flash` | 262K | | | | | | $0.10 | $0.30 |
88
89
  | `huggingface/stepfun-ai/Step-3.7-Flash` | 262K | | | | | | $0.20 | $1 |
89
90
  | `huggingface/tencent/Hy3` | 262K | | | | | | $0.14 | $0.58 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Charm Hyper logo](https://models.dev/logos/hyper.svg)Charm Hyper
4
4
 
5
- Access 25 Charm Hyper models through Mastra's model router. Authentication is handled automatically using the `HYPER_API_KEY` environment variable.
5
+ Access 26 Charm Hyper models through Mastra's model router. Authentication is handled automatically using the `HYPER_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Charm Hyper documentation](https://hyper.charm.land).
8
8
 
@@ -41,13 +41,14 @@ for await (const chunk of stream) {
41
41
  | `hyper/deepseek-v4-pro` | 1.0M | | | | | | $2 | $5 |
42
42
  | `hyper/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
43
43
  | `hyper/gemma-4-26b-a4b-it` | 256K | | | | | | $0.12 | $0.42 |
44
+ | `hyper/glm-5` | 203K | | | | | | $0.76 | $2 |
44
45
  | `hyper/glm-5.1` | 203K | | | | | | $2 | $5 |
45
46
  | `hyper/glm-5.2` | 1.0M | | | | | | $1 | $4 |
46
- | `hyper/gpt-oss-120b` | 131K | | | | | | $0.18 | $0.71 |
47
+ | `hyper/gpt-oss-120b` | 131K | | | | | | $0.16 | $0.65 |
47
48
  | `hyper/kimi-k2.5` | 262K | | | | | | $0.56 | $3 |
48
49
  | `hyper/kimi-k2.6` | 262K | | | | | | $0.95 | $4 |
49
50
  | `hyper/kimi-k2.7-code` | 256K | | | | | | $0.95 | $4 |
50
- | `hyper/kimi-k3` | 1.0M | | | | | | $3 | $16 |
51
+ | `hyper/kimi-k3` | 1.0M | | | | | | $3 | $15 |
51
52
  | `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.61 | $0.84 |
52
53
  | `hyper/llama-4-maverick-17b-128e-instruct-fp8` | 430K | | | | | | $0.27 | $0.90 |
53
54
  | `hyper/minimax-m2.7` | 262K | | | | | | $0.43 | $2 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![InferX logo](https://models.dev/logos/inferx.svg)InferX
4
4
 
5
- Access 6 InferX models through Mastra's model router. Authentication is handled automatically using the `INFERX_API_KEY` environment variable.
5
+ Access 12 InferX models through Mastra's model router. Authentication is handled automatically using the `INFERX_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [InferX documentation](https://model.inferx.net/endpoints).
8
8
 
@@ -17,7 +17,7 @@ const agent = new Agent({
17
17
  id: "my-agent",
18
18
  name: "My Agent",
19
19
  instructions: "You are a helpful assistant",
20
- model: "inferx/google/gemma-4-31b-it-fp8"
20
+ model: "inferx/Agents-A1"
21
21
  });
22
22
 
23
23
  // Generate a response
@@ -34,14 +34,20 @@ for await (const chunk of stream) {
34
34
 
35
35
  ## Models
36
36
 
37
- | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
- | ------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
- | `inferx/google/gemma-4-31b-it-fp8` | 262K | | | | | | — | — |
40
- | `inferx/qwen/qwen3.5-122b-a10b-nvfp4` | 256K | | | | | | — | — |
41
- | `inferx/qwen/qwen3.6-27b-fp8` | 262K | | | | | | — | — |
42
- | `inferx/qwen/qwen3.6-35b-a3b-fp8` | 262K | | | | | | — | — |
43
- | `inferx/qwen3-coder-next-fp8` | 256K | | | | | | — | — |
44
- | `inferx/qwen3-coder-next-fp8-1m` | 1.0M | | | | | | — | — |
37
+ | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
+ | ----------------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
+ | `inferx/Agents-A1` | 262K | | | | | | — | — |
40
+ | `inferx/deepseek-v4-flash` | 1.0M | | | | | | — | — |
41
+ | `inferx/Devstral-2-123B-Instruct-2512-int4-AutoRound` | 128K | | | | | | — | — |
42
+ | `inferx/gemma-4-31B-it-fp8` | 262K | | | | | | — | — |
43
+ | `inferx/mimo-v25` | 1.0M | | | | | | — | — |
44
+ | `inferx/Ornith-1.0-35B-FP8` | 262K | | | | | | — | — |
45
+ | `inferx/Qwen3-Coder-Next-FP8` | 256K | | | | | | — | — |
46
+ | `inferx/Qwen3-Coder-Next-FP8-no-thinking` | 260K | | | | | | — | — |
47
+ | `inferx/Qwen3-Embedding-8B` | 33K | | | | | | — | — |
48
+ | `inferx/Qwen3.6-27B-FP8` | 262K | | | | | | — | — |
49
+ | `inferx/Qwen3.6-35B-A3B-FP8` | 262K | | | | | | — | — |
50
+ | `inferx/Qwen3.6-35B-A3B-fp8-no-thinking` | 262K | | | | | | — | — |
45
51
 
46
52
  ## Advanced configuration
47
53
 
@@ -53,7 +59,7 @@ const agent = new Agent({
53
59
  name: "custom-agent",
54
60
  model: {
55
61
  url: "https://model.inferx.net/endpoints/v1",
56
- id: "inferx/google/gemma-4-31b-it-fp8",
62
+ id: "inferx/Agents-A1",
57
63
  apiKey: process.env.INFERX_API_KEY,
58
64
  headers: {
59
65
  "X-Custom-Header": "value"
@@ -71,8 +77,8 @@ const agent = new Agent({
71
77
  model: ({ requestContext }) => {
72
78
  const useAdvanced = requestContext.task === "complex";
73
79
  return useAdvanced
74
- ? "inferx/qwen3-coder-next-fp8-1m"
75
- : "inferx/google/gemma-4-31b-it-fp8";
80
+ ? "inferx/mimo-v25"
81
+ : "inferx/Agents-A1";
76
82
  }
77
83
  });
78
84
  ```
@@ -40,7 +40,7 @@ for await (const chunk of stream) {
40
40
  | `kilo/~anthropic/claude-haiku-latest` | 200K | | | | | | $1 | $5 |
41
41
  | `kilo/~anthropic/claude-opus-latest` | 1.0M | | | | | | $5 | $25 |
42
42
  | `kilo/~anthropic/claude-sonnet-latest` | 1.0M | | | | | | $2 | $10 |
43
- | `kilo/~deepseek/deepseek-v4-flash-latest` | 262K | | | | | | $0.07 | $0.14 |
43
+ | `kilo/~deepseek/deepseek-v4-flash-latest` | 262K | | | | | | $0.06 | $0.12 |
44
44
  | `kilo/~google/gemini-flash-latest` | 1.0M | | | | | | $0.38 | $2 |
45
45
  | `kilo/~google/gemini-pro-latest` | 1.0M | | | | | | $2 | $12 |
46
46
  | `kilo/~moonshotai/kimi-latest` | 1.0M | | | | | | $3 | $14 |
@@ -298,12 +298,12 @@ for await (const chunk of stream) {
298
298
  | `kilo/qwen/qwen-plus` | 1.0M | | | | | | $0.26 | $0.78 |
299
299
  | `kilo/qwen/qwen-plus-2025-07-28` | 1.0M | | | | | | $0.26 | $0.78 |
300
300
  | `kilo/qwen/qwen-plus-2025-07-28:thinking` | 1.0M | | | | | | $0.26 | $0.78 |
301
- | `kilo/qwen/qwen2.5-vl-72b-instruct` | 32K | | | | | | $0.25 | $0.75 |
301
+ | `kilo/qwen/qwen2.5-vl-72b-instruct` | 128K | | | | | | $0.80 | $1 |
302
302
  | `kilo/qwen/qwen3-14b` | 41K | | | | | | $0.23 | $0.91 |
303
303
  | `kilo/qwen/qwen3-235b-a22b` | 131K | | | | | | $0.46 | $2 |
304
304
  | `kilo/qwen/qwen3-235b-a22b-2507` | 262K | | | | | | $0.15 | $0.60 |
305
305
  | `kilo/qwen/qwen3-235b-a22b-thinking-2507` | 131K | | | | | | $0.23 | $2 |
306
- | `kilo/qwen/qwen3-30b-a3b` | 41K | | | | | | $0.13 | $0.52 |
306
+ | `kilo/qwen/qwen3-30b-a3b` | 131K | | | | | | $0.13 | $0.52 |
307
307
  | `kilo/qwen/qwen3-30b-a3b-instruct-2507` | 128K | | | | | | $0.13 | $0.52 |
308
308
  | `kilo/qwen/qwen3-30b-a3b-thinking-2507` | 82K | | | | | | $0.20 | $2 |
309
309
  | `kilo/qwen/qwen3-32b` | 41K | | | | | | $0.08 | $0.28 |
@@ -333,7 +333,7 @@ for await (const chunk of stream) {
333
333
  | `kilo/qwen/qwen3.5-plus-02-15` | 1.0M | | | | | | $0.26 | $2 |
334
334
  | `kilo/qwen/qwen3.5-plus-20260420` | 1.0M | | | | | | $0.30 | $2 |
335
335
  | `kilo/qwen/qwen3.6-27b` | 262K | | | | | | $0.45 | $3 |
336
- | `kilo/qwen/qwen3.6-35b-a3b` | 262K | | | | | | $0.15 | $1 |
336
+ | `kilo/qwen/qwen3.6-35b-a3b` | 262K | | | | | | $0.14 | $1 |
337
337
  | `kilo/qwen/qwen3.6-flash` | 1.0M | | | | | | $0.19 | $1 |
338
338
  | `kilo/qwen/qwen3.6-max-preview` | 262K | | | | | | $1 | $6 |
339
339
  | `kilo/qwen/qwen3.6-plus` | 1.0M | | | | | | $0.33 | $2 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![LLMTR logo](https://models.dev/logos/llmtr.svg)LLMTR
4
4
 
5
- Access 6 LLMTR models through Mastra's model router. Authentication is handled automatically using the `LLMTR_API_KEY` environment variable.
5
+ Access 32 LLMTR models through Mastra's model router. Authentication is handled automatically using the `LLMTR_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [LLMTR documentation](https://llmtr.com/docs).
8
8
 
@@ -34,14 +34,39 @@ for await (const chunk of stream) {
34
34
 
35
35
  ## Models
36
36
 
37
- | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
- | --------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
- | `llmtr/gemma-4` | 33K | | | | | | $5 | $10 |
40
- | `llmtr/magibu-11b-v8` | 8K | | | | | | | |
41
- | `llmtr/medgemma-4b` | 8K | | | | | | $3 | $5 |
42
- | `llmtr/qwen3-6-35b` | 16K | | | | | | $5 | $10 |
43
- | `llmtr/sincap` | 128K | | | | | | | |
44
- | `llmtr/trendyol-7b` | 33K | | | | | | | |
37
+ | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
+ | --------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
+ | `llmtr/gemma-4` | 131K | | | | | | $2 | $5 |
40
+ | `llmtr/google/gemini-2.5-flash-lite` | 1.0M | | | | | | $0.10 | $0.10 |
41
+ | `llmtr/magibu-11b-v8` | 8K | | | | | | $0.10 | $0.50 |
42
+ | `llmtr/meta/muse-spark-1.2-contributor` | 1.0M | | | | | | $0.10 | $0.20 |
43
+ | `llmtr/mimo/mimo-v2.5` | 1.0M | | | | | | $0.14 | $0.28 |
44
+ | `llmtr/mimo/mimo-v2.5-pro` | 1.0M | | | | | | $0.43 | $0.87 |
45
+ | `llmtr/mistral/voxtral-small-latest` | 32K | | | | | | $0.10 | $0.30 |
46
+ | `llmtr/muse-glimmer-30b-tr` | 131K | | | | | | $2 | $5 |
47
+ | `llmtr/perplexity/sonar-deep-research` | 128K | | | | | | $2 | $8 |
48
+ | `llmtr/poolside/laguna-xs-2.1` | 262K | | | | | | — | — |
49
+ | `llmtr/publicai/apertus-70b-instruct` | 66K | | | | | | $0.82 | $3 |
50
+ | `llmtr/publicai/apertus-8b-instruct` | 66K | | | | | | $0.10 | $0.20 |
51
+ | `llmtr/qwen/qwen-flash` | 1.0M | | | | | | $0.05 | $0.40 |
52
+ | `llmtr/qwen/qwen-plus` | 1.0M | | | | | | $0.40 | $1 |
53
+ | `llmtr/qwen/qwen3-coder-flash` | 1.0M | | | | | | $0.30 | $2 |
54
+ | `llmtr/qwen/qwen3-coder-plus` | 1.0M | | | | | | $1 | $5 |
55
+ | `llmtr/qwen/qwen3-max` | 256K | | | | | | $1 | $6 |
56
+ | `llmtr/qwen/qwen3-vl-plus` | 256K | | | | | | $0.20 | $2 |
57
+ | `llmtr/qwen/qwen3.5-397b-a17b` | 256K | | | | | | $0.60 | $4 |
58
+ | `llmtr/qwen/qwen3.5-plus` | 1.0M | | | | | | $0.40 | $2 |
59
+ | `llmtr/qwen/qwen3.6-flash` | 1.0M | | | | | | $0.25 | $2 |
60
+ | `llmtr/qwen/qwen3.6-plus` | 1.0M | | | | | | $0.50 | $3 |
61
+ | `llmtr/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.40 | $2 |
62
+ | `llmtr/qwen3-6-35b` | 16K | | | | | | $5 | $10 |
63
+ | `llmtr/sakana/fugu-ultra` | 1.0M | | | | | | $5 | $30 |
64
+ | `llmtr/thinkingmachines/inkling` | 262K | | | | | | $2 | $5 |
65
+ | `llmtr/thinkingmachines/inkling-small` | 262K | | | | | | $0.58 | $1 |
66
+ | `llmtr/trendyol-asure-12b` | 41K | | | | | | $0.10 | $0.50 |
67
+ | `llmtr/upstage/solar-pro2` | 66K | | | | | | $0.15 | $0.60 |
68
+ | `llmtr/upstage/solar-pro3` | 131K | | | | | | $0.15 | $0.60 |
69
+ | `llmtr/upstage/solar-pro4` | 524K | | | | | | $0.03 | $0.12 |
45
70
 
46
71
  ## Advanced configuration
47
72
 
@@ -71,7 +96,7 @@ const agent = new Agent({
71
96
  model: ({ requestContext }) => {
72
97
  const useAdvanced = requestContext.task === "complex";
73
98
  return useAdvanced
74
- ? "llmtr/trendyol-7b"
99
+ ? "llmtr/upstage/solar-pro4"
75
100
  : "llmtr/gemma-4";
76
101
  }
77
102
  });
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![NanoGPT logo](https://models.dev/logos/nano-gpt.svg)NanoGPT
4
4
 
5
- Access 599 NanoGPT models through Mastra's model router. Authentication is handled automatically using the `NANO_GPT_API_KEY` environment variable.
5
+ Access 601 NanoGPT models through Mastra's model router. Authentication is handled automatically using the `NANO_GPT_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [NanoGPT documentation](https://docs.nano-gpt.com).
8
8
 
@@ -150,20 +150,16 @@ for await (const chunk of stream) {
150
150
  | `nano-gpt/deepseek/deepseek-v3.2:thinking` | 163K | | | | | | $0.28 | $0.42 |
151
151
  | `nano-gpt/deepseek/deepseek-v4-flash` | 1.0M | | | | | | $0.14 | $0.28 |
152
152
  | `nano-gpt/deepseek/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
153
- | `nano-gpt/deepseek/deepseek-v4-flash-0731-cheaper` | 1.0M | | | | | | $0.14 | $0.28 |
154
- | `nano-gpt/deepseek/deepseek-v4-flash-0731-cheaper:thinking` | 1.0M | | | | | | $0.14 | $0.28 |
155
153
  | `nano-gpt/deepseek/deepseek-v4-flash-0731:thinking` | 1.0M | | | | | | $0.14 | $0.28 |
156
154
  | `nano-gpt/deepseek/deepseek-v4-flash-latest` | 1.0M | | | | | | $0.14 | $0.28 |
157
155
  | `nano-gpt/deepseek/deepseek-v4-flash:thinking` | 1.0M | | | | | | $0.14 | $0.28 |
158
156
  | `nano-gpt/deepseek/deepseek-v4-pro` | 1.0M | | | | | | $1 | $2 |
159
157
  | `nano-gpt/deepseek/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $3 |
160
158
  | `nano-gpt/deepseek/deepseek-v4-pro-0813:thinking` | 1.0M | | | | | | $1 | $3 |
161
- | `nano-gpt/deepseek/deepseek-v4-pro-cheaper` | 1.0M | | | | | | $0.43 | $0.87 |
162
- | `nano-gpt/deepseek/deepseek-v4-pro-cheaper:thinking` | 1.0M | | | | | | $0.43 | $0.87 |
163
159
  | `nano-gpt/deepseek/deepseek-v4-pro:thinking` | 1.0M | | | | | | $1 | $2 |
164
160
  | `nano-gpt/dmind/dmind-1-mini` | 33K | | | | | | $0.20 | $0.40 |
165
161
  | `nano-gpt/Doctor-Shotgun/MS3.2-24B-Magnum-Diamond` | 16K | | | | | | $0.49 | $0.49 |
166
- | `nano-gpt/dots-studio/dots-3-note-preview:free` | 393K | | | | | | | |
162
+ | `nano-gpt/dots-studio/dots-3-note-preview` | 393K | | | | | | $0.10 | $0.20 |
167
163
  | `nano-gpt/doubao-1.5-pro-256k` | 256K | | | | | | $0.80 | $1 |
168
164
  | `nano-gpt/doubao-1.5-pro-32k` | 32K | | | | | | $0.13 | $0.33 |
169
165
  | `nano-gpt/doubao-1.5-vision-pro-32k` | 32K | | | | | | $0.46 | $1 |
@@ -474,14 +470,17 @@ for await (const chunk of stream) {
474
470
  | `nano-gpt/qwen3-max-2026-01-23` | 256K | | | | | | $1 | $6 |
475
471
  | `nano-gpt/qwen3-vl-235b-a22b-instruct-original` | 33K | | | | | | $0.50 | $1 |
476
472
  | `nano-gpt/qwen3-vl-235b-a22b-thinking` | 33K | | | | | | $0.50 | $6 |
473
+ | `nano-gpt/qwen3.5-0.8b` | 262K | | | | | | $0.06 | $0.12 |
477
474
  | `nano-gpt/qwen3.5-122b-a10b` | 131K | | | | | | $0.44 | $3 |
478
475
  | `nano-gpt/qwen3.5-122b-a10b:thinking` | 131K | | | | | | $0.44 | $3 |
479
476
  | `nano-gpt/qwen3.5-27b` | 260K | | | | | | $0.27 | $2 |
480
477
  | `nano-gpt/Qwen3.5-27B-BlueStar-v3-Derestricted` | 262K | | | | | | $0.31 | $0.31 |
481
478
  | `nano-gpt/Qwen3.5-27B-Queen-Derestricted` | 262K | | | | | | $0.31 | $0.31 |
482
479
  | `nano-gpt/qwen3.5-27b:thinking` | 260K | | | | | | $0.27 | $2 |
480
+ | `nano-gpt/qwen3.5-2b` | 262K | | | | | | $0.08 | $0.16 |
483
481
  | `nano-gpt/qwen3.5-35b-a3b` | 260K | | | | | | $0.23 | $2 |
484
482
  | `nano-gpt/qwen3.5-35b-a3b:thinking` | 260K | | | | | | $0.23 | $2 |
483
+ | `nano-gpt/qwen3.5-4b` | 262K | | | | | | $0.10 | $0.20 |
485
484
  | `nano-gpt/qwen3.5-flash` | 992K | | | | | | $0.10 | $0.40 |
486
485
  | `nano-gpt/qwen3.5-flash:thinking` | 992K | | | | | | $0.10 | $0.40 |
487
486
  | `nano-gpt/qwen3.5-omni-flash` | 49K | | | | | | — | — |
@@ -493,6 +492,8 @@ for await (const chunk of stream) {
493
492
  | `nano-gpt/qwen3.7-max:thinking` | 1.0M | | | | | | $3 | $8 |
494
493
  | `nano-gpt/qwen3.7-plus` | 992K | | | | | | $0.40 | $2 |
495
494
  | `nano-gpt/qwen3.7-plus:thinking` | 984K | | | | | | $0.40 | $2 |
495
+ | `nano-gpt/qwen3.8-27b` | 262K | | | | | | $0.40 | $3 |
496
+ | `nano-gpt/qwen3.8-27b:thinking` | 262K | | | | | | $0.40 | $3 |
496
497
  | `nano-gpt/qwen3.8-max` | 991K | | | | | | $2 | $6 |
497
498
  | `nano-gpt/qwen3.8-max:thinking` | 991K | | | | | | $2 | $6 |
498
499
  | `nano-gpt/ReadyArt/MS3.2-The-Omega-Directive-24B-Unslop-v2.0` | 16K | | | | | | $0.50 | $0.50 |
@@ -526,9 +527,10 @@ for await (const chunk of stream) {
526
527
  | `nano-gpt/stepfun-ai/step-3.5-flash` | 256K | | | | | | $0.10 | $0.30 |
527
528
  | `nano-gpt/stepfun-ai/step-3.5-flash-2603` | 256K | | | | | | $0.10 | $0.30 |
528
529
  | `nano-gpt/stepfun/step-3.7-flash:thinking` | 262K | | | | | | $0.20 | $1 |
529
- | `nano-gpt/TEE/deepseek-v3.1` | 164K | | | | | | $1 | $3 |
530
530
  | `nano-gpt/TEE/deepseek-v3.2` | 164K | | | | | | $0.50 | $1 |
531
531
  | `nano-gpt/TEE/deepseek-v4-flash` | 1.0M | | | | | | $0.20 | $0.40 |
532
+ | `nano-gpt/TEE/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
533
+ | `nano-gpt/TEE/deepseek-v4-pro-0813:thinking` | 1.0M | | | | | | $1 | $4 |
532
534
  | `nano-gpt/TEE/gemma-3-27b-it` | 131K | | | | | | $0.20 | $0.80 |
533
535
  | `nano-gpt/TEE/gemma-4-26b-a4b-uncensored` | 66K | | | | | | $0.15 | $0.70 |
534
536
  | `nano-gpt/TEE/gemma-4-31b-it` | 262K | | | | | | $0.15 | $0.46 |
@@ -69,9 +69,9 @@ for await (const chunk of stream) {
69
69
  | `ofox/bailian/qwen3.7-plus` | 1.0M | | | | | | $0.40 | $2 |
70
70
  | `ofox/bailian/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
71
71
  | `ofox/deepseek/deepseek-v3.2` | 128K | | | | | | $0.29 | $0.43 |
72
- | `ofox/deepseek/deepseek-v4-flash` | 1.0M | | | | | | $0.14 | $0.28 |
73
- | `ofox/deepseek/deepseek-v4-pro` | 1.0M | | | | | | $0.45 | $0.88 |
74
- | `ofox/deepseek/deepseek-v4-pro-0813` | 1.0M | | | | | | $0.45 | $0.88 |
72
+ | `ofox/deepseek/deepseek-v4-flash` | 1.0M | | | | | | $0.44 | $1 |
73
+ | `ofox/deepseek/deepseek-v4-pro` | 1.0M | | | | | | $1 | $4 |
74
+ | `ofox/deepseek/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
75
75
  | `ofox/google/gemini-2.5-flash` | 1.0M | | | | | | $0.30 | $3 |
76
76
  | `ofox/google/gemini-2.5-flash-lite` | 1.0M | | | | | | $0.10 | $0.40 |
77
77
  | `ofox/google/gemini-2.5-pro` | 1.0M | | | | | | $1 | $10 |
@@ -36,8 +36,8 @@ for await (const chunk of stream) {
36
36
 
37
37
  | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
38
  | ------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
- | `opencode-go/deepseek-v4-flash` | 1.0M | | | | | | $0.07 | $0.14 |
40
- | `opencode-go/deepseek-v4-pro` | 1.0M | | | | | | $0.43 | $0.87 |
39
+ | `opencode-go/deepseek-v4-flash` | 1.0M | | | | | | $0.22 | $0.66 |
40
+ | `opencode-go/deepseek-v4-pro` | 1.0M | | | | | | $0.66 | $2 |
41
41
  | `opencode-go/glm-5.1` | 203K | | | | | | $1 | $4 |
42
42
  | `opencode-go/glm-5.2` | 1.0M | | | | | | $1 | $4 |
43
43
  | `opencode-go/glm-5.3` | 1.0M | | | | | | $1 | $4 |
@@ -0,0 +1,76 @@
1
+ > Discover all available pages from the documentation index: https://mastra.ai/llms.txt
2
+
3
+ # ![RunInfra logo](https://models.dev/logos/runinfra.svg)RunInfra
4
+
5
+ Access 4 RunInfra models through Mastra's model router. Authentication is handled automatically using the `RUNINFRA_GATEWAY_KEY` environment variable.
6
+
7
+ Learn more in the [RunInfra documentation](https://runinfra.ai/docs).
8
+
9
+ ```bash
10
+ RUNINFRA_GATEWAY_KEY=your-api-key
11
+ ```
12
+
13
+ ```typescript
14
+ import { Agent } from "@mastra/core/agent";
15
+
16
+ const agent = new Agent({
17
+ id: "my-agent",
18
+ name: "My Agent",
19
+ instructions: "You are a helpful assistant",
20
+ model: "runinfra/Inferact/Qwen3.8-2.4T-A95B-NVFP4"
21
+ });
22
+
23
+ // Generate a response
24
+ const response = await agent.generate("Hello!");
25
+
26
+ // Stream a response
27
+ const stream = await agent.stream("Tell me a story");
28
+ for await (const chunk of stream) {
29
+ console.log(chunk);
30
+ }
31
+ ```
32
+
33
+ > **Note:** Mastra uses the OpenAI-compatible `/chat/completions` endpoint. Some provider-specific features may not be available. Check the [RunInfra documentation](https://runinfra.ai/docs) for details.
34
+
35
+ ## Models
36
+
37
+ | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
+ | ------------------------------------------------------------ | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
+ | `runinfra/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.13 | $0.27 |
40
+ | `runinfra/Inferact/Qwen3.8-2.4T-A95B-NVFP4` | 262K | | | | | | $2 | $6 |
41
+ | `runinfra/nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-BF16` | 262K | | | | | | $0.05 | $0.15 |
42
+ | `runinfra/Qwen/Qwen3.8-27B` | 262K | | | | | | $0.10 | $0.40 |
43
+
44
+ ## Advanced configuration
45
+
46
+ ### Custom headers
47
+
48
+ ```typescript
49
+ const agent = new Agent({
50
+ id: "custom-agent",
51
+ name: "custom-agent",
52
+ model: {
53
+ url: "https://api.runinfra.ai/v1",
54
+ id: "runinfra/Inferact/Qwen3.8-2.4T-A95B-NVFP4",
55
+ apiKey: process.env.RUNINFRA_GATEWAY_KEY,
56
+ headers: {
57
+ "X-Custom-Header": "value"
58
+ }
59
+ }
60
+ });
61
+ ```
62
+
63
+ ### Dynamic model selection
64
+
65
+ ```typescript
66
+ const agent = new Agent({
67
+ id: "dynamic-agent",
68
+ name: "Dynamic Agent",
69
+ model: ({ requestContext }) => {
70
+ const useAdvanced = requestContext.task === "complex";
71
+ return useAdvanced
72
+ ? "runinfra/nvidia/NVIDIA-Nemotron-3.5-Lightning-30B-A3B-BF16"
73
+ : "runinfra/Inferact/Qwen3.8-2.4T-A95B-NVFP4";
74
+ }
75
+ });
76
+ ```
@@ -0,0 +1,85 @@
1
+ > Discover all available pages from the documentation index: https://mastra.ai/llms.txt
2
+
3
+ # ![SCNet Token Plan logo](https://models.dev/logos/scnet-token-plan.svg)SCNet Token Plan
4
+
5
+ Access 13 SCNet Token Plan models through Mastra's model router. Authentication is handled automatically using the `SCNET_API_KEY` environment variable.
6
+
7
+ Learn more in the [SCNet Token Plan documentation](https://www.scnet.cn/ac/openapi/doc/2.0/moduleapi/plans/token-plan.html).
8
+
9
+ ```bash
10
+ SCNET_API_KEY=your-api-key
11
+ ```
12
+
13
+ ```typescript
14
+ import { Agent } from "@mastra/core/agent";
15
+
16
+ const agent = new Agent({
17
+ id: "my-agent",
18
+ name: "My Agent",
19
+ instructions: "You are a helpful assistant",
20
+ model: "scnet-token-plan/DeepSeek-V3.2"
21
+ });
22
+
23
+ // Generate a response
24
+ const response = await agent.generate("Hello!");
25
+
26
+ // Stream a response
27
+ const stream = await agent.stream("Tell me a story");
28
+ for await (const chunk of stream) {
29
+ console.log(chunk);
30
+ }
31
+ ```
32
+
33
+ > **Note:** Mastra uses the OpenAI-compatible `/chat/completions` endpoint. Some provider-specific features may not be available. Check the [SCNet Token Plan documentation](https://www.scnet.cn/ac/openapi/doc/2.0/moduleapi/plans/token-plan.html) for details.
34
+
35
+ ## Models
36
+
37
+ | Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
38
+ | ------------------------------------ | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
+ | `scnet-token-plan/DeepSeek-V3.2` | 128K | | | | | | — | — |
40
+ | `scnet-token-plan/DeepSeek-V4-Flash` | 1.0M | | | | | | — | — |
41
+ | `scnet-token-plan/GLM-5` | 205K | | | | | | — | — |
42
+ | `scnet-token-plan/GLM-5.1` | 200K | | | | | | — | — |
43
+ | `scnet-token-plan/GLM-5.2` | 1.0M | | | | | | — | — |
44
+ | `scnet-token-plan/Kimi-K2.5` | 262K | | | | | | — | — |
45
+ | `scnet-token-plan/Kimi-K2.6` | 262K | | | | | | — | — |
46
+ | `scnet-token-plan/Kimi-K2.7-Code` | 262K | | | | | | — | — |
47
+ | `scnet-token-plan/Kimi-K3` | 1.0M | | | | | | — | — |
48
+ | `scnet-token-plan/MiMo-V2.5-Pro` | 1.0M | | | | | | — | — |
49
+ | `scnet-token-plan/MiniMax-M2.5` | 205K | | | | | | — | — |
50
+ | `scnet-token-plan/MiniMax-M2.7` | 205K | | | | | | — | — |
51
+ | `scnet-token-plan/MiniMax-M3` | 512K | | | | | | — | — |
52
+
53
+ ## Advanced configuration
54
+
55
+ ### Custom headers
56
+
57
+ ```typescript
58
+ const agent = new Agent({
59
+ id: "custom-agent",
60
+ name: "custom-agent",
61
+ model: {
62
+ url: "https://api.scnet.cn/api/llm/v1",
63
+ id: "scnet-token-plan/DeepSeek-V3.2",
64
+ apiKey: process.env.SCNET_API_KEY,
65
+ headers: {
66
+ "X-Custom-Header": "value"
67
+ }
68
+ }
69
+ });
70
+ ```
71
+
72
+ ### Dynamic model selection
73
+
74
+ ```typescript
75
+ const agent = new Agent({
76
+ id: "dynamic-agent",
77
+ name: "Dynamic Agent",
78
+ model: ({ requestContext }) => {
79
+ const useAdvanced = requestContext.task === "complex";
80
+ return useAdvanced
81
+ ? "scnet-token-plan/MiniMax-M3"
82
+ : "scnet-token-plan/DeepSeek-V3.2";
83
+ }
84
+ });
85
+ ```
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Together AI logo](https://models.dev/logos/togetherai.svg)Together AI
4
4
 
5
- Access 35 Together AI models through Mastra's model router. Authentication is handled automatically using the `TOGETHER_API_KEY` environment variable.
5
+ Access 36 Together AI models through Mastra's model router. Authentication is handled automatically using the `TOGETHER_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Together AI documentation](https://docs.together.ai/docs/serverless-models).
8
8
 
@@ -37,6 +37,7 @@ for await (const chunk of stream) {
37
37
  | `togetherai/deepcogito/cogito-v2-1-671b` | 164K | | | | | | $1 | $1 |
38
38
  | `togetherai/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
39
39
  | `togetherai/deepseek-ai/DeepSeek-V4-Pro` | 512K | | | | | | $2 | $3 |
40
+ | `togetherai/deepseek-ai/DeepSeek-V4-Pro-0813` | 1.0M | | | | | | $1 | $4 |
40
41
  | `togetherai/google/gemma-3n-E4B-it` | 33K | | | | | | $0.06 | $0.12 |
41
42
  | `togetherai/google/gemma-4-31B-it` | 262K | | | | | | $0.39 | $0.97 |
42
43
  | `togetherai/LiquidAI/LFM2-24B-A2B` | 33K | | | | | | $0.03 | $0.12 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Umans AI Coding Plan logo](https://models.dev/logos/umans-ai-coding-plan.svg)Umans AI Coding Plan
4
4
 
5
- Access 8 Umans AI Coding Plan models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_CODING_PLAN_API_KEY` environment variable.
5
+ Access 9 Umans AI Coding Plan models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_CODING_PLAN_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Umans AI Coding Plan documentation](https://app.umans.ai/offers/code/docs).
8
8
 
@@ -38,6 +38,7 @@ for await (const chunk of stream) {
38
38
  | --------------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
39
  | `umans-ai-coding-plan/umans-coder` | 262K | | | | | | — | — |
40
40
  | `umans-ai-coding-plan/umans-deepseek-v4-flash-0731` | 1.0M | | | | | | — | — |
41
+ | `umans-ai-coding-plan/umans-deepseek-v4-pro-0813` | 1.0M | | | | | | — | — |
41
42
  | `umans-ai-coding-plan/umans-flash` | 262K | | | | | | — | — |
42
43
  | `umans-ai-coding-plan/umans-glm-5.1` | 205K | | | | | | — | — |
43
44
  | `umans-ai-coding-plan/umans-glm-5.2` | 406K | | | | | | — | — |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![Umans AI logo](https://models.dev/logos/umans-ai.svg)Umans AI
4
4
 
5
- Access 7 Umans AI models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_API_KEY` environment variable.
5
+ Access 8 Umans AI models through Mastra's model router. Authentication is handled automatically using the `UMANS_AI_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [Umans AI documentation](https://app.umans.ai/offers/code/docs/orgs).
8
8
 
@@ -38,6 +38,7 @@ for await (const chunk of stream) {
38
38
  | --------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
39
39
  | `umans-ai/umans-coder` | 262K | | | | | | $0.95 | $4 |
40
40
  | `umans-ai/umans-deepseek-v4-flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
41
+ | `umans-ai/umans-deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
41
42
  | `umans-ai/umans-flash` | 262K | | | | | | $0.15 | $1 |
42
43
  | `umans-ai/umans-glm-5.1` | 205K | | | | | | $1 | $4 |
43
44
  | `umans-ai/umans-glm-5.2` | 406K | | | | | | $1 | $4 |
@@ -2,7 +2,7 @@
2
2
 
3
3
  # ![xAI logo](https://models.dev/logos/xai.svg)xAI
4
4
 
5
- Access 11 xAI models through Mastra's model router. Authentication is handled automatically using the `XAI_API_KEY` environment variable.
5
+ Access 12 xAI models through Mastra's model router. Authentication is handled automatically using the `XAI_API_KEY` environment variable.
6
6
 
7
7
  Learn more in the [xAI documentation](https://docs.x.ai/docs/models).
8
8
 
@@ -42,6 +42,7 @@ for await (const chunk of stream) {
42
42
  | `xai/grok-4.6` | 500K | | | | | | $2 | $6 |
43
43
  | `xai/grok-build-0.1` | 256K | | | | | | $1 | $2 |
44
44
  | `xai/grok-imagine-image` | 8K | | | | | | — | — |
45
+ | `xai/grok-imagine-image-2.0` | 8K | | | | | | — | — |
45
46
  | `xai/grok-imagine-image-quality` | 8K | | | | | | — | — |
46
47
  | `xai/grok-imagine-video` | 1K | | | | | | — | — |
47
48
  | `xai/grok-imagine-video-1.5` | 1K | | | | | | — | — |
@@ -24,6 +24,7 @@ Direct access to individual AI model providers. Each provider offers unique mode
24
24
  - [Alibaba Token Plan](https://mastra.ai/models/providers/alibaba-token-plan)
25
25
  - [Alibaba Token Plan (China)](https://mastra.ai/models/providers/alibaba-token-plan-cn)
26
26
  - [Ambient](https://mastra.ai/models/providers/ambient)
27
+ - [AMD](https://mastra.ai/models/providers/amd)
27
28
  - [AnyAPI](https://mastra.ai/models/providers/anyapi)
28
29
  - [Atomic Chat](https://mastra.ai/models/providers/atomic-chat)
29
30
  - [Auriko](https://mastra.ai/models/providers/auriko)
@@ -126,9 +127,11 @@ Direct access to individual AI model providers. Each provider offers unique mode
126
127
  - [Regolo AI](https://mastra.ai/models/providers/regolo-ai)
127
128
  - [Requesty](https://mastra.ai/models/providers/requesty)
128
129
  - [routing.run](https://mastra.ai/models/providers/routing-run)
130
+ - [RunInfra](https://mastra.ai/models/providers/runinfra)
129
131
  - [Sakana AI](https://mastra.ai/models/providers/sakana)
130
132
  - [Sarvam AI](https://mastra.ai/models/providers/sarvam)
131
133
  - [Scaleway](https://mastra.ai/models/providers/scaleway)
134
+ - [SCNet Token Plan](https://mastra.ai/models/providers/scnet-token-plan)
132
135
  - [SCX.ai](https://mastra.ai/models/providers/scx)
133
136
  - [SiliconFlow](https://mastra.ai/models/providers/siliconflow)
134
137
  - [SiliconFlow (China)](https://mastra.ai/models/providers/siliconflow-cn)
package/CHANGELOG.md CHANGED
@@ -1,5 +1,12 @@
1
1
  # @mastra/mcp-docs-server
2
2
 
3
+ ## 1.2.17-alpha.12
4
+
5
+ ### Patch Changes
6
+
7
+ - Updated dependencies [[`940bf5c`](https://github.com/mastra-ai/mastra/commit/940bf5ccf04f2c9ebd8a1390431733222a03b1cd)]:
8
+ - @mastra/core@1.60.0-alpha.7
9
+
3
10
  ## 1.2.17-alpha.10
4
11
 
5
12
  ### Patch Changes
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@mastra/mcp-docs-server",
3
- "version": "1.2.17-alpha.11",
3
+ "version": "1.2.17-alpha.13",
4
4
  "description": "MCP server for accessing Mastra.ai documentation, changelogs, and news.",
5
5
  "type": "module",
6
6
  "main": "dist/index.js",
@@ -28,7 +28,7 @@
28
28
  "jsdom": "^26.1.0",
29
29
  "local-pkg": "^1.1.2",
30
30
  "zod": "^4.4.3",
31
- "@mastra/core": "1.60.0-alpha.6",
31
+ "@mastra/core": "1.60.0-alpha.7",
32
32
  "@mastra/mcp": "^1.17.0-alpha.0"
33
33
  },
34
34
  "devDependencies": {
@@ -47,7 +47,7 @@
47
47
  "vitest": "4.1.10",
48
48
  "@internal/lint": "0.0.123",
49
49
  "@internal/types-builder": "0.0.98",
50
- "@mastra/core": "1.60.0-alpha.6"
50
+ "@mastra/core": "1.60.0-alpha.7"
51
51
  },
52
52
  "homepage": "https://mastra.ai",
53
53
  "repository": {