@mastra/mcp-docs-server 1.2.17-alpha.7 → 1.2.17
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/.docs/course/02-agent-tools-mcp/32-conclusion.md +1 -1
- package/.docs/docs/agents/code-mode.md +3 -3
- package/.docs/docs/agents/guardrails.md +1 -1
- package/.docs/docs/agents/{agent-approval.md → human-in-the-loop.md} +5 -5
- package/.docs/docs/agents/networks.md +3 -3
- package/.docs/docs/agents/overview.md +7 -7
- package/.docs/docs/agents/processors.md +3 -3
- package/.docs/docs/agents/{using-tools.md → tools.md} +7 -7
- package/.docs/docs/{server/auth → auth}/custom-auth-provider.md +1 -1
- package/.docs/docs/{server/auth → auth}/fga.md +27 -1
- package/.docs/docs/{server/auth.md → auth/overview.md} +4 -4
- package/.docs/docs/{server/auth → auth}/simple-auth.md +1 -1
- package/.docs/docs/{server/auth → auth}/workers.md +2 -2
- package/.docs/docs/{capabilities/channels.md → channels.md} +2 -2
- package/.docs/docs/{agents → connections}/a2a.md +2 -2
- package/.docs/docs/{agents → connections}/acp.md +18 -6
- package/.docs/docs/{mcp/overview.md → connections/mcp.md} +1 -1
- package/.docs/docs/connections/overview.md +5 -5
- package/.docs/docs/{agents → connections}/sdk-agents.md +3 -1
- package/.docs/docs/deployment/cloud-providers.md +1 -0
- package/.docs/docs/deployment/mastra-server.md +2 -2
- package/.docs/docs/deployment/overview.md +2 -1
- package/.docs/docs/deployment/sandbox.md +3 -3
- package/.docs/docs/deployment/workers.md +4 -4
- package/.docs/docs/guides/multi-agent-systems.md +7 -7
- package/.docs/docs/guides/streaming.md +1 -1
- package/.docs/docs/harness/agent-controller.md +5 -3
- package/.docs/docs/{long-running-agents → harness}/background-tasks.md +5 -5
- package/.docs/docs/{long-running-agents → harness}/durable-agents.md +18 -2
- package/.docs/docs/{long-running-agents → harness}/goals.md +6 -6
- package/.docs/docs/harness/overview.md +11 -10
- package/.docs/docs/{long-running-agents → harness}/schedules.md +5 -5
- package/.docs/docs/{long-running-agents → harness}/signal-providers.md +4 -4
- package/.docs/docs/mastra-platform/deploy.md +1 -1
- package/.docs/docs/mastra-platform/overview.md +1 -1
- package/.docs/docs/mastra-platform/server.md +1 -1
- package/.docs/docs/memory/message-history.md +1 -1
- package/.docs/docs/memory/overview.md +4 -4
- package/.docs/docs/memory/working-memory.md +1 -1
- package/.docs/docs/observability/integrations/exporters/mastra-storage.md +1 -1
- package/.docs/docs/{workspace → sandbox}/filesystem.md +2 -2
- package/.docs/docs/{workspace → sandbox}/lsp.md +3 -3
- package/.docs/docs/{workspace/sandbox.md → sandbox/overview.md} +4 -3
- package/.docs/docs/{workspace → sandbox}/search.md +2 -2
- package/.docs/docs/{workspace → sandbox}/skills.md +5 -5
- package/.docs/docs/server/custom-api-routes.md +2 -2
- package/.docs/docs/server/mastra-client.md +2 -2
- package/.docs/docs/server/{mastra-server.md → overview.md} +3 -3
- package/.docs/docs/server/pubsub.md +2 -2
- package/.docs/docs/server/server-adapters.md +4 -4
- package/.docs/docs/{agents/skills.md → skills.md} +4 -4
- package/.docs/docs/{storage/overview.md → storage.md} +2 -1
- package/.docs/docs/studio/auth.md +4 -4
- package/.docs/docs/studio/overview.md +2 -2
- package/.docs/docs/{capabilities/subagents.md → subagents.md} +35 -5
- package/.docs/docs/workflows/agents-and-tools.md +1 -1
- package/.docs/docs/workflows/control-flow.md +0 -4
- package/.docs/docs/workflows/human-in-the-loop.md +0 -4
- package/.docs/docs/workflows/overview.md +1 -1
- package/.docs/docs/workflows/scheduled-workflows.md +2 -2
- package/.docs/docs/workflows/snapshots.md +1 -1
- package/.docs/docs/workflows/suspend-and-resume.md +0 -4
- package/.docs/integrations/agentic-ui/ai-sdk-ui.md +1 -1
- package/.docs/integrations/agentic-ui/copilotkit.md +1 -1
- package/.docs/integrations/auth/google.md +2 -2
- package/.docs/integrations/auth/workos.md +1 -1
- package/.docs/integrations/browsers/agent-browser.md +2 -2
- package/.docs/integrations/browsers/browser-viewer.md +6 -6
- package/.docs/integrations/browsers/firecrawl.md +1 -1
- package/.docs/integrations/browsers/stagehand.md +2 -2
- package/.docs/integrations/channels/discord.md +2 -2
- package/.docs/integrations/channels/github.md +1 -1
- package/.docs/integrations/channels/imessage.md +4 -4
- package/.docs/integrations/channels/slack.md +5 -5
- package/.docs/integrations/channels/teams.md +2 -2
- package/.docs/integrations/channels/telegram.md +2 -2
- package/.docs/integrations/channels/whatsapp.md +2 -2
- package/.docs/integrations/databases/postgresql.md +1 -0
- package/.docs/integrations/deploy/amazon-ec2.md +2 -2
- package/.docs/integrations/deploy/aws-lambda.md +3 -3
- package/.docs/integrations/deploy/azure-app-services.md +2 -2
- package/.docs/integrations/deploy/cloudflare.md +2 -2
- package/.docs/integrations/deploy/digital-ocean.md +3 -3
- package/.docs/integrations/deploy/kubernetes.md +11 -11
- package/.docs/integrations/deploy/netlify.md +3 -3
- package/.docs/integrations/deploy/render.md +389 -0
- package/.docs/integrations/deploy/vercel.md +2 -2
- package/.docs/integrations/file-storage/amazon-s3.md +1 -1
- package/.docs/integrations/file-storage/azure-blob.md +1 -1
- package/.docs/integrations/file-storage/google-cloud-storage.md +1 -1
- package/.docs/integrations/file-storage/mesa.md +2 -2
- package/.docs/integrations/file-storage/vercel-files.md +1 -1
- package/.docs/integrations/frameworks/astro.md +6 -2
- package/.docs/integrations/frameworks/electron.md +1 -1
- package/.docs/integrations/frameworks/express.md +1 -1
- package/.docs/integrations/frameworks/hono.md +1 -1
- package/.docs/integrations/frameworks/nestjs.md +1 -1
- package/.docs/integrations/frameworks/next-js.md +6 -2
- package/.docs/integrations/frameworks/nuxt.md +1 -1
- package/.docs/integrations/frameworks/sveltekit.md +1 -1
- package/.docs/integrations/frameworks/vite-react.md +6 -2
- package/.docs/integrations/sandboxes/agentcore.md +1 -1
- package/.docs/integrations/sandboxes/apple-container.md +1 -1
- package/.docs/integrations/sandboxes/cloudflare-sandbox.md +118 -0
- package/.docs/integrations/sandboxes/daytona.md +1 -1
- package/.docs/integrations/sandboxes/docker.md +4 -3
- package/.docs/integrations/sandboxes/e2b.md +1 -1
- package/.docs/integrations/sandboxes/modal.md +1 -1
- package/.docs/integrations/sandboxes/railway.md +11 -0
- package/.docs/integrations.md +4 -0
- package/.docs/models/environment-variables.md +8 -2
- package/.docs/models/gateways/merge-gateway.md +212 -0
- package/.docs/models/gateways/openrouter.md +5 -1
- package/.docs/models/gateways/vercel.md +22 -1
- package/.docs/models/gateways.md +1 -0
- package/.docs/models/index.md +1 -1
- package/.docs/models/providers/alibaba-token-plan-cn.md +2 -1
- package/.docs/models/providers/alibaba-token-plan.md +2 -1
- package/.docs/models/providers/ambient.md +2 -2
- package/.docs/models/providers/amd.md +73 -0
- package/.docs/models/providers/arcee.md +79 -0
- package/.docs/models/providers/baseten.md +1 -1
- package/.docs/models/providers/cerebras.md +2 -3
- package/.docs/models/providers/chutes.md +2 -1
- package/.docs/models/providers/cloudflare-workers-ai.md +3 -2
- package/.docs/models/providers/cortecs.md +2 -1
- package/.docs/models/providers/crof.md +1 -1
- package/.docs/models/providers/crossmodel.md +3 -2
- package/.docs/models/providers/deepinfra.md +5 -1
- package/.docs/models/providers/digitalocean.md +3 -2
- package/.docs/models/providers/edenai.md +19 -5
- package/.docs/models/providers/empiriolabs.md +10 -1
- package/.docs/models/providers/hetzner.md +6 -8
- package/.docs/models/providers/huggingface.md +4 -1
- package/.docs/models/providers/hyper.md +9 -8
- package/.docs/models/providers/inferx.md +19 -13
- package/.docs/models/providers/jalapeno.md +89 -0
- package/.docs/models/providers/kilo.md +17 -14
- package/.docs/models/providers/kosmik.md +73 -0
- package/.docs/models/providers/llmgateway.md +3 -3
- package/.docs/models/providers/llmtr.md +35 -10
- package/.docs/models/providers/nano-gpt.md +18 -19
- package/.docs/models/providers/ofox.md +3 -3
- package/.docs/models/providers/opencode-go.md +2 -2
- package/.docs/models/providers/requesty.md +3 -3
- package/.docs/models/providers/runinfra.md +76 -0
- package/.docs/models/providers/sakana.md +3 -2
- package/.docs/models/providers/scnet-token-plan.md +85 -0
- package/.docs/models/providers/scx-ai.md +76 -0
- package/.docs/models/providers/togetherai.md +2 -1
- package/.docs/models/providers/umans-ai-coding-plan.md +2 -1
- package/.docs/models/providers/umans-ai.md +2 -1
- package/.docs/models/providers/vivgrid.md +2 -1
- package/.docs/models/providers/wandb.md +2 -1
- package/.docs/models/providers/xai.md +2 -1
- package/.docs/models/providers.md +7 -2
- package/.docs/reference/acp/acp-agent.md +2 -2
- package/.docs/reference/acp/create-acp-tool.md +1 -1
- package/.docs/reference/agent-controller/agent-controller-class.md +2 -2
- package/.docs/reference/agents/agent.md +2 -2
- package/.docs/reference/agents/channels.md +2 -2
- package/.docs/reference/agents/createSkill.md +2 -2
- package/.docs/reference/agents/durable-agent.md +1 -1
- package/.docs/reference/agents/generate.md +3 -1
- package/.docs/reference/agents/getSkill.md +1 -1
- package/.docs/reference/agents/listSkills.md +1 -1
- package/.docs/reference/agents/listSuspendedRuns.md +6 -6
- package/.docs/reference/agents/listTools.md +2 -2
- package/.docs/reference/agents/network.md +3 -1
- package/.docs/reference/ai-sdk/chat-route.md +1 -1
- package/.docs/reference/ai-sdk/handle-chat-stream.md +1 -1
- package/.docs/reference/ai-sdk/handle-network-stream.md +2 -2
- package/.docs/reference/ai-sdk/handle-workflow-stream.md +1 -1
- package/.docs/reference/ai-sdk/network-route.md +2 -2
- package/.docs/reference/ai-sdk/to-ai-sdk-messages.md +1 -1
- package/.docs/reference/ai-sdk/to-ai-sdk-stream.md +1 -1
- package/.docs/reference/ai-sdk/workflow-route.md +1 -1
- package/.docs/reference/auth/fga.md +7 -5
- package/.docs/reference/auth/jwt.md +1 -1
- package/.docs/reference/browser/agent-browser.md +2 -2
- package/.docs/reference/browser/browser-viewer.md +2 -2
- package/.docs/reference/browser/firecrawl-browser.md +1 -1
- package/.docs/reference/browser/mastra-browser.md +1 -1
- package/.docs/reference/browser/stagehand-browser.md +2 -2
- package/.docs/reference/build-with-ai.md +2 -2
- package/.docs/reference/channels/channel-provider.md +1 -1
- package/.docs/reference/channels/slack-provider.md +1 -1
- package/.docs/reference/cli/mastra.md +2 -0
- package/.docs/reference/client-js/agents.md +24 -4
- package/.docs/reference/coding-agent/create-coding-agent.md +142 -13
- package/.docs/reference/configuration.md +6 -6
- package/.docs/reference/core/getEditor.md +1 -1
- package/.docs/reference/core/getMCPServer.md +1 -1
- package/.docs/reference/core/getMCPServerById.md +1 -1
- package/.docs/reference/core/getTool.md +1 -1
- package/.docs/reference/core/getToolById.md +1 -1
- package/.docs/reference/core/listMCPServers.md +1 -1
- package/.docs/reference/core/listTools.md +1 -1
- package/.docs/reference/core/removeWorkspace.md +1 -1
- package/.docs/reference/editor/mastra-editor.md +2 -2
- package/.docs/reference/editor/prompt-blocks.md +2 -2
- package/.docs/reference/editor/tool-provider.md +108 -1
- package/.docs/reference/editor/tools.md +1 -1
- package/.docs/reference/editor/versioning.md +3 -3
- package/.docs/reference/evals/prompt-alignment.md +18 -0
- package/.docs/reference/evals/rubric.md +1 -1
- package/.docs/reference/file-based-agents/memory.md +2 -2
- package/.docs/reference/file-based-agents/server.md +3 -3
- package/.docs/reference/file-based-agents/skills.md +1 -1
- package/.docs/reference/file-based-agents/storage.md +3 -3
- package/.docs/reference/file-based-agents/subagents.md +1 -1
- package/.docs/reference/file-based-agents/workspace.md +3 -3
- package/.docs/reference/index.md +1 -0
- package/.docs/reference/manual-install.md +3 -3
- package/.docs/reference/memory/memory-class.md +1 -0
- package/.docs/reference/memory/settled.md +57 -0
- package/.docs/reference/migrations/network-to-supervisor.md +2 -2
- package/.docs/reference/processors/provider-history-compat.md +6 -5
- package/.docs/reference/processors/skill-search-processor.md +3 -1
- package/.docs/reference/processors/token-limiter-processor.md +4 -0
- package/.docs/reference/processors/tool-call-filter.md +7 -7
- package/.docs/reference/processors/tool-search-processor.md +1 -1
- package/.docs/reference/project-structure.md +1 -1
- package/.docs/reference/pubsub/lease-provider.md +3 -3
- package/.docs/reference/pubsub/redis-streams.md +1 -1
- package/.docs/reference/rag/graph-rag.md +71 -8
- package/.docs/reference/rag/retrieval.md +26 -18
- package/.docs/reference/schedules/overview.md +1 -1
- package/.docs/reference/streaming/ChunkType.md +2 -2
- package/.docs/reference/streaming/agents/stream.md +29 -4
- package/.docs/reference/streaming/agents/streamUntilIdle.md +1 -1
- package/.docs/reference/tools/ask-user-tool.md +1 -1
- package/.docs/reference/tools/create-code-mode.md +1 -1
- package/.docs/reference/tools/create-tool.md +4 -4
- package/.docs/reference/tools/mcp-client.md +2 -0
- package/.docs/reference/tools/mcp-server.md +97 -4
- package/.docs/reference/tools/submit-plan-tool.md +1 -1
- package/.docs/reference/tools/task-tools.md +2 -2
- package/.docs/reference/vectors/vectorize.md +12 -2
- package/.docs/reference/workers/overview.md +2 -2
- package/.docs/reference/workspace/local-filesystem.md +1 -1
- package/.docs/reference/workspace/local-sandbox.md +4 -3
- package/.docs/reference/workspace/platform-sandbox.md +11 -0
- package/.docs/reference/workspace/process-manager.md +20 -4
- package/.docs/reference/workspace/sandbox.md +1 -1
- package/.docs/reference/workspace/workspace-class.md +5 -5
- package/CHANGELOG.md +80 -0
- package/package.json +6 -6
- package/.docs/models/providers/merge-gateway.md +0 -265
- /package/.docs/docs/{server/auth → auth}/composite-auth.md +0 -0
- /package/.docs/docs/{server/auth → auth}/jwt.md +0 -0
- /package/.docs/docs/{browser/overview.md → browser.md} +0 -0
- /package/.docs/docs/{getting-started/develop.md → develop.md} +0 -0
- /package/.docs/docs/{long-running-agents → harness}/signals.md +0 -0
- /package/.docs/docs/{editor/overview.md → studio/editor.md} +0 -0
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# EmpirioLabs AI
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 53 EmpirioLabs AI models through Mastra's model router. Authentication is handled automatically using the `EMPIRIOLABS_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [EmpirioLabs AI documentation](https://docs.empiriolabs.ai).
|
|
8
8
|
|
|
@@ -41,6 +41,8 @@ for await (const chunk of stream) {
|
|
|
41
41
|
| `empiriolabs/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
42
42
|
| `empiriolabs/deepseek-v4-pro` | 1.0M | | | | | | $2 | $3 |
|
|
43
43
|
| `empiriolabs/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
44
|
+
| `empiriolabs/fugu-ultra-v1-0` | 1.0M | | | | | | $8 | $45 |
|
|
45
|
+
| `empiriolabs/fugu-ultra-v1-1` | 1.0M | | | | | | $5 | $30 |
|
|
44
46
|
| `empiriolabs/gemma-4-26b-a4b` | 262K | | | | | | $0.05 | $0.29 |
|
|
45
47
|
| `empiriolabs/glm-4-5-flash` | 200K | | | | | | — | — |
|
|
46
48
|
| `empiriolabs/glm-4-7-flash` | 200K | | | | | | — | — |
|
|
@@ -57,7 +59,9 @@ for await (const chunk of stream) {
|
|
|
57
59
|
| `empiriolabs/minimax-m3` | 1.0M | | | | | | $0.23 | $0.90 |
|
|
58
60
|
| `empiriolabs/mistral-medium-3` | 130K | | | | | | — | — |
|
|
59
61
|
| `empiriolabs/mistral-small-4` | 256K | | | | | | $0.15 | $0.60 |
|
|
62
|
+
| `empiriolabs/muse-glimmer-30b` | 131K | | | | | | $0.20 | $0.80 |
|
|
60
63
|
| `empiriolabs/muse-spark-1-1` | 1.0M | | | | | | $1 | $4 |
|
|
64
|
+
| `empiriolabs/muse-spark-1-2` | 1.0M | | | | | | $1 | $4 |
|
|
61
65
|
| `empiriolabs/qwen3-5-122b-a10b` | 256K | | | | | | $0.12 | $0.92 |
|
|
62
66
|
| `empiriolabs/qwen3-5-27b` | 256K | | | | | | $0.09 | $0.69 |
|
|
63
67
|
| `empiriolabs/qwen3-5-35b-a3b` | 256K | | | | | | $0.06 | $0.46 |
|
|
@@ -74,9 +78,14 @@ for await (const chunk of stream) {
|
|
|
74
78
|
| `empiriolabs/qwen3-7-flash` | 1.0M | | | | | | $0.03 | $0.13 |
|
|
75
79
|
| `empiriolabs/qwen3-7-max` | 1.0M | | | | | | $3 | $8 |
|
|
76
80
|
| `empiriolabs/qwen3-7-plus` | 1.0M | | | | | | $0.40 | $2 |
|
|
81
|
+
| `empiriolabs/qwen3-8-27b` | 262K | | | | | | $0.17 | $0.50 |
|
|
77
82
|
| `empiriolabs/qwen3-8-max` | 1.0M | | | | | | $2 | $6 |
|
|
78
83
|
| `empiriolabs/qwen3-max` | 256K | | | | | | $1 | $6 |
|
|
79
84
|
| `empiriolabs/seed-2-0-code` | 256K | | | | | | $0.40 | $2 |
|
|
85
|
+
| `empiriolabs/seed-2-0-lite` | 256K | | | | | | $0.31 | $3 |
|
|
86
|
+
| `empiriolabs/seed-2-0-mini` | 256K | | | | | | $0.12 | $0.50 |
|
|
87
|
+
| `empiriolabs/seed-2-0-pro` | 256K | | | | | | $0.63 | $4 |
|
|
88
|
+
| `empiriolabs/seed-2-1-turbo` | 256K | | | | | | $0.63 | $3 |
|
|
80
89
|
| `empiriolabs/step-3-5-flash` | 256K | | | | | | $0.10 | $0.30 |
|
|
81
90
|
| `empiriolabs/step-3-5-flash-2603` | 256K | | | | | | $0.10 | $0.30 |
|
|
82
91
|
| `empiriolabs/step-3-7-flash` | 256K | | | | | | $0.20 | $1 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Hetzner
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 2 Hetzner models through Mastra's model router. Authentication is handled automatically using the `HETZNER_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Hetzner documentation](https://experiments.hetzner.com).
|
|
8
8
|
|
|
@@ -17,7 +17,7 @@ const agent = new Agent({
|
|
|
17
17
|
id: "my-agent",
|
|
18
18
|
name: "My Agent",
|
|
19
19
|
instructions: "You are a helpful assistant",
|
|
20
|
-
model: "hetzner/
|
|
20
|
+
model: "hetzner/Qwen/Qwen3.6-35B-A3B-FP8"
|
|
21
21
|
});
|
|
22
22
|
|
|
23
23
|
// Generate a response
|
|
@@ -36,10 +36,8 @@ for await (const chunk of stream) {
|
|
|
36
36
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ---------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
|
-
| `hetzner/DeepSeek-V4-Flash-0731` | 512K | | | | | | — | — |
|
|
40
|
-
| `hetzner/GLM-5.2-NVFP4` | 512K | | | | | | — | — |
|
|
41
|
-
| `hetzner/Kimi-K2.7-Code` | 262K | | | | | | — | — |
|
|
42
39
|
| `hetzner/Qwen/Qwen3.6-35B-A3B-FP8` | 262K | | | | | | — | — |
|
|
40
|
+
| `hetzner/Qwen3.8-27B` | 262K | | | | | | — | — |
|
|
43
41
|
|
|
44
42
|
## Advanced configuration
|
|
45
43
|
|
|
@@ -51,7 +49,7 @@ const agent = new Agent({
|
|
|
51
49
|
name: "custom-agent",
|
|
52
50
|
model: {
|
|
53
51
|
url: "https://inference.hetzner.com/api/v1",
|
|
54
|
-
id: "hetzner/
|
|
52
|
+
id: "hetzner/Qwen/Qwen3.6-35B-A3B-FP8",
|
|
55
53
|
apiKey: process.env.HETZNER_API_KEY,
|
|
56
54
|
headers: {
|
|
57
55
|
"X-Custom-Header": "value"
|
|
@@ -69,8 +67,8 @@ const agent = new Agent({
|
|
|
69
67
|
model: ({ requestContext }) => {
|
|
70
68
|
const useAdvanced = requestContext.task === "complex";
|
|
71
69
|
return useAdvanced
|
|
72
|
-
? "hetzner/
|
|
73
|
-
: "hetzner/
|
|
70
|
+
? "hetzner/Qwen3.8-27B"
|
|
71
|
+
: "hetzner/Qwen/Qwen3.6-35B-A3B-FP8";
|
|
74
72
|
}
|
|
75
73
|
});
|
|
76
74
|
```
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Hugging Face
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 68 Hugging Face models through Mastra's model router. Authentication is handled automatically using the `HF_TOKEN` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Hugging Face documentation](https://huggingface.co).
|
|
8
8
|
|
|
@@ -77,6 +77,8 @@ for await (const chunk of stream) {
|
|
|
77
77
|
| `huggingface/Qwen/Qwen3-Embedding-8B` | 32K | | | | | | $0.01 | — |
|
|
78
78
|
| `huggingface/Qwen/Qwen3-Next-80B-A3B-Instruct` | 262K | | | | | | $0.25 | $1 |
|
|
79
79
|
| `huggingface/Qwen/Qwen3-Next-80B-A3B-Thinking` | 262K | | | | | | $0.30 | $2 |
|
|
80
|
+
| `huggingface/Qwen/Qwen3-VL-235B-A22B-Instruct` | 131K | | | | | | $0.30 | $2 |
|
|
81
|
+
| `huggingface/Qwen/Qwen3-VL-235B-A22B-Thinking` | 131K | | | | | | $0.98 | $4 |
|
|
80
82
|
| `huggingface/Qwen/Qwen3.5-122B-A10B` | 262K | | | | | | $0.40 | $3 |
|
|
81
83
|
| `huggingface/Qwen/Qwen3.5-27B` | 262K | | | | | | $0.30 | $2 |
|
|
82
84
|
| `huggingface/Qwen/Qwen3.5-35B-A3B` | 262K | | | | | | $0.25 | $2 |
|
|
@@ -84,6 +86,7 @@ for await (const chunk of stream) {
|
|
|
84
86
|
| `huggingface/Qwen/Qwen3.5-9B` | 262K | | | | | | $0.17 | $0.25 |
|
|
85
87
|
| `huggingface/Qwen/Qwen3.6-27B` | 262K | | | | | | $0.47 | $3 |
|
|
86
88
|
| `huggingface/Qwen/Qwen3.6-35B-A3B` | 262K | | | | | | $0.15 | $0.95 |
|
|
89
|
+
| `huggingface/Qwen/Qwen3.8-2.4T-A95B` | 262K | | | | | | $3 | $6 |
|
|
87
90
|
| `huggingface/stepfun-ai/Step-3.5-Flash` | 262K | | | | | | $0.10 | $0.30 |
|
|
88
91
|
| `huggingface/stepfun-ai/Step-3.7-Flash` | 262K | | | | | | $0.20 | $1 |
|
|
89
92
|
| `huggingface/tencent/Hy3` | 262K | | | | | | $0.14 | $0.58 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Charm Hyper
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 26 Charm Hyper models through Mastra's model router. Authentication is handled automatically using the `HYPER_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Charm Hyper documentation](https://hyper.charm.land).
|
|
8
8
|
|
|
@@ -37,20 +37,21 @@ for await (const chunk of stream) {
|
|
|
37
37
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
38
|
| ---------------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
39
|
| `hyper/deepseek-v4-flash` | 1.0M | | | | | | $0.20 | $0.40 |
|
|
40
|
-
| `hyper/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.
|
|
40
|
+
| `hyper/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.13 | $0.26 |
|
|
41
41
|
| `hyper/deepseek-v4-pro` | 1.0M | | | | | | $2 | $5 |
|
|
42
42
|
| `hyper/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
43
43
|
| `hyper/gemma-4-26b-a4b-it` | 256K | | | | | | $0.12 | $0.42 |
|
|
44
|
-
| `hyper/glm-5
|
|
44
|
+
| `hyper/glm-5` | 203K | | | | | | $0.84 | $3 |
|
|
45
|
+
| `hyper/glm-5.1` | 203K | | | | | | $1 | $4 |
|
|
45
46
|
| `hyper/glm-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
46
|
-
| `hyper/gpt-oss-120b` | 131K | | | | | | $0.
|
|
47
|
-
| `hyper/kimi-k2.5` | 262K | | | | | | $0.
|
|
47
|
+
| `hyper/gpt-oss-120b` | 131K | | | | | | $0.16 | $0.65 |
|
|
48
|
+
| `hyper/kimi-k2.5` | 262K | | | | | | $0.57 | $3 |
|
|
48
49
|
| `hyper/kimi-k2.6` | 262K | | | | | | $0.95 | $4 |
|
|
49
50
|
| `hyper/kimi-k2.7-code` | 256K | | | | | | $0.95 | $4 |
|
|
50
|
-
| `hyper/kimi-k3` | 1.0M | | | | | | $3 | $
|
|
51
|
-
| `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.
|
|
51
|
+
| `hyper/kimi-k3` | 1.0M | | | | | | $3 | $15 |
|
|
52
|
+
| `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.64 | $0.77 |
|
|
52
53
|
| `hyper/llama-4-maverick-17b-128e-instruct-fp8` | 430K | | | | | | $0.27 | $0.90 |
|
|
53
|
-
| `hyper/minimax-m2.7` | 262K | | | | | | $0.
|
|
54
|
+
| `hyper/minimax-m2.7` | 262K | | | | | | $0.47 | $2 |
|
|
54
55
|
| `hyper/minimax-m3` | 512K | | | | | | $0.33 | $1 |
|
|
55
56
|
| `hyper/qwen3-coder-480b-a35b-instruct-int4-mixed-ar` | 106K | | | | | | $0.45 | $2 |
|
|
56
57
|
| `hyper/qwen3-next-80b-a3b-instruct` | 262K | | | | | | $0.12 | $1 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# InferX
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 12 InferX models through Mastra's model router. Authentication is handled automatically using the `INFERX_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [InferX documentation](https://model.inferx.net/endpoints).
|
|
8
8
|
|
|
@@ -17,7 +17,7 @@ const agent = new Agent({
|
|
|
17
17
|
id: "my-agent",
|
|
18
18
|
name: "My Agent",
|
|
19
19
|
instructions: "You are a helpful assistant",
|
|
20
|
-
model: "inferx/
|
|
20
|
+
model: "inferx/Agents-A1"
|
|
21
21
|
});
|
|
22
22
|
|
|
23
23
|
// Generate a response
|
|
@@ -34,14 +34,20 @@ for await (const chunk of stream) {
|
|
|
34
34
|
|
|
35
35
|
## Models
|
|
36
36
|
|
|
37
|
-
| Model
|
|
38
|
-
|
|
|
39
|
-
| `inferx/
|
|
40
|
-
| `inferx/
|
|
41
|
-
| `inferx/
|
|
42
|
-
| `inferx/
|
|
43
|
-
| `inferx/
|
|
44
|
-
| `inferx/
|
|
37
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
|
+
| ----------------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
|
+
| `inferx/Agents-A1` | 262K | | | | | | — | — |
|
|
40
|
+
| `inferx/deepseek-v4-flash` | 1.0M | | | | | | — | — |
|
|
41
|
+
| `inferx/Devstral-2-123B-Instruct-2512-int4-AutoRound` | 128K | | | | | | — | — |
|
|
42
|
+
| `inferx/gemma-4-31B-it-fp8` | 262K | | | | | | — | — |
|
|
43
|
+
| `inferx/mimo-v25` | 1.0M | | | | | | — | — |
|
|
44
|
+
| `inferx/Ornith-1.0-35B-FP8` | 262K | | | | | | — | — |
|
|
45
|
+
| `inferx/Qwen3-Coder-Next-FP8` | 256K | | | | | | — | — |
|
|
46
|
+
| `inferx/Qwen3-Coder-Next-FP8-no-thinking` | 260K | | | | | | — | — |
|
|
47
|
+
| `inferx/Qwen3-Embedding-8B` | 33K | | | | | | — | — |
|
|
48
|
+
| `inferx/Qwen3.6-27B-FP8` | 262K | | | | | | — | — |
|
|
49
|
+
| `inferx/Qwen3.6-35B-A3B-FP8` | 262K | | | | | | — | — |
|
|
50
|
+
| `inferx/Qwen3.6-35B-A3B-fp8-no-thinking` | 262K | | | | | | — | — |
|
|
45
51
|
|
|
46
52
|
## Advanced configuration
|
|
47
53
|
|
|
@@ -53,7 +59,7 @@ const agent = new Agent({
|
|
|
53
59
|
name: "custom-agent",
|
|
54
60
|
model: {
|
|
55
61
|
url: "https://model.inferx.net/endpoints/v1",
|
|
56
|
-
id: "inferx/
|
|
62
|
+
id: "inferx/Agents-A1",
|
|
57
63
|
apiKey: process.env.INFERX_API_KEY,
|
|
58
64
|
headers: {
|
|
59
65
|
"X-Custom-Header": "value"
|
|
@@ -71,8 +77,8 @@ const agent = new Agent({
|
|
|
71
77
|
model: ({ requestContext }) => {
|
|
72
78
|
const useAdvanced = requestContext.task === "complex";
|
|
73
79
|
return useAdvanced
|
|
74
|
-
? "inferx/
|
|
75
|
-
: "inferx/
|
|
80
|
+
? "inferx/mimo-v25"
|
|
81
|
+
: "inferx/Agents-A1";
|
|
76
82
|
}
|
|
77
83
|
});
|
|
78
84
|
```
|
|
@@ -0,0 +1,89 @@
|
|
|
1
|
+
> Discover all available pages from the documentation index: https://mastra.ai/llms.txt
|
|
2
|
+
|
|
3
|
+
# Jalapeno Cloud
|
|
4
|
+
|
|
5
|
+
Access 17 Jalapeno Cloud models through Mastra's model router. Authentication is handled automatically using the `JALAPENO_API_KEY` environment variable.
|
|
6
|
+
|
|
7
|
+
Learn more in the [Jalapeno Cloud documentation](https://www.jalapeno-cloud.ai/docs/).
|
|
8
|
+
|
|
9
|
+
```bash
|
|
10
|
+
JALAPENO_API_KEY=your-api-key
|
|
11
|
+
```
|
|
12
|
+
|
|
13
|
+
```typescript
|
|
14
|
+
import { Agent } from "@mastra/core/agent";
|
|
15
|
+
|
|
16
|
+
const agent = new Agent({
|
|
17
|
+
id: "my-agent",
|
|
18
|
+
name: "My Agent",
|
|
19
|
+
instructions: "You are a helpful assistant",
|
|
20
|
+
model: "jalapeno/DeepSeek-V4-Flash"
|
|
21
|
+
});
|
|
22
|
+
|
|
23
|
+
// Generate a response
|
|
24
|
+
const response = await agent.generate("Hello!");
|
|
25
|
+
|
|
26
|
+
// Stream a response
|
|
27
|
+
const stream = await agent.stream("Tell me a story");
|
|
28
|
+
for await (const chunk of stream) {
|
|
29
|
+
console.log(chunk);
|
|
30
|
+
}
|
|
31
|
+
```
|
|
32
|
+
|
|
33
|
+
> **Note:** Mastra uses the OpenAI-compatible `/chat/completions` endpoint. Some provider-specific features may not be available. Check the [Jalapeno Cloud documentation](https://www.jalapeno-cloud.ai/docs/) for details.
|
|
34
|
+
|
|
35
|
+
## Models
|
|
36
|
+
|
|
37
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
|
+
| -------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
|
+
| `jalapeno/DeepSeek-V4-Flash` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
40
|
+
| `jalapeno/DeepSeek-V4-Pro` | 1.0M | | | | | | $2 | $3 |
|
|
41
|
+
| `jalapeno/GLM-5.1` | 203K | | | | | | $1 | $4 |
|
|
42
|
+
| `jalapeno/GLM-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
43
|
+
| `jalapeno/Hy3` | 203K | | | | | | $0.14 | $0.58 |
|
|
44
|
+
| `jalapeno/Kimi-K2.5` | 262K | | | | | | $0.60 | $3 |
|
|
45
|
+
| `jalapeno/Kimi-K2.7-Code` | 271K | | | | | | $0.95 | $4 |
|
|
46
|
+
| `jalapeno/Kimi-K3` | 1.0M | | | | | | $3 | $15 |
|
|
47
|
+
| `jalapeno/MiniMax-M3` | 524K | | | | | | $0.30 | $1 |
|
|
48
|
+
| `jalapeno/Qwen3-Next-80B-A3B-Instruct` | 129K | | | | | | $0.15 | $2 |
|
|
49
|
+
| `jalapeno/Qwen3-Next-80B-A3B-Thinking` | 131K | | | | | | $0.15 | $2 |
|
|
50
|
+
| `jalapeno/Qwen3-VL-235B-A22B-Instruct` | 129K | | | | | | $0.30 | $2 |
|
|
51
|
+
| `jalapeno/Qwen3-VL-235B-A22B-Thinking` | 131K | | | | | | $0.98 | $4 |
|
|
52
|
+
| `jalapeno/Qwen3.5-122B-A10B` | 262K | | | | | | $0.40 | $3 |
|
|
53
|
+
| `jalapeno/Qwen3.5-27B` | 262K | | | | | | $0.30 | $2 |
|
|
54
|
+
| `jalapeno/Qwen3.5-35B-A3B` | 262K | | | | | | $0.25 | $2 |
|
|
55
|
+
| `jalapeno/Qwen3.5-397B-A17B` | 262K | | | | | | $0.60 | $4 |
|
|
56
|
+
|
|
57
|
+
## Advanced configuration
|
|
58
|
+
|
|
59
|
+
### Custom headers
|
|
60
|
+
|
|
61
|
+
```typescript
|
|
62
|
+
const agent = new Agent({
|
|
63
|
+
id: "custom-agent",
|
|
64
|
+
name: "custom-agent",
|
|
65
|
+
model: {
|
|
66
|
+
url: "https://api.jalapeno-cloud.ai/v1",
|
|
67
|
+
id: "jalapeno/DeepSeek-V4-Flash",
|
|
68
|
+
apiKey: process.env.JALAPENO_API_KEY,
|
|
69
|
+
headers: {
|
|
70
|
+
"X-Custom-Header": "value"
|
|
71
|
+
}
|
|
72
|
+
}
|
|
73
|
+
});
|
|
74
|
+
```
|
|
75
|
+
|
|
76
|
+
### Dynamic model selection
|
|
77
|
+
|
|
78
|
+
```typescript
|
|
79
|
+
const agent = new Agent({
|
|
80
|
+
id: "dynamic-agent",
|
|
81
|
+
name: "Dynamic Agent",
|
|
82
|
+
model: ({ requestContext }) => {
|
|
83
|
+
const useAdvanced = requestContext.task === "complex";
|
|
84
|
+
return useAdvanced
|
|
85
|
+
? "jalapeno/Qwen3.5-397B-A17B"
|
|
86
|
+
: "jalapeno/DeepSeek-V4-Flash";
|
|
87
|
+
}
|
|
88
|
+
});
|
|
89
|
+
```
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# Kilo Gateway
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 360 Kilo Gateway models through Mastra's model router. Authentication is handled automatically using the `KILO_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [Kilo Gateway documentation](https://kilo.ai).
|
|
8
8
|
|
|
@@ -40,11 +40,11 @@ for await (const chunk of stream) {
|
|
|
40
40
|
| `kilo/~anthropic/claude-haiku-latest` | 200K | | | | | | $1 | $5 |
|
|
41
41
|
| `kilo/~anthropic/claude-opus-latest` | 1.0M | | | | | | $5 | $25 |
|
|
42
42
|
| `kilo/~anthropic/claude-sonnet-latest` | 1.0M | | | | | | $2 | $10 |
|
|
43
|
-
| `kilo/~deepseek/deepseek-v4-flash-latest` |
|
|
43
|
+
| `kilo/~deepseek/deepseek-v4-flash-latest` | 262K | | | | | | $0.08 | $0.15 |
|
|
44
44
|
| `kilo/~google/gemini-flash-latest` | 1.0M | | | | | | $0.38 | $2 |
|
|
45
45
|
| `kilo/~google/gemini-pro-latest` | 1.0M | | | | | | $2 | $12 |
|
|
46
|
-
| `kilo/~moonshotai/kimi-latest` |
|
|
47
|
-
| `kilo/~openai/gpt-latest` | 1.1M | | | | | | $
|
|
46
|
+
| `kilo/~moonshotai/kimi-latest` | 975K | | | | | | $3 | $13 |
|
|
47
|
+
| `kilo/~openai/gpt-latest` | 1.1M | | | | | | $3 | $15 |
|
|
48
48
|
| `kilo/~openai/gpt-mini-latest` | 400K | | | | | | $0.75 | $5 |
|
|
49
49
|
| `kilo/~x-ai/grok-latest` | 500K | | | | | | $2 | $6 |
|
|
50
50
|
| `kilo/ai21/jamba-large-1.7` | 256K | | | | | | $2 | $8 |
|
|
@@ -108,6 +108,7 @@ for await (const chunk of stream) {
|
|
|
108
108
|
| `kilo/deepseek/deepseek-v4-pro` | 1.0M | | | | | | $2 | $3 |
|
|
109
109
|
| `kilo/deepseek/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
110
110
|
| `kilo/deepseek/deepseek-v4-pro:discounted` | 1.0M | | | | | | $0.43 | $0.87 |
|
|
111
|
+
| `kilo/dots-studio/dots-3-note-preview:free` | 512K | | | | | | — | — |
|
|
111
112
|
| `kilo/google/gemini-2.5-flash` | 1.0M | | | | | | $0.30 | $3 |
|
|
112
113
|
| `kilo/google/gemini-2.5-flash-image` | 33K | | | | | | $0.30 | $3 |
|
|
113
114
|
| `kilo/google/gemini-2.5-flash-lite` | 1.0M | | | | | | $0.10 | $0.40 |
|
|
@@ -127,7 +128,7 @@ for await (const chunk of stream) {
|
|
|
127
128
|
| `kilo/google/gemini-3.5-flash` | 1.0M | | | | | | $2 | $9 |
|
|
128
129
|
| `kilo/google/gemini-3.5-flash-lite` | 1.0M | | | | | | $0.30 | $3 |
|
|
129
130
|
| `kilo/google/gemini-3.6-flash` | 1.0M | | | | | | $0.75 | $4 |
|
|
130
|
-
| `kilo/google/gemini-3.7-flash` | 1.0M | | | | | | $
|
|
131
|
+
| `kilo/google/gemini-3.7-flash` | 1.0M | | | | | | $2 | $8 |
|
|
131
132
|
| `kilo/google/gemma-2-27b-it` | 8K | | | | | | $0.65 | $0.65 |
|
|
132
133
|
| `kilo/google/gemma-3-12b-it` | 131K | | | | | | $0.05 | $0.15 |
|
|
133
134
|
| `kilo/google/gemma-3-27b-it` | 131K | | | | | | $0.08 | $0.16 |
|
|
@@ -296,13 +297,13 @@ for await (const chunk of stream) {
|
|
|
296
297
|
| `kilo/qwen/qwen-2.5-coder-32b-instruct` | 33K | | | | | | $0.66 | $1 |
|
|
297
298
|
| `kilo/qwen/qwen-plus` | 1.0M | | | | | | $0.26 | $0.78 |
|
|
298
299
|
| `kilo/qwen/qwen-plus-2025-07-28` | 1.0M | | | | | | $0.26 | $0.78 |
|
|
299
|
-
| `kilo/qwen/qwen-plus-2025-07-28:thinking` | 1.0M | | | | | | $0.
|
|
300
|
-
| `kilo/qwen/qwen2.5-vl-72b-instruct` |
|
|
300
|
+
| `kilo/qwen/qwen-plus-2025-07-28:thinking` | 1.0M | | | | | | $0.26 | $0.78 |
|
|
301
|
+
| `kilo/qwen/qwen2.5-vl-72b-instruct` | 128K | | | | | | $0.80 | $1 |
|
|
301
302
|
| `kilo/qwen/qwen3-14b` | 41K | | | | | | $0.23 | $0.91 |
|
|
302
303
|
| `kilo/qwen/qwen3-235b-a22b` | 131K | | | | | | $0.46 | $2 |
|
|
303
304
|
| `kilo/qwen/qwen3-235b-a22b-2507` | 262K | | | | | | $0.15 | $0.60 |
|
|
304
305
|
| `kilo/qwen/qwen3-235b-a22b-thinking-2507` | 131K | | | | | | $0.23 | $2 |
|
|
305
|
-
| `kilo/qwen/qwen3-30b-a3b` |
|
|
306
|
+
| `kilo/qwen/qwen3-30b-a3b` | 131K | | | | | | $0.13 | $0.52 |
|
|
306
307
|
| `kilo/qwen/qwen3-30b-a3b-instruct-2507` | 128K | | | | | | $0.13 | $0.52 |
|
|
307
308
|
| `kilo/qwen/qwen3-30b-a3b-thinking-2507` | 82K | | | | | | $0.20 | $2 |
|
|
308
309
|
| `kilo/qwen/qwen3-32b` | 41K | | | | | | $0.08 | $0.28 |
|
|
@@ -331,8 +332,8 @@ for await (const chunk of stream) {
|
|
|
331
332
|
| `kilo/qwen/qwen3.5-flash-02-23` | 1.0M | | | | | | $0.07 | $0.26 |
|
|
332
333
|
| `kilo/qwen/qwen3.5-plus-02-15` | 1.0M | | | | | | $0.26 | $2 |
|
|
333
334
|
| `kilo/qwen/qwen3.5-plus-20260420` | 1.0M | | | | | | $0.30 | $2 |
|
|
334
|
-
| `kilo/qwen/qwen3.6-27b` |
|
|
335
|
-
| `kilo/qwen/qwen3.6-35b-a3b` | 262K | | | | | | $0.
|
|
335
|
+
| `kilo/qwen/qwen3.6-27b` | 131K | | | | | | $0.45 | $3 |
|
|
336
|
+
| `kilo/qwen/qwen3.6-35b-a3b` | 262K | | | | | | $0.14 | $1 |
|
|
336
337
|
| `kilo/qwen/qwen3.6-flash` | 1.0M | | | | | | $0.19 | $1 |
|
|
337
338
|
| `kilo/qwen/qwen3.6-max-preview` | 262K | | | | | | $1 | $6 |
|
|
338
339
|
| `kilo/qwen/qwen3.6-plus` | 1.0M | | | | | | $0.33 | $2 |
|
|
@@ -340,6 +341,7 @@ for await (const chunk of stream) {
|
|
|
340
341
|
| `kilo/qwen/qwen3.7-max` | 1.0M | | | | | | $1 | $4 |
|
|
341
342
|
| `kilo/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.32 | $1 |
|
|
342
343
|
| `kilo/qwen/qwen3.8-2.4t-a95b` | 1.0M | | | | | | $2 | $6 |
|
|
344
|
+
| `kilo/qwen/qwen3.8-27b` | 262K | | | | | | $0.45 | $3 |
|
|
343
345
|
| `kilo/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
|
|
344
346
|
| `kilo/rekaai/reka-edge` | 16K | | | | | | $0.10 | $0.10 |
|
|
345
347
|
| `kilo/rekaai/reka-flash-3` | 66K | | | | | | $0.10 | $0.20 |
|
|
@@ -360,7 +362,7 @@ for await (const chunk of stream) {
|
|
|
360
362
|
| `kilo/stepfun/step-3.7-flash` | 256K | | | | | | $0.20 | $1 |
|
|
361
363
|
| `kilo/stepfun/step-3.7-flash:free` | 262K | | | | | | — | — |
|
|
362
364
|
| `kilo/tencent/hunyuan-a13b-instruct` | 131K | | | | | | $0.14 | $0.57 |
|
|
363
|
-
| `kilo/tencent/hy3` | 262K | | | | | | $0.
|
|
365
|
+
| `kilo/tencent/hy3` | 262K | | | | | | $0.14 | $0.58 |
|
|
364
366
|
| `kilo/tencent/hy3-preview` | 262K | | | | | | $0.18 | $0.60 |
|
|
365
367
|
| `kilo/tencent/hy3:free` | 262K | | | | | | — | — |
|
|
366
368
|
| `kilo/thedrummer/cydonia-24b-v4.1` | 131K | | | | | | $0.30 | $0.50 |
|
|
@@ -380,18 +382,19 @@ for await (const chunk of stream) {
|
|
|
380
382
|
| `kilo/x-ai/grok-4.6` | 500K | | | | | | $2 | $6 |
|
|
381
383
|
| `kilo/x-ai/grok-build-0.1` | 256K | | | | | | $1 | $2 |
|
|
382
384
|
| `kilo/xiaomi/mimo-v2.5` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
383
|
-
| `kilo/xiaomi/mimo-v2.5-pro` | 1.0M | | | | | | $0.
|
|
385
|
+
| `kilo/xiaomi/mimo-v2.5-pro` | 1.0M | | | | | | $0.43 | $0.87 |
|
|
384
386
|
| `kilo/z-ai/glm-4.5` | 131K | | | | | | $0.60 | $2 |
|
|
385
387
|
| `kilo/z-ai/glm-4.5-air` | 131K | | | | | | $0.13 | $0.85 |
|
|
386
388
|
| `kilo/z-ai/glm-4.5v` | 66K | | | | | | $0.60 | $2 |
|
|
387
|
-
| `kilo/z-ai/glm-4.6` |
|
|
389
|
+
| `kilo/z-ai/glm-4.6` | 203K | | | | | | $0.50 | $2 |
|
|
388
390
|
| `kilo/z-ai/glm-4.6v` | 131K | | | | | | $0.30 | $0.90 |
|
|
389
391
|
| `kilo/z-ai/glm-4.7` | 203K | | | | | | $0.40 | $2 |
|
|
390
392
|
| `kilo/z-ai/glm-4.7-flash` | 203K | | | | | | $0.06 | $0.40 |
|
|
391
393
|
| `kilo/z-ai/glm-5` | 198K | | | | | | $0.60 | $2 |
|
|
392
394
|
| `kilo/z-ai/glm-5-turbo` | 203K | | | | | | $1 | $4 |
|
|
393
|
-
| `kilo/z-ai/glm-5.1` |
|
|
395
|
+
| `kilo/z-ai/glm-5.1` | 200K | | | | | | $1 | $4 |
|
|
394
396
|
| `kilo/z-ai/glm-5.2` | 1.0M | | | | | | $1 | $4 |
|
|
397
|
+
| `kilo/z-ai/glm-5.3` | 1.0M | | | | | | $1 | $4 |
|
|
395
398
|
| `kilo/z-ai/glm-5v-turbo` | 203K | | | | | | $1 | $4 |
|
|
396
399
|
|
|
397
400
|
## Advanced configuration
|
|
@@ -0,0 +1,73 @@
|
|
|
1
|
+
> Discover all available pages from the documentation index: https://mastra.ai/llms.txt
|
|
2
|
+
|
|
3
|
+
# Kosmik Compute
|
|
4
|
+
|
|
5
|
+
Access 1 Kosmik Compute model through Mastra's model router. Authentication is handled automatically using the `KOSMIK_API_KEY` environment variable.
|
|
6
|
+
|
|
7
|
+
Learn more in the [Kosmik Compute documentation](https://api.koscompute.com/docs/).
|
|
8
|
+
|
|
9
|
+
```bash
|
|
10
|
+
KOSMIK_API_KEY=your-api-key
|
|
11
|
+
```
|
|
12
|
+
|
|
13
|
+
```typescript
|
|
14
|
+
import { Agent } from "@mastra/core/agent";
|
|
15
|
+
|
|
16
|
+
const agent = new Agent({
|
|
17
|
+
id: "my-agent",
|
|
18
|
+
name: "My Agent",
|
|
19
|
+
instructions: "You are a helpful assistant",
|
|
20
|
+
model: "kosmik/qwen/qwen3.8-27b"
|
|
21
|
+
});
|
|
22
|
+
|
|
23
|
+
// Generate a response
|
|
24
|
+
const response = await agent.generate("Hello!");
|
|
25
|
+
|
|
26
|
+
// Stream a response
|
|
27
|
+
const stream = await agent.stream("Tell me a story");
|
|
28
|
+
for await (const chunk of stream) {
|
|
29
|
+
console.log(chunk);
|
|
30
|
+
}
|
|
31
|
+
```
|
|
32
|
+
|
|
33
|
+
> **Note:** Mastra uses the OpenAI-compatible `/chat/completions` endpoint. Some provider-specific features may not be available. Check the [Kosmik Compute documentation](https://api.koscompute.com/docs/) for details.
|
|
34
|
+
|
|
35
|
+
## Models
|
|
36
|
+
|
|
37
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
|
+
| ------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
|
+
| `kosmik/qwen/qwen3.8-27b` | 262K | | | | | | $0.35 | $2 |
|
|
40
|
+
|
|
41
|
+
## Advanced configuration
|
|
42
|
+
|
|
43
|
+
### Custom headers
|
|
44
|
+
|
|
45
|
+
```typescript
|
|
46
|
+
const agent = new Agent({
|
|
47
|
+
id: "custom-agent",
|
|
48
|
+
name: "custom-agent",
|
|
49
|
+
model: {
|
|
50
|
+
url: "https://api.koscompute.com/v1",
|
|
51
|
+
id: "kosmik/qwen/qwen3.8-27b",
|
|
52
|
+
apiKey: process.env.KOSMIK_API_KEY,
|
|
53
|
+
headers: {
|
|
54
|
+
"X-Custom-Header": "value"
|
|
55
|
+
}
|
|
56
|
+
}
|
|
57
|
+
});
|
|
58
|
+
```
|
|
59
|
+
|
|
60
|
+
### Dynamic model selection
|
|
61
|
+
|
|
62
|
+
```typescript
|
|
63
|
+
const agent = new Agent({
|
|
64
|
+
id: "dynamic-agent",
|
|
65
|
+
name: "Dynamic Agent",
|
|
66
|
+
model: ({ requestContext }) => {
|
|
67
|
+
const useAdvanced = requestContext.task === "complex";
|
|
68
|
+
return useAdvanced
|
|
69
|
+
? "kosmik/qwen/qwen3.8-27b"
|
|
70
|
+
: "kosmik/qwen/qwen3.8-27b";
|
|
71
|
+
}
|
|
72
|
+
});
|
|
73
|
+
```
|
|
@@ -55,7 +55,7 @@ for await (const chunk of stream) {
|
|
|
55
55
|
| `llmgateway/cosmos3-super-reasoner` | 262K | | | | | | $0.10 | $0.30 |
|
|
56
56
|
| `llmgateway/custom` | 128K | | | | | | — | — |
|
|
57
57
|
| `llmgateway/deepseek-v3.2` | 164K | | | | | | $0.26 | $0.38 |
|
|
58
|
-
| `llmgateway/deepseek-v4-flash` | 1.1M | | | | | | $0.
|
|
58
|
+
| `llmgateway/deepseek-v4-flash` | 1.1M | | | | | | $0.05 | $0.09 |
|
|
59
59
|
| `llmgateway/deepseek-v4-pro` | 1.1M | | | | | | $0.43 | $0.87 |
|
|
60
60
|
| `llmgateway/ernie-4.5-vl-424b-a47b` | 123K | | | | | | $0.42 | $1 |
|
|
61
61
|
| `llmgateway/fugu-ultra` | 1.0M | | | | | | $5 | $30 |
|
|
@@ -176,7 +176,7 @@ for await (const chunk of stream) {
|
|
|
176
176
|
| `llmgateway/nemotron-3-nano-30b` | 262K | | | | | | $0.06 | $0.24 |
|
|
177
177
|
| `llmgateway/nemotron-3-nano-omni` | 262K | | | | | | $0.06 | $0.24 |
|
|
178
178
|
| `llmgateway/nemotron-3-super-120b` | 262K | | | | | | $0.30 | $0.90 |
|
|
179
|
-
| `llmgateway/nemotron-3-ultra-550b` | 1.0M | | | | | | $0.50 | $
|
|
179
|
+
| `llmgateway/nemotron-3-ultra-550b` | 1.0M | | | | | | $0.50 | $2 |
|
|
180
180
|
| `llmgateway/o1` | 200K | | | | | | $15 | $60 |
|
|
181
181
|
| `llmgateway/o3` | 200K | | | | | | $2 | $8 |
|
|
182
182
|
| `llmgateway/o3-mini` | 200K | | | | | | $1 | $4 |
|
|
@@ -194,7 +194,7 @@ for await (const chunk of stream) {
|
|
|
194
194
|
| `llmgateway/qwen3-30b-a3b-instruct-2507` | 262K | | | | | | $0.10 | $0.30 |
|
|
195
195
|
| `llmgateway/qwen3-32b` | 41K | | | | | | $0.10 | $0.30 |
|
|
196
196
|
| `llmgateway/qwen3-coder-30b-a3b-instruct` | 262K | | | | | | $0.07 | $0.27 |
|
|
197
|
-
| `llmgateway/qwen3-coder-480b-a35b-instruct` | 262K | | | | | | $0.
|
|
197
|
+
| `llmgateway/qwen3-coder-480b-a35b-instruct` | 262K | | | | | | $0.38 | $2 |
|
|
198
198
|
| `llmgateway/qwen3-coder-flash` | 1.0M | | | | | | $0.30 | $2 |
|
|
199
199
|
| `llmgateway/qwen3-coder-next` | 262K | | | | | | $0.11 | $0.68 |
|
|
200
200
|
| `llmgateway/qwen3-coder-plus` | 1.0M | | | | | | $6 | $60 |
|
|
@@ -2,7 +2,7 @@
|
|
|
2
2
|
|
|
3
3
|
# LLMTR
|
|
4
4
|
|
|
5
|
-
Access
|
|
5
|
+
Access 32 LLMTR models through Mastra's model router. Authentication is handled automatically using the `LLMTR_API_KEY` environment variable.
|
|
6
6
|
|
|
7
7
|
Learn more in the [LLMTR documentation](https://llmtr.com/docs).
|
|
8
8
|
|
|
@@ -34,14 +34,39 @@ for await (const chunk of stream) {
|
|
|
34
34
|
|
|
35
35
|
## Models
|
|
36
36
|
|
|
37
|
-
| Model
|
|
38
|
-
|
|
|
39
|
-
| `llmtr/gemma-4`
|
|
40
|
-
| `llmtr/
|
|
41
|
-
| `llmtr/
|
|
42
|
-
| `llmtr/
|
|
43
|
-
| `llmtr/
|
|
44
|
-
| `llmtr/
|
|
37
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
38
|
+
| --------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
39
|
+
| `llmtr/gemma-4` | 131K | | | | | | $2 | $5 |
|
|
40
|
+
| `llmtr/google/gemini-2.5-flash-lite` | 1.0M | | | | | | $0.10 | $0.10 |
|
|
41
|
+
| `llmtr/magibu-11b-v8` | 8K | | | | | | $0.10 | $0.50 |
|
|
42
|
+
| `llmtr/meta/muse-spark-1.2-contributor` | 1.0M | | | | | | $0.10 | $0.20 |
|
|
43
|
+
| `llmtr/mimo/mimo-v2.5` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
44
|
+
| `llmtr/mimo/mimo-v2.5-pro` | 1.0M | | | | | | $0.43 | $0.87 |
|
|
45
|
+
| `llmtr/mistral/voxtral-small-latest` | 32K | | | | | | $0.10 | $0.30 |
|
|
46
|
+
| `llmtr/muse-glimmer-30b-tr` | 131K | | | | | | $2 | $5 |
|
|
47
|
+
| `llmtr/perplexity/sonar-deep-research` | 128K | | | | | | $2 | $8 |
|
|
48
|
+
| `llmtr/poolside/laguna-xs-2.1` | 262K | | | | | | — | — |
|
|
49
|
+
| `llmtr/publicai/apertus-70b-instruct` | 66K | | | | | | $0.82 | $3 |
|
|
50
|
+
| `llmtr/publicai/apertus-8b-instruct` | 66K | | | | | | $0.10 | $0.20 |
|
|
51
|
+
| `llmtr/qwen/qwen-flash` | 1.0M | | | | | | $0.05 | $0.40 |
|
|
52
|
+
| `llmtr/qwen/qwen-plus` | 1.0M | | | | | | $0.40 | $1 |
|
|
53
|
+
| `llmtr/qwen/qwen3-coder-flash` | 1.0M | | | | | | $0.30 | $2 |
|
|
54
|
+
| `llmtr/qwen/qwen3-coder-plus` | 1.0M | | | | | | $1 | $5 |
|
|
55
|
+
| `llmtr/qwen/qwen3-max` | 256K | | | | | | $1 | $6 |
|
|
56
|
+
| `llmtr/qwen/qwen3-vl-plus` | 256K | | | | | | $0.20 | $2 |
|
|
57
|
+
| `llmtr/qwen/qwen3.5-397b-a17b` | 256K | | | | | | $0.60 | $4 |
|
|
58
|
+
| `llmtr/qwen/qwen3.5-plus` | 1.0M | | | | | | $0.40 | $2 |
|
|
59
|
+
| `llmtr/qwen/qwen3.6-flash` | 1.0M | | | | | | $0.25 | $2 |
|
|
60
|
+
| `llmtr/qwen/qwen3.6-plus` | 1.0M | | | | | | $0.50 | $3 |
|
|
61
|
+
| `llmtr/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.40 | $2 |
|
|
62
|
+
| `llmtr/qwen3-6-35b` | 16K | | | | | | $5 | $10 |
|
|
63
|
+
| `llmtr/sakana/fugu-ultra` | 1.0M | | | | | | $5 | $30 |
|
|
64
|
+
| `llmtr/thinkingmachines/inkling` | 262K | | | | | | $2 | $5 |
|
|
65
|
+
| `llmtr/thinkingmachines/inkling-small` | 262K | | | | | | $0.58 | $1 |
|
|
66
|
+
| `llmtr/trendyol-asure-12b` | 41K | | | | | | $0.10 | $0.50 |
|
|
67
|
+
| `llmtr/upstage/solar-pro2` | 66K | | | | | | $0.15 | $0.60 |
|
|
68
|
+
| `llmtr/upstage/solar-pro3` | 131K | | | | | | $0.15 | $0.60 |
|
|
69
|
+
| `llmtr/upstage/solar-pro4` | 524K | | | | | | $0.03 | $0.12 |
|
|
45
70
|
|
|
46
71
|
## Advanced configuration
|
|
47
72
|
|
|
@@ -71,7 +96,7 @@ const agent = new Agent({
|
|
|
71
96
|
model: ({ requestContext }) => {
|
|
72
97
|
const useAdvanced = requestContext.task === "complex";
|
|
73
98
|
return useAdvanced
|
|
74
|
-
? "llmtr/
|
|
99
|
+
? "llmtr/upstage/solar-pro4"
|
|
75
100
|
: "llmtr/gemma-4";
|
|
76
101
|
}
|
|
77
102
|
});
|