@mastra/mcp-docs-server 1.2.23-alpha.5 → 1.2.23-alpha.8
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/.docs/docs/agents/code-mode.md +1 -1
- package/.docs/docs/agents/human-in-the-loop.md +1 -1
- package/.docs/docs/agents/networks.md +1 -1
- package/.docs/docs/agents/processors.md +1 -1
- package/.docs/docs/agents/structured-output.md +1 -1
- package/.docs/docs/auth/fga.md +1 -1
- package/.docs/docs/channels.md +2 -2
- package/.docs/docs/connections/mcp.md +1 -1
- package/.docs/docs/datasets/running-experiments.md +1 -1
- package/.docs/docs/deployment/sandbox.md +2 -2
- package/.docs/docs/deployment/workers.md +2 -2
- package/.docs/docs/evals/custom-scorers.md +1 -1
- package/.docs/docs/evals/multi-turn.md +1 -1
- package/.docs/docs/evals/overview.md +2 -2
- package/.docs/docs/evals/quick-checks.md +1 -1
- package/.docs/docs/evals/vitest-integration.md +136 -0
- package/.docs/docs/guides/context-engineering.md +1 -1
- package/.docs/docs/guides/multi-agent-systems.md +1 -1
- package/.docs/docs/guides/streaming.md +72 -52
- package/.docs/docs/harness/agent-controller.md +1 -1
- package/.docs/docs/harness/background-tasks.md +1 -1
- package/.docs/docs/harness/durable-agents.md +1 -1
- package/.docs/docs/harness/schedules.md +1 -1
- package/.docs/docs/harness/signal-providers.md +1 -1
- package/.docs/docs/harness/signals.md +1 -1
- package/.docs/docs/index.md +1 -1
- package/.docs/docs/mastra-platform/deploy.md +15 -15
- package/.docs/docs/mastra-platform/environments.md +2 -2
- package/.docs/docs/mastra-platform/github.md +2 -2
- package/.docs/docs/mastra-platform/regions.md +1 -1
- package/.docs/docs/mastra-platform/server.md +4 -4
- package/.docs/docs/mastra-platform/studio.md +1 -1
- package/.docs/docs/mastra-platform/trace-intelligence.md +1 -1
- package/.docs/docs/mastra-platform/workspaces.md +1 -1
- package/.docs/docs/memory/message-history.md +3 -3
- package/.docs/docs/memory/observational-memory.md +18 -18
- package/.docs/docs/memory/overview.md +1 -1
- package/.docs/docs/memory/working-memory.md +1 -1
- package/.docs/docs/observability/feedback.md +1 -1
- package/.docs/docs/observability/logging.md +1 -1
- package/.docs/docs/observability/tracing/overview.md +1 -1
- package/.docs/docs/sandbox/lsp.md +1 -1
- package/.docs/docs/sandbox/overview.md +1 -1
- package/.docs/docs/server/mastra-client.md +1 -1
- package/.docs/docs/server/overview.md +1 -1
- package/.docs/docs/server/pubsub.md +1 -1
- package/.docs/docs/server/request-context.md +2 -2
- package/.docs/docs/server/server-adapters.md +1 -1
- package/.docs/docs/skills.md +1 -1
- package/.docs/docs/studio/deployment.md +1 -1
- package/.docs/docs/studio/editor.md +1 -1
- package/.docs/docs/studio/overview.md +1 -1
- package/.docs/docs/subagents.md +2 -2
- package/.docs/docs/workflows/control-flow.md +1 -1
- package/.docs/docs/workflows/overview.md +1 -1
- package/.docs/docs/workflows/scheduled-workflows.md +1 -1
- package/.docs/docs/workflows/suspend-and-resume.md +2 -2
- package/.docs/integrations/sandboxes/agentcore.md +2 -0
- package/.docs/integrations/sandboxes/apple-container.md +4 -2
- package/.docs/integrations/sandboxes/blaxel.md +2 -0
- package/.docs/integrations/sandboxes/cloudflare-sandbox.md +1 -1
- package/.docs/integrations/sandboxes/daytona.md +2 -0
- package/.docs/integrations/sandboxes/docker.md +3 -1
- package/.docs/integrations/sandboxes/e2b.md +2 -0
- package/.docs/integrations/sandboxes/modal.md +3 -1
- package/.docs/integrations/sandboxes/railway.md +2 -0
- package/.docs/integrations/sandboxes/vercel.md +4 -0
- package/.docs/models/gateways/merge-gateway.md +2 -1
- package/.docs/models/gateways/netlify.md +1 -1
- package/.docs/models/gateways/openrouter.md +1 -1
- package/.docs/models/gateways/vercel.md +3 -1
- package/.docs/models/index.md +1 -1
- package/.docs/models/providers/abliteration-ai.md +7 -6
- package/.docs/models/providers/chutes.md +1 -1
- package/.docs/models/providers/coralbricks.md +4 -4
- package/.docs/models/providers/cortecs.md +3 -3
- package/.docs/models/providers/crossmodel.md +2 -2
- package/.docs/models/providers/edenai.md +7 -5
- package/.docs/models/providers/google.md +1 -2
- package/.docs/models/providers/groq.md +2 -1
- package/.docs/models/providers/hyper.md +8 -6
- package/.docs/models/providers/iteracompute.md +8 -7
- package/.docs/models/providers/kilo.md +23 -24
- package/.docs/models/providers/llmgateway-providers.md +2 -26
- package/.docs/models/providers/llmgateway.md +3 -14
- package/.docs/models/providers/nano-gpt.md +5 -3
- package/.docs/models/providers/requesty.md +4 -4
- package/.docs/models/providers/vancine.md +13 -11
- package/.docs/reference/agent-controller/agent-controller-class.md +2 -2
- package/.docs/reference/agent-controller/session.md +3 -3
- package/.docs/reference/agents/durable-agent.md +77 -9
- package/.docs/reference/agents/getDefaultGenerateOptions.md +1 -1
- package/.docs/reference/agents/listSuspendedRuns.md +2 -2
- package/.docs/reference/ai-sdk/chat-route.md +1 -1
- package/.docs/reference/ai-sdk/network-route.md +1 -1
- package/.docs/reference/ai-sdk/workflow-route.md +1 -1
- package/.docs/reference/browser/browser-viewer.md +1 -1
- package/.docs/reference/cli/mastra.md +4 -4
- package/.docs/reference/core/mastra-class.md +1 -1
- package/.docs/reference/datasets/createExperiment.md +1 -1
- package/.docs/reference/editor/tool-provider.md +1 -1
- package/.docs/reference/editor/versioning.md +1 -1
- package/.docs/reference/evals/multi-turn-judge.md +1 -1
- package/.docs/reference/evals/rubric.md +1 -1
- package/.docs/reference/file-based-agents/schedules.md +2 -2
- package/.docs/reference/file-based-agents/workspace.md +1 -1
- package/.docs/reference/manual-install.md +3 -3
- package/.docs/reference/memory/observational-memory.md +4 -4
- package/.docs/reference/memory/settled.md +1 -1
- package/.docs/reference/migrations/mastra-cloud.md +9 -9
- package/.docs/reference/migrations/upgrade-to-v1/overview.md +1 -1
- package/.docs/reference/observability/tracing/exporters/cloud-exporter.md +1 -1
- package/.docs/reference/processors/processor-interface.md +1 -1
- package/.docs/reference/processors/regex-filter-processor.md +3 -3
- package/.docs/reference/processors/token-cost-control.md +2 -2
- package/.docs/reference/processors/token-limiter-processor.md +1 -1
- package/.docs/reference/processors/tool-search-processor.md +1 -1
- package/.docs/reference/processors/working-memory-processor.md +1 -1
- package/.docs/reference/pubsub/base.md +2 -2
- package/.docs/reference/pubsub/lease-provider.md +2 -2
- package/.docs/reference/rag/vector-databases.md +33 -33
- package/.docs/reference/server/create-route.md +1 -1
- package/.docs/reference/signals/task-signal-provider.md +1 -1
- package/.docs/reference/storage/composite.md +1 -1
- package/.docs/reference/storage/retention.md +4 -4
- package/.docs/reference/streaming/ChunkType.md +1 -1
- package/.docs/reference/tools/isolated-vm-transport.md +1 -1
- package/.docs/reference/tools/mcp-client.md +2 -2
- package/.docs/reference/vectors/couchbase.md +1 -1
- package/.docs/reference/vectors/mongodb.md +2 -2
- package/.docs/reference/voice/overview.md +1 -1
- package/.docs/reference/workflows/workflow-methods/foreach.md +1 -1
- package/.docs/reference/workspace/platform-sandbox.md +3 -1
- package/.docs/reference/workspace/process-manager.md +1 -1
- package/.docs/reference/workspace/sandbox.md +20 -3
- package/.docs/reference/workspace/workspace-class.md +3 -3
- package/package.json +6 -7
- package/CHANGELOG.md +0 -5936
package/.docs/docs/subagents.md
CHANGED
|
@@ -138,7 +138,7 @@ Called after a delegation finishes. Use it to inspect results or provide feedbac
|
|
|
138
138
|
- Return `{ feedback: '...' }`: Add feedback that gets saved to the parent agent's memory and is visible to subsequent iterations
|
|
139
139
|
- Return `{ resultText: '...' }`: Replace the tool result text the parent model sees for this delegation, within the current run
|
|
140
140
|
|
|
141
|
-
|
|
141
|
+
Set `resultText` when the subagent's own result would give the parent a misleading signal. For example, empty text from a subagent stopped on a tool-calls step looks like a successful but empty delegation to the parent model. Unlike `feedback` on the next turn, `resultText` affects the parent's reasoning immediately.
|
|
142
142
|
|
|
143
143
|
```typescript
|
|
144
144
|
const stream = await parentAgent.stream('Research AI trends', {
|
|
@@ -196,7 +196,7 @@ await parentAgent.generate('Research AI trends', { maxSteps: 10, requestContext
|
|
|
196
196
|
const hookErrors = requestContext.get('__mastra_delegationHookErrors') ?? []
|
|
197
197
|
```
|
|
198
198
|
|
|
199
|
-
When `hookErrorStrategy` is `'throw'`, a throwing `onDelegationStart` blocks the subagent from running at all, and a throwing `messageFilter` or `onDelegationComplete` surfaces to the parent as a failed tool call. If `onDelegationComplete` throws while handling a failed delegation, the original delegation error
|
|
199
|
+
When `hookErrorStrategy` is `'throw'`, a throwing `onDelegationStart` blocks the subagent from running at all, and a throwing `messageFilter` or `onDelegationComplete` surfaces to the parent as a failed tool call. If `onDelegationComplete` throws while handling a failed delegation, the original delegation error remains the surfaced error. The hook isn't invoked again for its own failure.
|
|
200
200
|
|
|
201
201
|
## Message filtering
|
|
202
202
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Control flow
|
|
6
6
|
|
|
7
|
-
Workflows
|
|
7
|
+
Workflows give you control over predefined tasks by dividing each task into **steps**. Those steps can run sequentially or in parallel, and they can follow conditional paths.
|
|
8
8
|
|
|
9
9
|
Each step connects to the next in the workflow through defined schemas that keep data controlled and consistent.
|
|
10
10
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Workflows
|
|
6
6
|
|
|
7
|
-
Workflows let you define complex sequences
|
|
7
|
+
Workflows let you define complex task sequences with clear, structured steps instead of relying on one agent to reason through the entire process. You control the task breakdown along with the flow and timing of data. Workflows use the built-in execution engine by default, while managed infrastructure is available through [workflow runners](https://mastra.ai/docs/deployment/workflow-runners) such as Inngest.
|
|
8
8
|
|
|
9
9
|
**For AI agents:** Run `npx mastra api workflow run start test-workflow '{"inputData":{"message":"Hello, world!"}}'` to start this workflow directly instead of writing a temporary script. The command requires a running Mastra server; start the local server with `npx mastra dev`, or pass the reachable server's base URL with `--url`. Run `npx mastra api workflow run start --schema` before constructing different input. Install Mastra's skill with `npx skills add mastra-ai/skills --skill mastra` for complete API CLI discovery, targeting, schema, authentication, and error-handling guidance.
|
|
10
10
|
|
|
@@ -44,7 +44,7 @@ A separate "register schedule" call won't happen. The scheduler reads `schedule`
|
|
|
44
44
|
|
|
45
45
|
A workflow that declares `schedule` is auto-promoted to the **evented execution engine**. The public API (`workflow.start()`, `workflow.startAsync()`, `streamLegacy()`, `resume()`) is unchanged, `EventedWorkflow extends Workflow` and overrides each method with matching signatures. From your code, scheduled fires and manual runs are indistinguishable.
|
|
46
46
|
|
|
47
|
-
The promotion
|
|
47
|
+
The promotion means evented runs require a storage adapter with concurrent-update support, such as `@mastra/libsql`; otherwise, `createRun()` throws a clear error that points to the `schedule` field. Switch adapters or remove the schedule.
|
|
48
48
|
|
|
49
49
|
## Single schedule
|
|
50
50
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Suspend and resume
|
|
6
6
|
|
|
7
|
-
|
|
7
|
+
Pause a workflow at any step to collect additional data, wait for an API callback, throttle a costly operation, or request [human-in-the-loop](https://mastra.ai/docs/workflows/human-in-the-loop) input. Suspension saves the current execution state as a snapshot. Later, resume from a [specific step ID](https://mastra.ai/docs/workflows/snapshots) to restore the exact captured state. [Snapshots](https://mastra.ai/docs/workflows/snapshots) are stored in your configured storage provider and persist across deployments and application restarts.
|
|
8
8
|
|
|
9
9
|
## Pausing a workflow with `suspend()`
|
|
10
10
|
|
|
@@ -98,7 +98,7 @@ const stream = run.resume({
|
|
|
98
98
|
})
|
|
99
99
|
```
|
|
100
100
|
|
|
101
|
-
You can call `resume()` from anywhere in your application,
|
|
101
|
+
You can call `resume()` from anywhere in your application, such as an HTTP endpoint or event handler. Timers and code responding to [human input](https://mastra.ai/docs/workflows/human-in-the-loop) can also resume a workflow.
|
|
102
102
|
|
|
103
103
|
```typescript
|
|
104
104
|
const midnight = new Date()
|
|
@@ -92,6 +92,8 @@ if (!result?.success) {
|
|
|
92
92
|
|
|
93
93
|
**accept** (`string`): Accept header used for command event streams. (Default: `application/vnd.amazon.eventstream`)
|
|
94
94
|
|
|
95
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given, applied as a shell cd prefix. A per-command cwd always wins. The path is shell-quoted, so \~ is not expanded — use an absolute path.
|
|
96
|
+
|
|
95
97
|
**commandTimeout** (`number`): Default command timeout in milliseconds. (Default: `300000`)
|
|
96
98
|
|
|
97
99
|
**stopSessionOnLifecycle** (`boolean`): Whether stop() and destroy() should call StopRuntimeSession. Defaults to false because AgentCore Runtime sessions are often shared with agent invocations outside the sandbox instance. (Default: `false`)
|
|
@@ -125,7 +125,9 @@ console.log(response.text)
|
|
|
125
125
|
|
|
126
126
|
**labels** (`Record<string, string>`): Additional container labels. Mastra labels (mastra.sandbox, mastra.sandbox.id) are always included.
|
|
127
127
|
|
|
128
|
-
**
|
|
128
|
+
**workingDirectory** (`string`): Working directory inside the container, used as the default for command execution when no per-command cwd is given. A per-command cwd always wins. (Default: `'/workspace'`)
|
|
129
|
+
|
|
130
|
+
**workingDir** (`string`): Deprecated alias for workingDirectory. When both are set, workingDirectory wins.
|
|
129
131
|
|
|
130
132
|
**timeout** (`number`): Default command timeout in milliseconds. (Default: `300000 (5 minutes)`)
|
|
131
133
|
|
|
@@ -198,7 +200,7 @@ const sandbox = new AppleContainerSandbox({
|
|
|
198
200
|
})
|
|
199
201
|
```
|
|
200
202
|
|
|
201
|
-
These options are only applied when a new container is created. If the sandbox reconnects to an existing container with the same name, destroy and recreate the sandbox to apply changed runtime options. Apple `--tmpfs` accepts container paths only, such as `/tmp`; it doesn't accept Docker-style option specs like `/tmp:rw,size=256m`. When `readonlyRootfs` is enabled, make sure `
|
|
203
|
+
These options are only applied when a new container is created. If the sandbox reconnects to an existing container with the same name, destroy and recreate the sandbox to apply changed runtime options. Apple `--tmpfs` accepts container paths only, such as `/tmp`; it doesn't accept Docker-style option specs like `/tmp:rw,size=256m`. When `readonlyRootfs` is enabled, make sure `workingDirectory` points to a path supplied by the image or a bind mount. A writable tmpfs is also supported.
|
|
202
204
|
|
|
203
205
|
## Security model
|
|
204
206
|
|
|
@@ -83,6 +83,8 @@ const workspace = new Workspace({
|
|
|
83
83
|
|
|
84
84
|
**env** (`Record<string, string>`): Environment variables to set in the sandbox.
|
|
85
85
|
|
|
86
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins. Use an absolute path.
|
|
87
|
+
|
|
86
88
|
**labels** (`Record<string, string>`): Custom labels for the sandbox.
|
|
87
89
|
|
|
88
90
|
**runtimes** (`SandboxRuntime[]`): Supported runtimes. Valid values: 'node', 'python', 'bash', 'ruby', 'go', 'rust', 'java', 'cpp', 'r'. (Default: `['node', 'python', 'bash']`)
|
|
@@ -103,7 +103,7 @@ await workspace.sandbox?.writeFiles?.([
|
|
|
103
103
|
|
|
104
104
|
**env** (`Record<string, string>`): Environment variables applied to every command.
|
|
105
105
|
|
|
106
|
-
**workingDirectory** (`string`):
|
|
106
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins.
|
|
107
107
|
|
|
108
108
|
**commandTimeout** (`number`): Default command timeout in milliseconds. (Default: `300000`)
|
|
109
109
|
|
|
@@ -353,6 +353,8 @@ For Daytona-specific desktop APIs (regions, compressed screenshots, screen recor
|
|
|
353
353
|
|
|
354
354
|
**env** (`Record<string, string>`): Environment variables to set in the sandbox. (Default: `{}`)
|
|
355
355
|
|
|
356
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins. When set, the sandbox skips its automatic working-directory probe; when omitted, the probe fills the workingDirectory getter after start.
|
|
357
|
+
|
|
356
358
|
**labels** (`Record<string, string>`): Custom metadata labels. (Default: `{}`)
|
|
357
359
|
|
|
358
360
|
**name** (`string`): Sandbox display name. (Default: `Sandbox id`)
|
|
@@ -99,7 +99,9 @@ const agent = new Agent({
|
|
|
99
99
|
|
|
100
100
|
**tmpfs** (`Record<string, string>`): tmpfs mount paths and options. Maps to Docker HostConfig.Tmpfs.
|
|
101
101
|
|
|
102
|
-
**
|
|
102
|
+
**workingDirectory** (`string`): Working directory inside the container, used as the default for command execution when no per-command cwd is given. A per-command cwd always wins. (Default: `'/workspace'`)
|
|
103
|
+
|
|
104
|
+
**workingDir** (`string`): Deprecated alias for workingDirectory. When both are set, workingDirectory wins.
|
|
103
105
|
|
|
104
106
|
**labels** (`Record<string, string>`): Additional container labels. Mastra labels (mastra.sandbox, mastra.sandbox.id) are always included.
|
|
105
107
|
|
|
@@ -68,6 +68,8 @@ const agent = new Agent({
|
|
|
68
68
|
|
|
69
69
|
**env** (`Record<string, string>`): Environment variables to set in the sandbox
|
|
70
70
|
|
|
71
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins; when neither is set, commands run from the sandbox home directory. Use an absolute path.
|
|
72
|
+
|
|
71
73
|
**id** (`string`): Unique identifier for this sandbox instance (Default: `Auto-generated`)
|
|
72
74
|
|
|
73
75
|
**sandboxId** (`string`): Persisted E2B provider sandbox ID to reattach to deterministically. When set, start() connects to this exact sandbox (resuming it if paused) instead of discovering by logical id metadata. Only a typed "sandbox gone" error (not found, killed, or not running) falls through to the usual logical-id lookup and create ladder; auth, quota, rate-limit, timeout, and network errors propagate without creating a new sandbox. A sandbox tagged with a different logical id is refused. Read the resolved provider ID from the sandboxId property after start.
|
|
@@ -88,7 +88,9 @@ Get your credentials from the [Modal dashboard](https://modal.com/settings/token
|
|
|
88
88
|
|
|
89
89
|
**env** (`Record<string, string>`): Environment variables baked into the sandbox at create time.
|
|
90
90
|
|
|
91
|
-
**
|
|
91
|
+
**workingDirectory** (`string`): Default working directory inside the sandbox, used for command execution when no per-command cwd is given. A per-command cwd always wins.
|
|
92
|
+
|
|
93
|
+
**workdir** (`string`): Deprecated alias for workingDirectory. When both are set, workingDirectory wins.
|
|
92
94
|
|
|
93
95
|
**tokenId** (`string`): Modal token ID. Falls back to MODAL\_TOKEN\_ID environment variable.
|
|
94
96
|
|
|
@@ -221,6 +221,8 @@ const result = await sandbox.executeCommand('cat', ['/tmp/state.txt'])
|
|
|
221
221
|
|
|
222
222
|
**env** (`Record<string, string>`): Environment variables baked into the sandbox, available to every command. (Default: `{}`)
|
|
223
223
|
|
|
224
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins. Use an absolute path.
|
|
225
|
+
|
|
224
226
|
**template** (`SandboxTemplate | (base: SandboxTemplate) => SandboxTemplate`): Provision the sandbox from a custom base image built with the Railway template builder. Accepts a builder callback or a pre-built template. Ignored when sandboxId is set.
|
|
225
227
|
|
|
226
228
|
**timeout** (`number`): Default execution timeout in milliseconds applied to commands that do not specify their own timeout. Commands run until they exit when omitted.
|
|
@@ -156,6 +156,8 @@ Both callbacks are optional and can be used independently.
|
|
|
156
156
|
|
|
157
157
|
**env** (`Record<string, string>`): Default environment variables inherited by all commands. (Default: `{}`)
|
|
158
158
|
|
|
159
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins. Use an absolute path.
|
|
160
|
+
|
|
159
161
|
**metadata** (`Record<string, unknown>`): Custom metadata surfaced via getInfo(). (Default: `{}`)
|
|
160
162
|
|
|
161
163
|
**instructions** (`string | ((opts) => string)`): Override the default instructions returned by getInstructions(). Pass a string to replace them, or a function to extend the defaults.
|
|
@@ -301,6 +303,8 @@ const workspace = new Workspace({
|
|
|
301
303
|
|
|
302
304
|
**env** (`Record<string, string>`): Environment variables baked into the deployed function source. These are embedded at deploy time, not set dynamically.
|
|
303
305
|
|
|
306
|
+
**workingDirectory** (`string`): Default directory for command execution when no per-command cwd is given. A per-command cwd always wins. When omitted, commands fall back to /tmp and sandbox.workingDirectory stays undefined.
|
|
307
|
+
|
|
304
308
|
**commandTimeout** (`number`): Per-invocation command timeout in milliseconds. (Default: `55000`)
|
|
305
309
|
|
|
306
310
|
**instructions** (`string | ((opts) => string)`): Custom instructions that override the default instructions returned by getInstructions(). Pass a string to fully replace, or a function to extend the defaults.
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Merge Gateway
|
|
6
6
|
|
|
7
|
-
Merge Gateway aggregates models from multiple providers with enhanced features like rate limiting and failover. Access
|
|
7
|
+
Merge Gateway aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 177 models through Mastra's model router.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Merge Gateway documentation](https://docs.merge.dev/merge-gateway).
|
|
10
10
|
|
|
@@ -67,6 +67,7 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
67
67
|
| `deepseek/deepseek-v3.2` |
|
|
68
68
|
| `deepseek/deepseek-v4-flash` |
|
|
69
69
|
| `deepseek/deepseek-v4-flash-0731` |
|
|
70
|
+
| `deepseek/deepseek-v4-flash-0731-fast` |
|
|
70
71
|
| `deepseek/deepseek-v4-pro` |
|
|
71
72
|
| `deepseek/deepseek-v4-pro-0423` |
|
|
72
73
|
| `deepseek/deepseek-v4-pro-0813` |
|
|
@@ -111,7 +111,6 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
111
111
|
| `openrouter/~x-ai/grok-latest` |
|
|
112
112
|
| `openrouter/~z-ai/glm-latest` |
|
|
113
113
|
| `openrouter/anthracite-org/magnum-v4-72b` |
|
|
114
|
-
| `openrouter/arcee-ai/trinity-large-thinking` |
|
|
115
114
|
| `openrouter/baidu/ernie-4.5-vl-424b-a47b` |
|
|
116
115
|
| `openrouter/bytedance-seed/seed-1.6` |
|
|
117
116
|
| `openrouter/bytedance-seed/seed-1.6-flash` |
|
|
@@ -143,6 +142,7 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
143
142
|
| `openrouter/google/gemma-4-31b-it` |
|
|
144
143
|
| `openrouter/gryphe/mythomax-l2-13b` |
|
|
145
144
|
| `openrouter/ibm-granite/granite-4.1-8b` |
|
|
145
|
+
| `openrouter/ibm-granite/granite-4.2-8b` |
|
|
146
146
|
| `openrouter/inception/mercury-2` |
|
|
147
147
|
| `openrouter/inclusionai/ling-3.0-flash` |
|
|
148
148
|
| `openrouter/inclusionai/ling-3.0-flash-fin:free` |
|
|
@@ -140,10 +140,10 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
140
140
|
| `gryphe/mythomax-l2-13b` |
|
|
141
141
|
| `ibm-granite/granite-4.0-h-micro` |
|
|
142
142
|
| `ibm-granite/granite-4.1-8b` |
|
|
143
|
+
| `ibm-granite/granite-4.2-8b` |
|
|
143
144
|
| `inception/mercury-2` |
|
|
144
145
|
| `inclusionai/ling-3.0-flash` |
|
|
145
146
|
| `inclusionai/ling-3.0-flash-fin:free` |
|
|
146
|
-
| `kwaipilot/kat-coder-air-v2.5` |
|
|
147
147
|
| `kwaipilot/kat-coder-pro-v2` |
|
|
148
148
|
| `kwaipilot/kat-coder-pro-v2.5` |
|
|
149
149
|
| `liquid/lfm-2.5-2.6b:free` |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Vercel
|
|
6
6
|
|
|
7
|
-
Vercel aggregates models from multiple providers with enhanced features like rate limiting and failover. Access
|
|
7
|
+
Vercel aggregates models from multiple providers with enhanced features like rate limiting and failover. Access 363 models through Mastra's model router.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Vercel documentation](https://ai-sdk.dev/providers/ai-sdk-providers).
|
|
10
10
|
|
|
@@ -69,6 +69,7 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
69
69
|
| `alibaba/qwen3.8-2.4t-a95b` |
|
|
70
70
|
| `alibaba/qwen3.8-27b` |
|
|
71
71
|
| `alibaba/qwen3.8-flash` |
|
|
72
|
+
| `alibaba/qwen3.8-flash-next` |
|
|
72
73
|
| `alibaba/qwen3.8-max` |
|
|
73
74
|
| `alibaba/wan-v2.5-t2v-preview` |
|
|
74
75
|
| `alibaba/wan-v2.6-i2v` |
|
|
@@ -384,6 +385,7 @@ ANTHROPIC_API_KEY=ant-...
|
|
|
384
385
|
| `voyage/voyage-law-2` |
|
|
385
386
|
| `xiaomi/mimo-v2.5` |
|
|
386
387
|
| `xiaomi/mimo-v2.5-pro` |
|
|
388
|
+
| `xiaomi/mimo-v2.5-pro-ultraspeed` |
|
|
387
389
|
| `zai/glm-4.5` |
|
|
388
390
|
| `zai/glm-4.5-air` |
|
|
389
391
|
| `zai/glm-4.5v` |
|
package/.docs/models/index.md
CHANGED
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Model Providers
|
|
6
6
|
|
|
7
|
-
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to
|
|
7
|
+
Mastra provides a unified interface for working with LLMs across multiple providers, giving you access to 7023 models from 199 providers through a single API.
|
|
8
8
|
|
|
9
9
|
## Features
|
|
10
10
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# abliteration.ai
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 3 abliteration.ai models through Mastra's model router. Authentication is handled automatically using the `ABLIT_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [abliteration.ai documentation](https://docs.abliteration.ai/models).
|
|
10
10
|
|
|
@@ -36,10 +36,11 @@ for await (const chunk of stream) {
|
|
|
36
36
|
|
|
37
37
|
## Models
|
|
38
38
|
|
|
39
|
-
| Model
|
|
40
|
-
|
|
|
41
|
-
| `abliteration-ai/abliterated-model`
|
|
42
|
-
| `abliteration-ai/abliterated-model-large`
|
|
39
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
40
|
+
| -------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
41
|
+
| `abliteration-ai/abliterated-model` | 150K | | | | | | $3 | $3 |
|
|
42
|
+
| `abliteration-ai/abliterated-model-large` | 1.0M | | | | | | $5 | $5 |
|
|
43
|
+
| `abliteration-ai/abliterated-model-large-v2` | 1.0M | | | | | | $5 | $5 |
|
|
43
44
|
|
|
44
45
|
## Advanced configuration
|
|
45
46
|
|
|
@@ -69,7 +70,7 @@ const agent = new Agent({
|
|
|
69
70
|
model: ({ requestContext }) => {
|
|
70
71
|
const useAdvanced = requestContext.task === "complex";
|
|
71
72
|
return useAdvanced
|
|
72
|
-
? "abliteration-ai/abliterated-model-large"
|
|
73
|
+
? "abliteration-ai/abliterated-model-large-v2"
|
|
73
74
|
: "abliteration-ai/abliterated-model";
|
|
74
75
|
}
|
|
75
76
|
});
|
|
@@ -48,7 +48,7 @@ for await (const chunk of stream) {
|
|
|
48
48
|
| `chutes/Qwen/Qwen3-32B-TEE` | 41K | | | | | | $0.10 | $0.42 |
|
|
49
49
|
| `chutes/Qwen/Qwen3.5-397B-A17B-TEE` | 262K | | | | | | $0.45 | $3 |
|
|
50
50
|
| `chutes/Qwen/Qwen3.6-27B-TEE` | 262K | | | | | | $0.30 | $2 |
|
|
51
|
-
| `chutes/Qwen/Qwen3.8-27B-TEE` | 262K | | | | | | $0.
|
|
51
|
+
| `chutes/Qwen/Qwen3.8-27B-TEE` | 262K | | | | | | $0.32 | $3 |
|
|
52
52
|
| `chutes/unsloth/Mistral-Nemo-Instruct-2407-TEE` | 131K | | | | | | $0.02 | $0.10 |
|
|
53
53
|
| `chutes/zai-org/GLM-5.1-TEE` | 203K | | | | | | $0.98 | $3 |
|
|
54
54
|
| `chutes/zai-org/GLM-5.2-TEE` | 1.0M | | | | | | $1 | $4 |
|
|
@@ -19,7 +19,7 @@ const agent = new Agent({
|
|
|
19
19
|
id: "my-agent",
|
|
20
20
|
name: "My Agent",
|
|
21
21
|
instructions: "You are a helpful assistant",
|
|
22
|
-
model: "coralbricks/glm-5.
|
|
22
|
+
model: "coralbricks/glm-5.3-fp4"
|
|
23
23
|
});
|
|
24
24
|
|
|
25
25
|
// Generate a response
|
|
@@ -38,7 +38,7 @@ for await (const chunk of stream) {
|
|
|
38
38
|
|
|
39
39
|
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
40
40
|
| -------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
41
|
-
| `coralbricks/glm-5.
|
|
41
|
+
| `coralbricks/glm-5.3-fp4` | 1.0M | | | | | | $1 | $4 |
|
|
42
42
|
| `coralbricks/gpt-oss-120b` | 131K | | | | | | $0.12 | $0.60 |
|
|
43
43
|
| `coralbricks/kimi-k3` | 1.0M | | | | | | $3 | $15 |
|
|
44
44
|
|
|
@@ -52,7 +52,7 @@ const agent = new Agent({
|
|
|
52
52
|
name: "custom-agent",
|
|
53
53
|
model: {
|
|
54
54
|
url: "https://inference.coralbricks.ai/v1",
|
|
55
|
-
id: "coralbricks/glm-5.
|
|
55
|
+
id: "coralbricks/glm-5.3-fp4",
|
|
56
56
|
apiKey: process.env.CORAL_API_KEY,
|
|
57
57
|
headers: {
|
|
58
58
|
"X-Custom-Header": "value"
|
|
@@ -71,7 +71,7 @@ const agent = new Agent({
|
|
|
71
71
|
const useAdvanced = requestContext.task === "complex";
|
|
72
72
|
return useAdvanced
|
|
73
73
|
? "coralbricks/kimi-k3"
|
|
74
|
-
: "coralbricks/glm-5.
|
|
74
|
+
: "coralbricks/glm-5.3-fp4";
|
|
75
75
|
}
|
|
76
76
|
});
|
|
77
77
|
```
|
|
@@ -55,7 +55,8 @@ for await (const chunk of stream) {
|
|
|
55
55
|
| `cortecs/deepseek-v3.2` | 164K | | | | | | $0.30 | $0.49 |
|
|
56
56
|
| `cortecs/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.13 | $0.28 |
|
|
57
57
|
| `cortecs/deepseek-v4-pro` | 1.0M | | | | | | $2 | $3 |
|
|
58
|
-
| `cortecs/
|
|
58
|
+
| `cortecs/deepseek-v4-pro-0813` | 1.0M | | | | | | $2 | $4 |
|
|
59
|
+
| `cortecs/devstral-2512` | 256K | | | | | | $0.48 | $2 |
|
|
59
60
|
| `cortecs/gemini-2.5-flash` | 1.0M | | | | | | $0.30 | $2 |
|
|
60
61
|
| `cortecs/gemini-2.5-pro` | 1.0M | | | | | | $1 | $10 |
|
|
61
62
|
| `cortecs/gemini-3.1-flash-lite` | 1.0M | | | | | | $0.27 | $2 |
|
|
@@ -104,7 +105,7 @@ for await (const chunk of stream) {
|
|
|
104
105
|
| `cortecs/minicpm-v-4.5` | 32K | | | | | | $0.65 | $1 |
|
|
105
106
|
| `cortecs/minimax-m2` | 400K | | | | | | $0.35 | $1 |
|
|
106
107
|
| `cortecs/minimax-m2.1` | 196K | | | | | | $0.36 | $1 |
|
|
107
|
-
| `cortecs/minimax-m2.5` |
|
|
108
|
+
| `cortecs/minimax-m2.5` | 196K | | | | | | $0.30 | $1 |
|
|
108
109
|
| `cortecs/minimax-m2.7` | 197K | | | | | | $0.67 | $3 |
|
|
109
110
|
| `cortecs/minimax-m3` | 1.0M | | | | | | $0.40 | $2 |
|
|
110
111
|
| `cortecs/ministral-14b-2512` | 256K | | | | | | $0.22 | $0.22 |
|
|
@@ -114,7 +115,6 @@ for await (const chunk of stream) {
|
|
|
114
115
|
| `cortecs/mistral-7b-instruct-v0.3` | 127K | | | | | | $0.11 | $0.11 |
|
|
115
116
|
| `cortecs/mistral-large-2402` | 32K | | | | | | $4 | $13 |
|
|
116
117
|
| `cortecs/mistral-large-2512` | 256K | | | | | | $0.56 | $2 |
|
|
117
|
-
| `cortecs/mistral-medium-2508` | 128K | | | | | | $0.45 | $2 |
|
|
118
118
|
| `cortecs/mistral-medium-3.5` | 256K | | | | | | $1 | $7 |
|
|
119
119
|
| `cortecs/mistral-nemo-instruct-2407` | 128K | | | | | | $0.14 | $0.14 |
|
|
120
120
|
| `cortecs/mistral-small-2503` | 128K | | | | | | $0.11 | $0.33 |
|
|
@@ -75,8 +75,8 @@ for await (const chunk of stream) {
|
|
|
75
75
|
| `crossmodel/qwen/qwen3.6-flash` | 1.0M | | | | | | $0.19 | $1 |
|
|
76
76
|
| `crossmodel/qwen/qwen3.6-plus` | 1.0M | | | | | | $0.32 | $2 |
|
|
77
77
|
| `crossmodel/qwen/qwen3.7-flash` | 1.0M | | | | | | $0.04 | $0.13 |
|
|
78
|
-
| `crossmodel/qwen/qwen3.7-max` | 1.0M | | | | | | $2 | $
|
|
79
|
-
| `crossmodel/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.
|
|
78
|
+
| `crossmodel/qwen/qwen3.7-max` | 1.0M | | | | | | $2 | $6 |
|
|
79
|
+
| `crossmodel/qwen/qwen3.7-plus` | 1.0M | | | | | | $0.32 | $1 |
|
|
80
80
|
| `crossmodel/qwen/qwen3.8-flash` | 1.0M | | | | | | $0.13 | $0.43 |
|
|
81
81
|
| `crossmodel/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
|
|
82
82
|
| `crossmodel/tencent/hy3` | 262K | | | | | | $0.16 | $0.64 |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Eden AI
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 240 Eden AI models through Mastra's model router. Authentication is handled automatically using the `EDENAI_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Eden AI documentation](https://docs.edenai.co).
|
|
10
10
|
|
|
@@ -107,11 +107,12 @@ for await (const chunk of stream) {
|
|
|
107
107
|
| `edenai/fireworks_ai/accounts/fireworks/models/gpt-oss-120b` | 131K | | | | | | $0.15 | $0.60 |
|
|
108
108
|
| `edenai/fireworks_ai/accounts/fireworks/models/muse-glimmer-30b` | 131K | | | | | | $0.35 | $2 |
|
|
109
109
|
| `edenai/fireworks_ai/gpt-oss-120b` | 131K | | | | | | $0.15 | $0.60 |
|
|
110
|
-
| `edenai/flexai/deepseek-v4-flash-0731` | 786K | | | | | | $0.
|
|
110
|
+
| `edenai/flexai/deepseek-v4-flash-0731` | 786K | | | | | | $0.03 | $0.10 |
|
|
111
111
|
| `edenai/flexai/gpt-oss-120b` | 131K | | | | | | $0.04 | $0.10 |
|
|
112
112
|
| `edenai/flexai/gpt-oss-20b` | 131K | | | | | | $0.02 | $0.10 |
|
|
113
113
|
| `edenai/flexai/Muse-Glimmer-30B` | 131K | | | | | | $0.30 | $1 |
|
|
114
114
|
| `edenai/flexai/Nemotron-3-Super-120B-A12B` | 262K | | | | | | $0.09 | $0.40 |
|
|
115
|
+
| `edenai/flexai/Step-3.7-Flash` | 262K | | | | | | $0.20 | $1 |
|
|
115
116
|
| `edenai/google/gemini-2.5-flash-image` | 33K | | | | | | $0.30 | $3 |
|
|
116
117
|
| `edenai/google/gemini-3-flash-preview` | 1.0M | | | | | | $0.50 | $3 |
|
|
117
118
|
| `edenai/google/gemini-3-pro-image` | 66K | | | | | | $2 | $12 |
|
|
@@ -131,8 +132,8 @@ for await (const chunk of stream) {
|
|
|
131
132
|
| `edenai/google/gemini-pro-latest` | 1.0M | | | | | | $2 | $12 |
|
|
132
133
|
| `edenai/groq/openai/gpt-oss-120b` | 131K | | | | | | $0.15 | $0.60 |
|
|
133
134
|
| `edenai/groq/openai/gpt-oss-20b` | 131K | | | | | | $0.07 | $0.30 |
|
|
134
|
-
| `edenai/ionos/meta-llama/Llama-3.3-70B-Instruct` | 128K | | | | | | $0.
|
|
135
|
-
| `edenai/ionos/openai/gpt-oss-120b` | 131K | | | | | | $0.17 | $0.
|
|
135
|
+
| `edenai/ionos/meta-llama/Llama-3.3-70B-Instruct` | 128K | | | | | | $0.75 | $0.75 |
|
|
136
|
+
| `edenai/ionos/openai/gpt-oss-120b` | 131K | | | | | | $0.17 | $0.75 |
|
|
136
137
|
| `edenai/minimax/MiniMax-M2` | 205K | | | | | | $0.30 | $1 |
|
|
137
138
|
| `edenai/minimax/MiniMax-M2.1` | 205K | | | | | | $0.30 | $1 |
|
|
138
139
|
| `edenai/minimax/MiniMax-M2.5` | 205K | | | | | | $0.30 | $1 |
|
|
@@ -152,6 +153,7 @@ for await (const chunk of stream) {
|
|
|
152
153
|
| `edenai/mistral/voxtral-small-latest` | 33K | | | | | | $0.10 | $0.40 |
|
|
153
154
|
| `edenai/moonshot/kimi-k2.6` | 262K | | | | | | $0.95 | $4 |
|
|
154
155
|
| `edenai/moonshot/kimi-k2.7-code` | 262K | | | | | | $0.95 | $4 |
|
|
156
|
+
| `edenai/moonshot/kimi-k2.7-code-highspeed` | 262K | | | | | | $2 | $8 |
|
|
155
157
|
| `edenai/moonshot/kimi-k3` | 1.0M | | | | | | $3 | $15 |
|
|
156
158
|
| `edenai/nebius/deepseek-ai/DeepSeek-V4-Flash-0731` | 1.0M | | | | | | $0.14 | $0.28 |
|
|
157
159
|
| `edenai/nebius/meta-llama/Llama-3.3-70B-Instruct` | 131K | | | | | | $0.13 | $0.40 |
|
|
@@ -223,7 +225,7 @@ for await (const chunk of stream) {
|
|
|
223
225
|
| `edenai/qwen/qwen3.8-flash` | 1.0M | | | | | | $0.16 | $0.47 |
|
|
224
226
|
| `edenai/qwen/qwen3.8-max` | 1.0M | | | | | | $2 | $6 |
|
|
225
227
|
| `edenai/qwen/qwq-plus` | 131K | | | | | | $0.80 | $2 |
|
|
226
|
-
| `edenai/scaleway/deepseek-v4-flash-0731` | 256K | | | | | | $0.
|
|
228
|
+
| `edenai/scaleway/deepseek-v4-flash-0731` | 256K | | | | | | $0.46 | $0.93 |
|
|
227
229
|
| `edenai/scaleway/gpt-oss-120b` | 128K | | | | | | $0.17 | $0.70 |
|
|
228
230
|
| `edenai/scaleway/llama-3.3-70b-instruct` | 128K | | | | | | $1 | $1 |
|
|
229
231
|
| `edenai/tensorx/deepseek/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.25 | $0.30 |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Google
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 38 Google models through Mastra's model router. Authentication is handled automatically using one of the following environment variables: `GOOGLE_API_KEY`, `GOOGLE_GENERATIVE_AI_API_KEY`.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Google documentation](https://ai.google.dev/gemini-api/docs/models).
|
|
10
10
|
|
|
@@ -67,7 +67,6 @@ for await (const chunk of stream) {
|
|
|
67
67
|
| `google/gemini-flash-latest` | 1.0M | | | | | | $0.75 | $4 |
|
|
68
68
|
| `google/gemini-flash-lite-latest` | 1.0M | | | | | | $0.30 | $3 |
|
|
69
69
|
| `google/gemini-omni-flash-preview` | 131K | | | | | | $2 | $18 |
|
|
70
|
-
| `google/gemini-robotics-er-1.6-preview` | 131K | | | | | | $1 | $5 |
|
|
71
70
|
| `google/gemma-4-26b-a4b-it` | 262K | | | | | | — | — |
|
|
72
71
|
| `google/gemma-4-31b-it` | 262K | | | | | | — | — |
|
|
73
72
|
| `google/lyria-3-clip-preview` | 1.0M | | | | | | — | — |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Groq
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 16 Groq models through Mastra's model router. Authentication is handled automatically using the `GROQ_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Groq documentation](https://console.groq.com/docs/models).
|
|
10
10
|
|
|
@@ -49,6 +49,7 @@ for await (const chunk of stream) {
|
|
|
49
49
|
| `groq/openai/gpt-oss-20b` | 131K | | | | | | $0.07 | $0.30 |
|
|
50
50
|
| `groq/openai/gpt-oss-safeguard-20b` | 131K | | | | | | $0.07 | $0.30 |
|
|
51
51
|
| `groq/qwen/qwen3.6-27b` | 131K | | | | | | $0.60 | $3 |
|
|
52
|
+
| `groq/qwen/qwen3.8-27b` | 131K | | | | | | $0.80 | $4 |
|
|
52
53
|
| `groq/whisper-large-v3` | — | | | | | | — | — |
|
|
53
54
|
| `groq/whisper-large-v3-turbo` | — | | | | | | — | — |
|
|
54
55
|
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# Charm Hyper
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 31 Charm Hyper models through Mastra's model router. Authentication is handled automatically using the `HYPER_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [Charm Hyper documentation](https://hyper.charm.land).
|
|
10
10
|
|
|
@@ -42,18 +42,20 @@ for await (const chunk of stream) {
|
|
|
42
42
|
| `hyper/deepseek-v4-flash-0731` | 1.0M | | | | | | $0.44 | $1 |
|
|
43
43
|
| `hyper/deepseek-v4-pro` | 1.0M | | | | | | $2 | $5 |
|
|
44
44
|
| `hyper/deepseek-v4-pro-0813` | 1.0M | | | | | | $1 | $4 |
|
|
45
|
-
| `hyper/gemma-4-26b-a4b-it` | 256K | | | | | | $0.
|
|
46
|
-
| `hyper/glm-5` | 203K | | | | | | $0.
|
|
45
|
+
| `hyper/gemma-4-26b-a4b-it` | 256K | | | | | | $0.12 | $0.42 |
|
|
46
|
+
| `hyper/glm-5` | 203K | | | | | | $0.91 | $3 |
|
|
47
47
|
| `hyper/glm-5.1` | 203K | | | | | | $1 | $4 |
|
|
48
48
|
| `hyper/glm-5.2` | 1.0M | | | | | | $2 | $5 |
|
|
49
|
-
| `hyper/
|
|
49
|
+
| `hyper/glm-5.3` | 1.0M | | | | | | $2 | $5 |
|
|
50
|
+
| `hyper/glm-5.3-flash` | 1.0M | | | | | | $0.16 | $0.54 |
|
|
51
|
+
| `hyper/gpt-oss-120b` | 128K | | | | | | $0.19 | $0.70 |
|
|
50
52
|
| `hyper/kimi-k2.5` | 262K | | | | | | $0.54 | $3 |
|
|
51
53
|
| `hyper/kimi-k2.6` | 262K | | | | | | $1 | $4 |
|
|
52
54
|
| `hyper/kimi-k2.7-code` | 262K | | | | | | $1 | $4 |
|
|
53
55
|
| `hyper/kimi-k3` | 1.0M | | | | | | $3 | $16 |
|
|
54
56
|
| `hyper/llama-3.3-70b-instruct` | 128K | | | | | | $0.61 | $1 |
|
|
55
|
-
| `hyper/llama-4-maverick-17b-128e-instruct-fp8` | 430K | | | | | | $0.
|
|
56
|
-
| `hyper/minimax-m2.7` | 262K | | | | | | $0.
|
|
57
|
+
| `hyper/llama-4-maverick-17b-128e-instruct-fp8` | 430K | | | | | | $0.27 | $0.90 |
|
|
58
|
+
| `hyper/minimax-m2.7` | 262K | | | | | | $0.42 | $2 |
|
|
57
59
|
| `hyper/minimax-m3` | 512K | | | | | | $0.33 | $1 |
|
|
58
60
|
| `hyper/qwen3-coder-480b-a35b-instruct-int4-mixed-ar` | 106K | | | | | | $0.45 | $2 |
|
|
59
61
|
| `hyper/qwen3-next-80b-a3b-instruct` | 262K | | | | | | $0.12 | $1 |
|
|
@@ -4,7 +4,7 @@
|
|
|
4
4
|
|
|
5
5
|
# IteraCompute
|
|
6
6
|
|
|
7
|
-
Access
|
|
7
|
+
Access 2 IteraCompute models through Mastra's model router. Authentication is handled automatically using the `ITERACOMPUTE_API_KEY` environment variable.
|
|
8
8
|
|
|
9
9
|
Learn more in the [IteraCompute documentation](https://iteracompute.com/docs.html).
|
|
10
10
|
|
|
@@ -19,7 +19,7 @@ const agent = new Agent({
|
|
|
19
19
|
id: "my-agent",
|
|
20
20
|
name: "My Agent",
|
|
21
21
|
instructions: "You are a helpful assistant",
|
|
22
|
-
model: "iteracompute/iteracompute/
|
|
22
|
+
model: "iteracompute/iteracompute/ornith-1.5-35b-a3b"
|
|
23
23
|
});
|
|
24
24
|
|
|
25
25
|
// Generate a response
|
|
@@ -36,9 +36,10 @@ for await (const chunk of stream) {
|
|
|
36
36
|
|
|
37
37
|
## Models
|
|
38
38
|
|
|
39
|
-
| Model
|
|
40
|
-
|
|
|
41
|
-
| `iteracompute/iteracompute/
|
|
39
|
+
| Model | Context | Tools | Reasoning | Image | Audio | Video | Input $/1M | Output $/1M |
|
|
40
|
+
| ---------------------------------------------- | ------- | ----- | --------- | ----- | ----- | ----- | ---------- | ----------- |
|
|
41
|
+
| `iteracompute/iteracompute/ornith-1.5-35b-a3b` | 328K | | | | | | $0.30 | $3 |
|
|
42
|
+
| `iteracompute/iteracompute/qwen3.8-27b` | 328K | | | | | | $0.30 | $3 |
|
|
42
43
|
|
|
43
44
|
## Advanced configuration
|
|
44
45
|
|
|
@@ -50,7 +51,7 @@ const agent = new Agent({
|
|
|
50
51
|
name: "custom-agent",
|
|
51
52
|
model: {
|
|
52
53
|
url: "https://api.iteracompute.com/v1",
|
|
53
|
-
id: "iteracompute/iteracompute/
|
|
54
|
+
id: "iteracompute/iteracompute/ornith-1.5-35b-a3b",
|
|
54
55
|
apiKey: process.env.ITERACOMPUTE_API_KEY,
|
|
55
56
|
headers: {
|
|
56
57
|
"X-Custom-Header": "value"
|
|
@@ -69,7 +70,7 @@ const agent = new Agent({
|
|
|
69
70
|
const useAdvanced = requestContext.task === "complex";
|
|
70
71
|
return useAdvanced
|
|
71
72
|
? "iteracompute/iteracompute/qwen3.8-27b"
|
|
72
|
-
: "iteracompute/iteracompute/
|
|
73
|
+
: "iteracompute/iteracompute/ornith-1.5-35b-a3b";
|
|
73
74
|
}
|
|
74
75
|
});
|
|
75
76
|
```
|