corent-mcp 0.4.2 → 0.5.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/server.js +27 -4
- package/package.json +1 -1
package/dist/server.js
CHANGED
|
@@ -1,5 +1,5 @@
|
|
|
1
1
|
/**
|
|
2
|
-
* Shared Corent MCP server definition — the
|
|
2
|
+
* Shared Corent MCP server definition — the 8 tools, used by both the
|
|
3
3
|
* stdio entrypoint (index.ts, for local `npx` use) and the hosted HTTP
|
|
4
4
|
* entrypoint (http.ts). Keeping the tools in one place means the two
|
|
5
5
|
* transports can never drift apart.
|
|
@@ -17,7 +17,7 @@ import { McpServer } from "@modelcontextprotocol/sdk/server/mcp.js";
|
|
|
17
17
|
import { z } from "zod";
|
|
18
18
|
import { MEDIA_WIDGET_HTML } from "./widget-html.js";
|
|
19
19
|
export const DEFAULT_API_URL = "https://api.corent.tech";
|
|
20
|
-
export const SERVER_VERSION = "0.
|
|
20
|
+
export const SERVER_VERSION = "0.5.0";
|
|
21
21
|
// --- MCP Apps (SEP-1865): the media tools render an inline widget in hosts
|
|
22
22
|
// that support it (claude.ai web/desktop, others). Hosts without the extension
|
|
23
23
|
// ignore the _meta and fall back to plain text results, so this is additive.
|
|
@@ -212,7 +212,7 @@ export function createCorentServer(config = {}) {
|
|
|
212
212
|
resolution: z
|
|
213
213
|
.enum(["720p", "1080p", "4k"])
|
|
214
214
|
.optional()
|
|
215
|
-
.describe("Pixel resolution (default 720p). Tier-capped:
|
|
215
|
+
.describe("Pixel resolution (default 720p). Tier-capped: pro allows up to 1080p, max_pro up to 4k; above the tier's cap it is clamped down, never rejected. Higher resolutions cost more (1080p ~2.5x, 4k ~5.5x). GET /v1/tiers lists each tier's menu with prices."),
|
|
216
216
|
image_url: z.string().url().optional().describe("If set, animates this image instead of pure text-to-video"),
|
|
217
217
|
},
|
|
218
218
|
outputSchema: jobOutput,
|
|
@@ -221,6 +221,29 @@ export function createCorentServer(config = {}) {
|
|
|
221
221
|
method: "POST",
|
|
222
222
|
body: JSON.stringify({ prompt, tier, preference, aspect_ratio, duration_s, resolution, image_url }),
|
|
223
223
|
})));
|
|
224
|
+
server.registerTool("generate_speech", {
|
|
225
|
+
title: "Generate speech (voice)",
|
|
226
|
+
description: "Generate spoken audio (text to speech) from text. Synchronous: returns the finished audio URL and the exact charge. Billed per 1,000 characters of input text (a few cents for a short narration). Failed generations are never billed.",
|
|
227
|
+
inputSchema: {
|
|
228
|
+
text: z.string().min(1).max(5000).describe("The text to speak, max 5000 characters"),
|
|
229
|
+
voice_id: z
|
|
230
|
+
.string()
|
|
231
|
+
.regex(/^[A-Za-z0-9]{8,64}$/)
|
|
232
|
+
.optional()
|
|
233
|
+
.describe("Optional provider voice id; omit for the default narration voice"),
|
|
234
|
+
},
|
|
235
|
+
outputSchema: {
|
|
236
|
+
id: z.string().optional(),
|
|
237
|
+
status: z.string().optional(),
|
|
238
|
+
audio_url: z.string().optional(),
|
|
239
|
+
meta: mediaMeta.optional(),
|
|
240
|
+
error: z.string().optional(),
|
|
241
|
+
},
|
|
242
|
+
annotations: { readOnlyHint: false, destructiveHint: false, idempotentHint: false, openWorldHint: true },
|
|
243
|
+
}, wrap(async ({ text, voice_id }) => corent("/v1/audio/speech", {
|
|
244
|
+
method: "POST",
|
|
245
|
+
body: JSON.stringify({ text, voice_id }),
|
|
246
|
+
})));
|
|
224
247
|
server.registerTool("get_job", {
|
|
225
248
|
title: "Check job status",
|
|
226
249
|
_meta: WIDGET_TOOL_META,
|
|
@@ -248,7 +271,7 @@ export function createCorentServer(config = {}) {
|
|
|
248
271
|
// prompt. The agent never picks a model or tier. ---
|
|
249
272
|
server.registerTool("plan", {
|
|
250
273
|
title: "Plan (cost preview, no generation)",
|
|
251
|
-
description: "Preview how Corent would handle a plain-language request WITHOUT generating anything (costs a fraction of a cent). Corent decides whether it's an image or video, which tier, aspect ratio, and duration, and returns the plan plus an estimated cost. Use this to decide or confirm cost before spending. If the request is something Corent can't generate (
|
|
274
|
+
description: "Preview how Corent would handle a plain-language request WITHOUT generating anything (costs a fraction of a cent). Corent decides whether it's an image or video, which tier, aspect ratio, and duration, and returns the plan plus an estimated cost. Use this to decide or confirm cost before spending. The planner covers image and video; for spoken audio call generate_speech directly. If the request is something Corent can't generate (text documents, 3D, editing, real-world actions), can_fulfill is false with a reason.",
|
|
252
275
|
inputSchema: {
|
|
253
276
|
intent: z
|
|
254
277
|
.string()
|
package/package.json
CHANGED