mcp-scraper 0.61.0 → 0.62.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +1 -1
- package/dist/bin/api-server.cjs +3 -3
- package/dist/bin/api-server.cjs.map +1 -1
- package/dist/bin/api-server.js +1 -1
- package/dist/bin/mcp-scraper-cli.cjs +1 -1
- package/dist/bin/mcp-scraper-cli.cjs.map +1 -1
- package/dist/bin/mcp-scraper-cli.js +1 -1
- package/dist/bin/mcp-scraper-install.cjs +1 -1
- package/dist/bin/mcp-scraper-install.cjs.map +1 -1
- package/dist/bin/mcp-scraper-install.js +1 -1
- package/dist/bin/mcp-stdio-server.cjs +3 -3
- package/dist/bin/mcp-stdio-server.cjs.map +1 -1
- package/dist/bin/mcp-stdio-server.js +2 -2
- package/dist/chunk-RAPIYV2X.js +7 -0
- package/dist/chunk-RAPIYV2X.js.map +1 -0
- package/dist/{chunk-CIKVUFC5.js → chunk-Y2RIYBYR.js} +4 -4
- package/dist/chunk-Y2RIYBYR.js.map +1 -0
- package/dist/{server-LQ4EMLSN.js → server-D3QL46Q3.js} +3 -3
- package/package.json +1 -1
- package/dist/chunk-CIKVUFC5.js.map +0 -1
- package/dist/chunk-E6QGM3ZA.js +0 -7
- package/dist/chunk-E6QGM3ZA.js.map +0 -1
- /package/dist/{server-LQ4EMLSN.js.map → server-D3QL46Q3.js.map} +0 -0
|
@@ -13,7 +13,7 @@ import {
|
|
|
13
13
|
registerScheduledResultsMcpTools,
|
|
14
14
|
registerSerpIntelligenceCaptureTools,
|
|
15
15
|
resolveDeploymentProfile
|
|
16
|
-
} from "../chunk-
|
|
16
|
+
} from "../chunk-Y2RIYBYR.js";
|
|
17
17
|
import "../chunk-2K74LVV7.js";
|
|
18
18
|
import "../chunk-SQMKUPA5.js";
|
|
19
19
|
import "../chunk-57Y5QHI3.js";
|
|
@@ -22,7 +22,7 @@ import {
|
|
|
22
22
|
} from "../chunk-IMBMDZUO.js";
|
|
23
23
|
import {
|
|
24
24
|
PACKAGE_VERSION
|
|
25
|
-
} from "../chunk-
|
|
25
|
+
} from "../chunk-RAPIYV2X.js";
|
|
26
26
|
import "../chunk-2U2ZDDSI.js";
|
|
27
27
|
import "../chunk-B7QNUVBT.js";
|
|
28
28
|
import "../chunk-6OHW5XZN.js";
|
|
@@ -0,0 +1 @@
|
|
|
1
|
+
{"version":3,"sources":["../src/version.ts"],"sourcesContent":["export const PACKAGE_VERSION = '0.62.0'\n"],"mappings":";AAAO,IAAM,kBAAkB;","names":[]}
|
|
@@ -21,7 +21,7 @@ import {
|
|
|
21
21
|
} from "./chunk-57Y5QHI3.js";
|
|
22
22
|
import {
|
|
23
23
|
PACKAGE_VERSION
|
|
24
|
-
} from "./chunk-
|
|
24
|
+
} from "./chunk-RAPIYV2X.js";
|
|
25
25
|
import {
|
|
26
26
|
MC_PER_CREDIT,
|
|
27
27
|
PAA_BASE_CREDITS,
|
|
@@ -9577,7 +9577,7 @@ function registerPaaExtractorMcpTools(server, executor, options = {}) {
|
|
|
9577
9577
|
}, async (input) => formatSearchSerp(await executor.searchSerp(input), input));
|
|
9578
9578
|
server.registerTool("extract_url", {
|
|
9579
9579
|
title: "Single URL Extract",
|
|
9580
|
-
description: "Extract structured data from one public URL: content, schema, headings, metadata, screenshots, branding, featured image, or media assets. Wayback replay URLs automatically return the archived page copy without playback chrome. Use delivery:auto for bounded inline results with automatic artifact offload, delivery:artifact for a durable owner-scoped report, or delivery:memory to save the full page into hosted MCP Memory. preserveMedia is the preferred media-retention flag; depositToVault and downloadMedia remain temporary compatibility aliases.",
|
|
9580
|
+
description: "Extract structured data from one public URL: content, schema, headings, metadata, screenshots, branding, featured image, or media assets. Tries a plain HTTP fetch first and only opens a stealth browser when that fetch is blocked, hits a bot check, or returns unusably thin content \u2014 most calls never need the browser, and a bot_check_unresolved error means the browser already tried and failed, not that a retry will help. Wayback replay URLs automatically return the archived page copy without playback chrome. Use delivery:auto for bounded inline results with automatic artifact offload, delivery:artifact for a durable owner-scoped report, or delivery:memory to save the full page into hosted MCP Memory. preserveMedia is the preferred media-retention flag; depositToVault and downloadMedia remain temporary compatibility aliases.",
|
|
9581
9581
|
inputSchema: exposesLocalNetworkAccess ? ExtractUrlLocalInputSchema : ExtractUrlInputSchema,
|
|
9582
9582
|
outputSchema: recordOutputSchema("extract_url", ExtractUrlOutputSchema),
|
|
9583
9583
|
annotations: { ...liveWebToolAnnotations("Single URL Extract"), readOnlyHint: false }
|
|
@@ -9605,7 +9605,7 @@ function registerPaaExtractorMcpTools(server, executor, options = {}) {
|
|
|
9605
9605
|
}, async (input) => formatMapWaybackSnapshots(await executor.mapWaybackSnapshots(input), input, ctx));
|
|
9606
9606
|
server.registerTool("extract_site", {
|
|
9607
9607
|
title: "Multi-Page Site Content Crawl",
|
|
9608
|
-
description: `Crawl a public website and return page CONTENT (Markdown) across multiple pages. A Wayback replay URL produces one archived site snapshot. The optional wayback plan produces whole-site, single-page, or selected-page timelines across explicit months or a month range, all in one export with a capture matrix. Pass a new idempotencyKey for each intended crawl and reuse it only when retrying that call. Every MCP crawl starts a durable export; poll check_site_export for honest outcome counters and ${fileBehavior("the saved ZIP.", "the owner-scoped downloadable ZIP.")} Content only \u2014 for a technical SEO audit use audit_site instead.`,
|
|
9608
|
+
description: `Crawl a public website and return page CONTENT (Markdown) across multiple pages. Each page is fetched with a plain HTTP request first; a browser retries only the pages that were blocked or hit a bot check, so a per-page bot_check_unresolved result means that specific page already failed the browser retry too. A Wayback replay URL produces one archived site snapshot. The optional wayback plan produces whole-site, single-page, or selected-page timelines across explicit months or a month range, all in one export with a capture matrix. Pass a new idempotencyKey for each intended crawl and reuse it only when retrying that call. Every MCP crawl starts a durable export; poll check_site_export for honest outcome counters and ${fileBehavior("the saved ZIP.", "the owner-scoped downloadable ZIP.")} Content only \u2014 for a technical SEO audit use audit_site instead.`,
|
|
9609
9609
|
inputSchema: ExtractSiteInputSchema,
|
|
9610
9610
|
outputSchema: recordOutputSchema("extract_site", ExtractSiteOutputSchema),
|
|
9611
9611
|
annotations: { ...liveWebToolAnnotations("Multi-Page Site Content Crawl"), readOnlyHint: false }
|
|
@@ -16008,4 +16008,4 @@ export {
|
|
|
16008
16008
|
ScheduledResultsMcpExecutor,
|
|
16009
16009
|
registerScheduledResultsMcpTools
|
|
16010
16010
|
};
|
|
16011
|
-
//# sourceMappingURL=chunk-
|
|
16011
|
+
//# sourceMappingURL=chunk-Y2RIYBYR.js.map
|