mcp-scraper 0.61.0 → 0.62.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
@@ -13,7 +13,7 @@ import {
13
13
  registerScheduledResultsMcpTools,
14
14
  registerSerpIntelligenceCaptureTools,
15
15
  resolveDeploymentProfile
16
- } from "../chunk-CIKVUFC5.js";
16
+ } from "../chunk-Y2RIYBYR.js";
17
17
  import "../chunk-2K74LVV7.js";
18
18
  import "../chunk-SQMKUPA5.js";
19
19
  import "../chunk-57Y5QHI3.js";
@@ -22,7 +22,7 @@ import {
22
22
  } from "../chunk-IMBMDZUO.js";
23
23
  import {
24
24
  PACKAGE_VERSION
25
- } from "../chunk-E6QGM3ZA.js";
25
+ } from "../chunk-RAPIYV2X.js";
26
26
  import "../chunk-2U2ZDDSI.js";
27
27
  import "../chunk-B7QNUVBT.js";
28
28
  import "../chunk-6OHW5XZN.js";
@@ -0,0 +1,7 @@
1
+ // src/version.ts
2
+ var PACKAGE_VERSION = "0.62.0";
3
+
4
+ export {
5
+ PACKAGE_VERSION
6
+ };
7
+ //# sourceMappingURL=chunk-RAPIYV2X.js.map
@@ -0,0 +1 @@
1
+ {"version":3,"sources":["../src/version.ts"],"sourcesContent":["export const PACKAGE_VERSION = '0.62.0'\n"],"mappings":";AAAO,IAAM,kBAAkB;","names":[]}
@@ -21,7 +21,7 @@ import {
21
21
  } from "./chunk-57Y5QHI3.js";
22
22
  import {
23
23
  PACKAGE_VERSION
24
- } from "./chunk-E6QGM3ZA.js";
24
+ } from "./chunk-RAPIYV2X.js";
25
25
  import {
26
26
  MC_PER_CREDIT,
27
27
  PAA_BASE_CREDITS,
@@ -9577,7 +9577,7 @@ function registerPaaExtractorMcpTools(server, executor, options = {}) {
9577
9577
  }, async (input) => formatSearchSerp(await executor.searchSerp(input), input));
9578
9578
  server.registerTool("extract_url", {
9579
9579
  title: "Single URL Extract",
9580
- description: "Extract structured data from one public URL: content, schema, headings, metadata, screenshots, branding, featured image, or media assets. Wayback replay URLs automatically return the archived page copy without playback chrome. Use delivery:auto for bounded inline results with automatic artifact offload, delivery:artifact for a durable owner-scoped report, or delivery:memory to save the full page into hosted MCP Memory. preserveMedia is the preferred media-retention flag; depositToVault and downloadMedia remain temporary compatibility aliases.",
9580
+ description: "Extract structured data from one public URL: content, schema, headings, metadata, screenshots, branding, featured image, or media assets. Tries a plain HTTP fetch first and only opens a stealth browser when that fetch is blocked, hits a bot check, or returns unusably thin content \u2014 most calls never need the browser, and a bot_check_unresolved error means the browser already tried and failed, not that a retry will help. Wayback replay URLs automatically return the archived page copy without playback chrome. Use delivery:auto for bounded inline results with automatic artifact offload, delivery:artifact for a durable owner-scoped report, or delivery:memory to save the full page into hosted MCP Memory. preserveMedia is the preferred media-retention flag; depositToVault and downloadMedia remain temporary compatibility aliases.",
9581
9581
  inputSchema: exposesLocalNetworkAccess ? ExtractUrlLocalInputSchema : ExtractUrlInputSchema,
9582
9582
  outputSchema: recordOutputSchema("extract_url", ExtractUrlOutputSchema),
9583
9583
  annotations: { ...liveWebToolAnnotations("Single URL Extract"), readOnlyHint: false }
@@ -9605,7 +9605,7 @@ function registerPaaExtractorMcpTools(server, executor, options = {}) {
9605
9605
  }, async (input) => formatMapWaybackSnapshots(await executor.mapWaybackSnapshots(input), input, ctx));
9606
9606
  server.registerTool("extract_site", {
9607
9607
  title: "Multi-Page Site Content Crawl",
9608
- description: `Crawl a public website and return page CONTENT (Markdown) across multiple pages. A Wayback replay URL produces one archived site snapshot. The optional wayback plan produces whole-site, single-page, or selected-page timelines across explicit months or a month range, all in one export with a capture matrix. Pass a new idempotencyKey for each intended crawl and reuse it only when retrying that call. Every MCP crawl starts a durable export; poll check_site_export for honest outcome counters and ${fileBehavior("the saved ZIP.", "the owner-scoped downloadable ZIP.")} Content only \u2014 for a technical SEO audit use audit_site instead.`,
9608
+ description: `Crawl a public website and return page CONTENT (Markdown) across multiple pages. Each page is fetched with a plain HTTP request first; a browser retries only the pages that were blocked or hit a bot check, so a per-page bot_check_unresolved result means that specific page already failed the browser retry too. A Wayback replay URL produces one archived site snapshot. The optional wayback plan produces whole-site, single-page, or selected-page timelines across explicit months or a month range, all in one export with a capture matrix. Pass a new idempotencyKey for each intended crawl and reuse it only when retrying that call. Every MCP crawl starts a durable export; poll check_site_export for honest outcome counters and ${fileBehavior("the saved ZIP.", "the owner-scoped downloadable ZIP.")} Content only \u2014 for a technical SEO audit use audit_site instead.`,
9609
9609
  inputSchema: ExtractSiteInputSchema,
9610
9610
  outputSchema: recordOutputSchema("extract_site", ExtractSiteOutputSchema),
9611
9611
  annotations: { ...liveWebToolAnnotations("Multi-Page Site Content Crawl"), readOnlyHint: false }
@@ -16008,4 +16008,4 @@ export {
16008
16008
  ScheduledResultsMcpExecutor,
16009
16009
  registerScheduledResultsMcpTools
16010
16010
  };
16011
- //# sourceMappingURL=chunk-CIKVUFC5.js.map
16011
+ //# sourceMappingURL=chunk-Y2RIYBYR.js.map