mcp-scraper 0.33.7 → 0.34.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (42) hide show
  1. package/dist/bin/api-server.cjs +10063 -10012
  2. package/dist/bin/api-server.cjs.map +1 -1
  3. package/dist/bin/api-server.js +2 -2
  4. package/dist/bin/mcp-scraper-cli.cjs +2 -1
  5. package/dist/bin/mcp-scraper-cli.cjs.map +1 -1
  6. package/dist/bin/mcp-scraper-cli.js +3 -3
  7. package/dist/bin/mcp-scraper-install.cjs +1 -1
  8. package/dist/bin/mcp-scraper-install.cjs.map +1 -1
  9. package/dist/bin/mcp-scraper-install.js +1 -1
  10. package/dist/bin/mcp-stdio-server.cjs +15 -8
  11. package/dist/bin/mcp-stdio-server.cjs.map +1 -1
  12. package/dist/bin/mcp-stdio-server.js +3 -3
  13. package/dist/bin/paa-harvest.cjs +22 -0
  14. package/dist/bin/paa-harvest.cjs.map +1 -1
  15. package/dist/bin/paa-harvest.js +2 -2
  16. package/dist/{chunk-XVCETVYY.js → chunk-4KQVIYIH.js} +11 -3
  17. package/dist/chunk-4KQVIYIH.js.map +1 -0
  18. package/dist/{chunk-V73MPRU6.js → chunk-CB5C3BPB.js} +17 -1
  19. package/dist/chunk-CB5C3BPB.js.map +1 -0
  20. package/dist/{chunk-H6QECE2U.js → chunk-J32XSMJI.js} +16 -10
  21. package/dist/chunk-J32XSMJI.js.map +1 -0
  22. package/dist/chunk-NFMIHEYE.js +7 -0
  23. package/dist/chunk-NFMIHEYE.js.map +1 -0
  24. package/dist/{chunk-AT2PFWHW.js → chunk-ZID3WQID.js} +2 -2
  25. package/dist/index.cjs +22 -0
  26. package/dist/index.cjs.map +1 -1
  27. package/dist/index.d.cts +3 -0
  28. package/dist/index.d.ts +3 -0
  29. package/dist/index.js +2 -2
  30. package/dist/{server-B35CQJB3.js → server-LKVPEQFE.js} +130 -108
  31. package/dist/server-LKVPEQFE.js.map +1 -0
  32. package/dist/{worker-3LYNVGLF.js → worker-ZZHSYI3Y.js} +3 -3
  33. package/docs/mcp-tool-manifest.generated.json +36 -8
  34. package/package.json +1 -1
  35. package/dist/chunk-H6QECE2U.js.map +0 -1
  36. package/dist/chunk-HAMN42SP.js +0 -7
  37. package/dist/chunk-HAMN42SP.js.map +0 -1
  38. package/dist/chunk-V73MPRU6.js.map +0 -1
  39. package/dist/chunk-XVCETVYY.js.map +0 -1
  40. package/dist/server-B35CQJB3.js.map +0 -1
  41. /package/dist/{chunk-AT2PFWHW.js.map → chunk-ZID3WQID.js.map} +0 -0
  42. /package/dist/{worker-3LYNVGLF.js.map → worker-ZZHSYI3Y.js.map} +0 -0
@@ -6,13 +6,13 @@ import {
6
6
  } from "./chunk-3HBPKR5G.js";
7
7
  import {
8
8
  harvest
9
- } from "./chunk-XVCETVYY.js";
9
+ } from "./chunk-4KQVIYIH.js";
10
10
  import {
11
11
  browserServiceApiKey,
12
12
  runWithCostContext
13
13
  } from "./chunk-V36LS5YV.js";
14
14
  import "./chunk-K443GQY5.js";
15
- import "./chunk-V73MPRU6.js";
15
+ import "./chunk-CB5C3BPB.js";
16
16
  import {
17
17
  MC_COSTS,
18
18
  serpActualCostMc
@@ -139,4 +139,4 @@ export {
139
139
  startWorker,
140
140
  tickOnce
141
141
  };
142
- //# sourceMappingURL=worker-3LYNVGLF.js.map
142
+ //# sourceMappingURL=worker-ZZHSYI3Y.js.map
@@ -1,5 +1,5 @@
1
1
  {
2
- "generatedAt": "2026-07-24T22:48:53.625Z",
2
+ "generatedAt": "2026-07-25T03:25:17.950Z",
3
3
  "generatedFrom": "dist/bin/mcp-stdio-server.js",
4
4
  "serverInfo": {
5
5
  "name": "mcp-scraper",
@@ -16727,7 +16727,7 @@
16727
16727
  {
16728
16728
  "name": "reddit_trending",
16729
16729
  "title": "Reddit Trending",
16730
- "description": "Discover the top Reddit conversations about a topic from the last 30 days: searches Reddit (optionally one subreddit), ranks threads by engagement (upvotes + 2x comments), scrapes the top ones, and extracts the real questions people asked. Each scraped thread takes ~30s, so keep maxThreads small; for a wide scan set includeComments:false to rank cheaply first, then read the winners with reddit_thread. Not for reading one known thread URL — use reddit_thread for that.",
16730
+ "description": "Discover the top Reddit conversations about a topic from the last week or month: finds relevant recent threads via a Google site:reddit.com search (optionally scoped to one subreddit), scrapes them for real upvotes, comments, and the questions people asked, and ranks by engagement (upvotes + 2x comments). Scraping runs in parallel across the discovered threads; set includeComments:false for a fast, cheap discovery-only sweep (relevant thread list, no engagement stats, no per-thread billing) and then read the ones you want with reddit_thread. Not for reading one known thread URL — use reddit_thread for that.",
16731
16731
  "inputSchema": {
16732
16732
  "type": "object",
16733
16733
  "properties": {
@@ -16741,17 +16741,26 @@
16741
16741
  "minLength": 1,
16742
16742
  "description": "Bare subreddit name to scope the scan to one community, e.g. \"SEO\" (no r/ prefix, no URL). Omit to scan all of Reddit."
16743
16743
  },
16744
+ "window": {
16745
+ "type": "string",
16746
+ "enum": [
16747
+ "week",
16748
+ "month"
16749
+ ],
16750
+ "default": "month",
16751
+ "description": "How recent the threads must be: \"week\" or \"month\" (default). Applied via a Google time filter over reddit.com, so it reflects genuine recency."
16752
+ },
16744
16753
  "maxThreads": {
16745
16754
  "type": "integer",
16746
16755
  "minimum": 1,
16747
- "maximum": 25,
16748
- "default": 5,
16749
- "description": "Top-ranked threads to keep. Default 5 — each scraped thread adds ~30s and the request has a hard ~300s ceiling, so raise this above 5 only with includeComments:false."
16756
+ "maximum": 40,
16757
+ "default": 20,
16758
+ "description": "How many discovered threads to scrape and rank. Default 20 (scrape-all). Each scraped thread is billed like reddit_thread + its comments, so lower this to cap cost; raise toward 40 for a wider sweep. Scraping runs in parallel and stops early if it nears the request time limit (partial:true in the response)."
16750
16759
  },
16751
16760
  "includeComments": {
16752
16761
  "type": "boolean",
16753
16762
  "default": true,
16754
- "description": "Scrape each ranked thread for its body and comments to extract questions. Set false for a fast, cheap discovery-only sweep (ranking + metadata only), then call reddit_thread on the winners."
16763
+ "description": "Scrape each discovered thread for real upvotes, comments, and the questions people asked, then rank by engagement. Set false for a fast, cheap discovery-only sweep — returns the discovered threads (title + url) in relevance order with NO engagement stats and NO per-thread billing, so you can then call reddit_thread on the ones you want."
16755
16764
  },
16756
16765
  "maxCommentsPerThread": {
16757
16766
  "type": "integer",
@@ -16877,7 +16886,14 @@
16877
16886
  "type": "integer",
16878
16887
  "minimum": 0
16879
16888
  },
16880
- "searchUrl": {
16889
+ "candidatesFound": {
16890
+ "type": "integer",
16891
+ "minimum": 0
16892
+ },
16893
+ "partial": {
16894
+ "type": "boolean"
16895
+ },
16896
+ "searchQuery": {
16881
16897
  "type": "string"
16882
16898
  }
16883
16899
  },
@@ -16889,7 +16905,9 @@
16889
16905
  "rankedThreads",
16890
16906
  "questions",
16891
16907
  "threadsScraped",
16892
- "searchUrl"
16908
+ "candidatesFound",
16909
+ "partial",
16910
+ "searchQuery"
16893
16911
  ],
16894
16912
  "additionalProperties": false,
16895
16913
  "$schema": "http://json-schema.org/draft-07/schema#"
@@ -17679,6 +17697,16 @@
17679
17697
  "maximum": 2,
17680
17698
  "default": 1,
17681
17699
  "description": "Number of result pages to fetch (1–2)."
17700
+ },
17701
+ "recency": {
17702
+ "type": "string",
17703
+ "enum": [
17704
+ "day",
17705
+ "week",
17706
+ "month",
17707
+ "year"
17708
+ ],
17709
+ "description": "Restrict results to a recent time window (Google \"past day/week/month/year\" filter). Omit for all-time. Useful for \"what is being said this week\" style queries; pairs well with a site: operator in the query."
17682
17710
  }
17683
17711
  },
17684
17712
  "required": [
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "mcp-scraper",
3
- "version": "0.33.7",
3
+ "version": "0.34.0",
4
4
  "description": "MCP server for MCP Scraper web intelligence tools",
5
5
  "type": "module",
6
6
  "main": "./dist/index.cjs",