mcp-scraper 0.3.1 → 0.3.3

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (51) hide show
  1. package/README.md +4 -1
  2. package/dist/bin/api-server.cjs +2050 -832
  3. package/dist/bin/api-server.cjs.map +1 -1
  4. package/dist/bin/api-server.js +2 -2
  5. package/dist/bin/browser-agent-stdio-server.cjs +1 -1
  6. package/dist/bin/browser-agent-stdio-server.cjs.map +1 -1
  7. package/dist/bin/browser-agent-stdio-server.js +2 -2
  8. package/dist/bin/mcp-scraper-cli.cjs +1 -1
  9. package/dist/bin/mcp-scraper-cli.cjs.map +1 -1
  10. package/dist/bin/mcp-scraper-cli.js +1 -1
  11. package/dist/bin/mcp-scraper-combined-stdio-server.cjs +351 -4
  12. package/dist/bin/mcp-scraper-combined-stdio-server.cjs.map +1 -1
  13. package/dist/bin/mcp-scraper-combined-stdio-server.js +4 -4
  14. package/dist/bin/mcp-scraper-install.cjs +3 -3
  15. package/dist/bin/mcp-scraper-install.cjs.map +1 -1
  16. package/dist/bin/mcp-scraper-install.js +2 -2
  17. package/dist/bin/mcp-stdio-server.cjs +349 -2
  18. package/dist/bin/mcp-stdio-server.cjs.map +1 -1
  19. package/dist/bin/mcp-stdio-server.js +2 -2
  20. package/dist/bin/paa-harvest.cjs +3 -1
  21. package/dist/bin/paa-harvest.cjs.map +1 -1
  22. package/dist/bin/paa-harvest.js +1 -1
  23. package/dist/{chunk-SXTXMFEQ.js → chunk-AZ5PKU4F.js} +34 -1
  24. package/dist/chunk-AZ5PKU4F.js.map +1 -0
  25. package/dist/{chunk-Q4STSM63.js → chunk-C3FGVJWH.js} +3 -3
  26. package/dist/chunk-C3FGVJWH.js.map +1 -0
  27. package/dist/{chunk-55T4SRLJ.js → chunk-MGWGZBL5.js} +350 -3
  28. package/dist/chunk-MGWGZBL5.js.map +1 -0
  29. package/dist/chunk-PVXDEREW.js +7 -0
  30. package/dist/chunk-PVXDEREW.js.map +1 -0
  31. package/dist/{chunk-LICHCMV6.js → chunk-RLBJ3QNC.js} +2 -2
  32. package/dist/{chunk-IPW4LFOT.js → chunk-UWSG3C5J.js} +4 -2
  33. package/dist/chunk-UWSG3C5J.js.map +1 -0
  34. package/dist/index.cjs +3 -1
  35. package/dist/index.cjs.map +1 -1
  36. package/dist/index.js +1 -1
  37. package/dist/{server-AXPNL2RV.js → server-D3XHIEQN.js} +997 -253
  38. package/dist/server-D3XHIEQN.js.map +1 -0
  39. package/dist/{worker-SLQ375UG.js → worker-56IXWOQU.js} +3 -3
  40. package/docs/mcp-tool-craft-lint.generated.md +5 -2
  41. package/docs/mcp-tool-manifest.generated.json +80 -5
  42. package/package.json +1 -1
  43. package/dist/chunk-55T4SRLJ.js.map +0 -1
  44. package/dist/chunk-D4JDGKOV.js +0 -7
  45. package/dist/chunk-D4JDGKOV.js.map +0 -1
  46. package/dist/chunk-IPW4LFOT.js.map +0 -1
  47. package/dist/chunk-Q4STSM63.js.map +0 -1
  48. package/dist/chunk-SXTXMFEQ.js.map +0 -1
  49. package/dist/server-AXPNL2RV.js.map +0 -1
  50. /package/dist/{chunk-LICHCMV6.js.map → chunk-RLBJ3QNC.js.map} +0 -0
  51. /package/dist/{worker-SLQ375UG.js.map → worker-56IXWOQU.js.map} +0 -0
@@ -4,7 +4,7 @@ import {
4
4
  createHarvestAttemptRecorder,
5
5
  harvestProblemResponse,
6
6
  serializeHarvestProblem
7
- } from "./chunk-SXTXMFEQ.js";
7
+ } from "./chunk-AZ5PKU4F.js";
8
8
  import {
9
9
  claimPendingJob,
10
10
  completeJob,
@@ -15,7 +15,7 @@ import {
15
15
  } from "./chunk-DBQDG7EH.js";
16
16
  import {
17
17
  harvest
18
- } from "./chunk-IPW4LFOT.js";
18
+ } from "./chunk-UWSG3C5J.js";
19
19
  import "./chunk-M2S27J6Z.js";
20
20
  import {
21
21
  browserServiceApiKey
@@ -125,4 +125,4 @@ export {
125
125
  startWorker,
126
126
  tickOnce
127
127
  };
128
- //# sourceMappingURL=worker-SLQ375UG.js.map
128
+ //# sourceMappingURL=worker-56IXWOQU.js.map
@@ -1,8 +1,8 @@
1
1
  # MCP Tool Craft Lint
2
2
 
3
- Generated: 2026-06-17T19:27:45.886Z
3
+ Generated: 2026-06-19T20:46:06.138Z
4
4
 
5
- Tools: 45
5
+ Tools: 48
6
6
  Checks per tool: 10
7
7
  Failing tools: 0
8
8
 
@@ -19,6 +19,8 @@ Failing tools: 0
19
19
  | `facebook_ad_search` | 10/10 | none |
20
20
  | `facebook_ad_transcribe` | 10/10 | none |
21
21
  | `facebook_video_transcribe` | 10/10 | none |
22
+ | `instagram_profile_content` | 10/10 | none |
23
+ | `instagram_media_download` | 10/10 | none |
22
24
  | `maps_place_intel` | 10/10 | none |
23
25
  | `maps_search` | 10/10 | none |
24
26
  | `directory_workflow` | 10/10 | none |
@@ -51,5 +53,6 @@ Failing tools: 0
51
53
  | `browser_replay_annotate` | 10/10 | none |
52
54
  | `browser_close` | 10/10 | none |
53
55
  | `browser_list_sessions` | 10/10 | none |
56
+ | `browser_capture_fanout` | 10/10 | none |
54
57
  | `capture_serp_snapshot` | 10/10 | none |
55
58
  | `capture_serp_page_snapshots` | 10/10 | none |
@@ -1,12 +1,12 @@
1
1
  {
2
- "generatedAt": "2026-06-18T02:50:57.263Z",
2
+ "generatedAt": "2026-06-19T20:46:06.135Z",
3
3
  "counts": {
4
- "main_stdio": 22,
4
+ "main_stdio": 24,
5
5
  "browser_agent_stdio": 22,
6
- "combined_stdio": 44,
7
- "hosted_http": 24,
6
+ "combined_stdio": 46,
7
+ "hosted_http": 26,
8
8
  "hosted_only": 2,
9
- "unique_public": 46
9
+ "unique_public": 48
10
10
  },
11
11
  "surfaces": {
12
12
  "main_stdio": [
@@ -21,6 +21,8 @@
21
21
  "facebook_ad_search",
22
22
  "facebook_ad_transcribe",
23
23
  "facebook_video_transcribe",
24
+ "instagram_profile_content",
25
+ "instagram_media_download",
24
26
  "maps_place_intel",
25
27
  "maps_search",
26
28
  "directory_workflow",
@@ -69,6 +71,8 @@
69
71
  "facebook_ad_search",
70
72
  "facebook_ad_transcribe",
71
73
  "facebook_video_transcribe",
74
+ "instagram_profile_content",
75
+ "instagram_media_download",
72
76
  "maps_place_intel",
73
77
  "maps_search",
74
78
  "directory_workflow",
@@ -115,6 +119,8 @@
115
119
  "facebook_ad_search",
116
120
  "facebook_ad_transcribe",
117
121
  "facebook_video_transcribe",
122
+ "instagram_profile_content",
123
+ "instagram_media_download",
118
124
  "maps_place_intel",
119
125
  "maps_search",
120
126
  "directory_workflow",
@@ -145,6 +151,8 @@
145
151
  "facebook_ad_search",
146
152
  "facebook_ad_transcribe",
147
153
  "facebook_video_transcribe",
154
+ "instagram_profile_content",
155
+ "instagram_media_download",
148
156
  "maps_place_intel",
149
157
  "maps_search",
150
158
  "directory_workflow",
@@ -521,6 +529,73 @@
521
529
  "facebook-public-video-url"
522
530
  ]
523
531
  },
532
+ {
533
+ "name": "instagram_profile_content",
534
+ "surfaces": [
535
+ "main_stdio",
536
+ "combined_stdio",
537
+ "hosted_http"
538
+ ],
539
+ "tokenBudgetClass": "specialized",
540
+ "hasOutputSchema": true,
541
+ "hasStructuredContent": true,
542
+ "writesLocalFiles": true,
543
+ "symptomTriggers": [
544
+ "the user wants to find public Instagram posts, reels, or tv URLs from a person, creator, brand, or profile"
545
+ ],
546
+ "samplePrompts": [
547
+ "Find the content list for @nasaartemis",
548
+ "Get public reels from this Instagram profile",
549
+ "Inventory this creator profile before downloading media"
550
+ ],
551
+ "formatExamples": [
552
+ "handle=\"@nasaartemis\"",
553
+ "url=\"https://www.instagram.com/nasaartemis/\", maxItems=50",
554
+ "handle=\"@creator\", browserProfile=\"work-accounts\", maxItems=200, maxScrolls=60"
555
+ ],
556
+ "negativeSpace": "This discovers profile links only; use instagram_media_download for one selected post/reel. Private or deep historical collection requires a local authenticated browser profile from browser_profile_import/sync.",
557
+ "producesHandlesFor": [
558
+ "instagram_media_download"
559
+ ],
560
+ "consumesHandles": [
561
+ "instagram-handle",
562
+ "instagram-profile-url"
563
+ ]
564
+ },
565
+ {
566
+ "name": "instagram_media_download",
567
+ "surfaces": [
568
+ "main_stdio",
569
+ "combined_stdio",
570
+ "hosted_http"
571
+ ],
572
+ "tokenBudgetClass": "specialized",
573
+ "hasOutputSchema": true,
574
+ "hasStructuredContent": true,
575
+ "writesLocalFiles": true,
576
+ "symptomTriggers": [
577
+ "the user wants to download text, image, reel video/audio tracks, muxed video, or transcript from one Instagram post or reel"
578
+ ],
579
+ "samplePrompts": [
580
+ "Download this Instagram reel and caption",
581
+ "Get the image and text from this Instagram post",
582
+ "Transcribe this Instagram reel"
583
+ ],
584
+ "formatExamples": [
585
+ "url=\"https://www.instagram.com/reel/SHORTCODE/\"",
586
+ "includeTranscript=true, downloadAllTracks=false",
587
+ "url=\"https://www.instagram.com/reel/SHORTCODE/\", browserProfile=\"work-accounts\""
588
+ ],
589
+ "negativeSpace": "Use instagram_profile_content first when the user has a profile but has not selected a post/reel. Private or login-gated media requires a local authenticated browser profile from browser_profile_import/sync.",
590
+ "producesHandlesFor": [
591
+ "local-instagram-media",
592
+ "instagram-cdn-url"
593
+ ],
594
+ "consumesHandles": [
595
+ "instagram-post-url",
596
+ "instagram-reel-url"
597
+ ]
598
+ },
524
599
  {
525
600
  "name": "maps_place_intel",
526
601
  "surfaces": [
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "mcp-scraper",
3
- "version": "0.3.1",
3
+ "version": "0.3.3",
4
4
  "description": "MCP server for MCP Scraper web intelligence tools",
5
5
  "type": "module",
6
6
  "main": "./dist/index.cjs",