torusguard 2.1.0 → 2.1.2

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (152) hide show
  1. package/.torusguard/.manifest.json +55 -5
  2. package/.torusguard/core/__init__.py +146 -0
  3. package/.torusguard/core/agent_roles.py +104 -0
  4. package/.torusguard/core/ast_walker.py +283 -0
  5. package/.torusguard/core/authorization.py +218 -0
  6. package/.torusguard/core/browser_verifier.py +128 -0
  7. package/.torusguard/core/bundle.py +141 -0
  8. package/.torusguard/core/call_graph.py +184 -0
  9. package/.torusguard/core/clustering.py +275 -0
  10. package/.torusguard/core/confidence.py +120 -0
  11. package/.torusguard/core/cross_file_taint.py +101 -0
  12. package/.torusguard/core/exploit_checker.py +317 -0
  13. package/.torusguard/core/formatter.py +351 -0
  14. package/.torusguard/core/governance.py +210 -0
  15. package/.torusguard/core/identity.py +104 -0
  16. package/.torusguard/core/import_resolver.py +91 -0
  17. package/.torusguard/core/incremental.py +102 -0
  18. package/.torusguard/core/lifecycle.py +137 -0
  19. package/.torusguard/core/models.py +425 -0
  20. package/.torusguard/core/parallel.py +56 -0
  21. package/.torusguard/core/parser.py +202 -0
  22. package/.torusguard/core/rechecker.py +107 -0
  23. package/.torusguard/core/replay_trace.py +178 -0
  24. package/.torusguard/core/rules_registry.py +131 -0
  25. package/.torusguard/core/run_folder.py +60 -0
  26. package/.torusguard/core/run_manager.py +163 -0
  27. package/.torusguard/core/runtime_evidence.py +175 -0
  28. package/.torusguard/core/runtime_validator.py +246 -0
  29. package/.torusguard/core/safety_gate.py +139 -0
  30. package/.torusguard/core/sarif.py +189 -0
  31. package/.torusguard/core/stack_profiler.py +184 -0
  32. package/.torusguard/core/symbol_table.py +91 -0
  33. package/.torusguard/core/taint.py +133 -0
  34. package/.torusguard/core/taint_graph.py +235 -0
  35. package/.torusguard/core/taint_rules.py +268 -0
  36. package/.torusguard/core/v070_reporter.py +102 -0
  37. package/.torusguard/core/v070_workflow.py +339 -0
  38. package/.torusguard/core/v6_reporter.py +180 -0
  39. package/.torusguard/core/v6_workflow.py +221 -0
  40. package/.torusguard/core/watcher.py +58 -0
  41. package/.torusguard/custom_rules/README.md +24 -0
  42. package/.torusguard/rules/TG-INPUT-007-unvalidated-redirect.md +53 -0
  43. package/.torusguard/rules/TG-INPUT-008-insecure-deserialization.md +52 -0
  44. package/.torusguard/scripts/__pycache__/audit_runner.cpython-314.pyc +0 -0
  45. package/.torusguard/scripts/__pycache__/finding_scorer.cpython-314.pyc +0 -0
  46. package/.torusguard/scripts/__pycache__/manifest_builder.cpython-314.pyc +0 -0
  47. package/.torusguard/scripts/__pycache__/rules_sync.cpython-314.pyc +0 -0
  48. package/.torusguard/scripts/audit_runner.py +108 -10
  49. package/.torusguard/scripts/deliberation_tournament.py +136 -0
  50. package/.torusguard/scripts/finding_scorer.py +43 -13
  51. package/.torusguard/scripts/reachability_analyzer.py +165 -0
  52. package/.torusguard/scripts/skill_profiler.py +26 -0
  53. package/.torusguard/scripts/stride_generator.py +198 -0
  54. package/.torusguard/skills/torusguard/SKILL.md +6 -2
  55. package/.torusguard/skills/torusguard-audit/SKILL.md +109 -84
  56. package/.torusguard/skills/torusguard-review/SKILL.md +27 -0
  57. package/.torusguard/skills/torusguard-threatmodel/SKILL.md +27 -0
  58. package/.torusguard/workflows/audit.md +21 -17
  59. package/.torusguard/workflows/review.md +11 -0
  60. package/.torusguard/workflows/threatmodel.md +9 -0
  61. package/README.md +149 -48
  62. package/package.json +10 -2
  63. package/skills/torusguard/SKILL.md +6 -2
  64. package/skills/torusguard/__pycache__/bootstrap.cpython-314.pyc +0 -0
  65. package/skills/torusguard/bootstrap.py +3 -3
  66. package/skills/torusguard/payload/.manifest.json +56 -7
  67. package/skills/torusguard/payload/core/__init__.py +146 -0
  68. package/skills/torusguard/payload/core/agent_roles.py +104 -0
  69. package/skills/torusguard/payload/core/ast_walker.py +283 -0
  70. package/skills/torusguard/payload/core/authorization.py +218 -0
  71. package/skills/torusguard/payload/core/browser_verifier.py +128 -0
  72. package/skills/torusguard/payload/core/bundle.py +141 -0
  73. package/skills/torusguard/payload/core/call_graph.py +184 -0
  74. package/skills/torusguard/payload/core/clustering.py +275 -0
  75. package/skills/torusguard/payload/core/confidence.py +120 -0
  76. package/skills/torusguard/payload/core/cross_file_taint.py +101 -0
  77. package/skills/torusguard/payload/core/exploit_checker.py +317 -0
  78. package/skills/torusguard/payload/core/formatter.py +351 -0
  79. package/skills/torusguard/payload/core/governance.py +210 -0
  80. package/skills/torusguard/payload/core/identity.py +104 -0
  81. package/skills/torusguard/payload/core/import_resolver.py +91 -0
  82. package/skills/torusguard/payload/core/incremental.py +102 -0
  83. package/skills/torusguard/payload/core/lifecycle.py +137 -0
  84. package/skills/torusguard/payload/core/models.py +425 -0
  85. package/skills/torusguard/payload/core/parallel.py +56 -0
  86. package/skills/torusguard/payload/core/parser.py +202 -0
  87. package/skills/torusguard/payload/core/rechecker.py +107 -0
  88. package/skills/torusguard/payload/core/replay_trace.py +178 -0
  89. package/skills/torusguard/payload/core/rules_registry.py +131 -0
  90. package/skills/torusguard/payload/core/run_folder.py +60 -0
  91. package/skills/torusguard/payload/core/run_manager.py +163 -0
  92. package/skills/torusguard/payload/core/runtime_evidence.py +175 -0
  93. package/skills/torusguard/payload/core/runtime_validator.py +246 -0
  94. package/skills/torusguard/payload/core/safety_gate.py +139 -0
  95. package/skills/torusguard/payload/core/sarif.py +189 -0
  96. package/skills/torusguard/payload/core/stack_profiler.py +184 -0
  97. package/skills/torusguard/payload/core/symbol_table.py +91 -0
  98. package/skills/torusguard/payload/core/taint.py +133 -0
  99. package/skills/torusguard/payload/core/taint_graph.py +235 -0
  100. package/skills/torusguard/payload/core/taint_rules.py +268 -0
  101. package/skills/torusguard/payload/core/v070_reporter.py +102 -0
  102. package/skills/torusguard/payload/core/v070_workflow.py +339 -0
  103. package/skills/torusguard/payload/core/v6_reporter.py +180 -0
  104. package/skills/torusguard/payload/core/v6_workflow.py +221 -0
  105. package/skills/torusguard/payload/core/watcher.py +58 -0
  106. package/skills/torusguard/payload/custom_rules/README.md +15 -0
  107. package/skills/torusguard/payload/rules/TG-INPUT-007-unvalidated-redirect.md +53 -0
  108. package/skills/torusguard/payload/rules/TG-INPUT-008-insecure-deserialization.md +52 -0
  109. package/skills/torusguard/payload/rules/container/TG-CONT-001-root-user-execution.md +50 -50
  110. package/skills/torusguard/payload/rules/container/TG-CONT-002-docker-socket-mount.md +47 -47
  111. package/skills/torusguard/payload/rules/container/TG-CONT-003-privileged-container-mode.md +53 -53
  112. package/skills/torusguard/payload/rules/container/TG-CONT-004-build-arg-secret-exposure.md +43 -43
  113. package/skills/torusguard/payload/rules/git/TG-GIT-001-historical-secret-in-git-commit.md +44 -44
  114. package/skills/torusguard/payload/rules/git/TG-GIT-002-plaintext-credentials-in-git-config.md +41 -41
  115. package/skills/torusguard/payload/rules/git/TG-GIT-003-sensitive-tracked-file-gitignore-breach.md +40 -40
  116. package/skills/torusguard/payload/rules/rag/TG-RAG-001-untrusted-rag-context-injection.md +72 -72
  117. package/skills/torusguard/payload/rules/rag/TG-RAG-002-autonomous-llm-tool-unsandboxed-call.md +51 -51
  118. package/skills/torusguard/payload/rules/rag/TG-RAG-003-unpartitioned-vector-tenant-lookup.md +51 -51
  119. package/skills/torusguard/payload/rules/redos/TG-REDOS-001-catastrophic-exponential-backtracking.md +46 -46
  120. package/skills/torusguard/payload/rules/redos/TG-REDOS-002-unbounded-nested-quantifier.md +43 -43
  121. package/skills/torusguard/payload/scripts/audit_runner.py +108 -10
  122. package/skills/torusguard/payload/scripts/deliberation_tournament.py +127 -0
  123. package/skills/torusguard/payload/scripts/finding_scorer.py +43 -13
  124. package/skills/torusguard/payload/scripts/reachability_analyzer.py +164 -0
  125. package/skills/torusguard/payload/scripts/stride_generator.py +196 -0
  126. package/skills/torusguard/payload/skills/torusguard/SKILL.md +6 -2
  127. package/skills/torusguard/payload/skills/torusguard/bootstrap.py +3 -3
  128. package/skills/torusguard/payload/skills/torusguard-ai-guard/SKILL.md +95 -95
  129. package/skills/torusguard/payload/skills/torusguard-audit/SKILL.md +109 -84
  130. package/skills/torusguard/payload/skills/torusguard-container/SKILL.md +94 -94
  131. package/skills/torusguard/payload/skills/torusguard-git-mine/SKILL.md +92 -92
  132. package/skills/torusguard/payload/skills/torusguard-ocr-scan/SKILL.md +94 -94
  133. package/skills/torusguard/payload/skills/torusguard-redos/SKILL.md +91 -91
  134. package/skills/torusguard/payload/skills/torusguard-review/SKILL.md +27 -0
  135. package/skills/torusguard/payload/skills/torusguard-threatmodel/SKILL.md +27 -0
  136. package/skills/torusguard/payload/workflows/ai-guard.md +31 -31
  137. package/skills/torusguard/payload/workflows/audit.md +21 -17
  138. package/skills/torusguard/payload/workflows/container.md +29 -29
  139. package/skills/torusguard/payload/workflows/git-mine.md +25 -25
  140. package/skills/torusguard/payload/workflows/ocr-scan.md +25 -25
  141. package/skills/torusguard/payload/workflows/redos.md +27 -27
  142. package/skills/torusguard/payload/workflows/review.md +11 -0
  143. package/skills/torusguard/payload/workflows/threatmodel.md +9 -0
  144. package/skills/torusguard/payload/workflows/torusguard-audit.md +35 -55
  145. package/skills/torusguard/references/csharp-security.md +41 -41
  146. package/skills/torusguard/references/go-security.md +41 -41
  147. package/skills/torusguard/references/java-security.md +40 -40
  148. package/skills/torusguard/references/polyglot-security-matrix.md +25 -25
  149. package/skills/torusguard/references/rust-security.md +40 -40
  150. package/skills/torusguard-audit/SKILL.md +107 -83
  151. package/skills/torusguard-review/SKILL.md +27 -0
  152. package/skills/torusguard-threatmodel/SKILL.md +27 -0
@@ -1,107 +1,122 @@
1
1
  ---
2
2
  name: torusguard-audit
3
- description: Static AST security scanning, line-shift invariant fingerprinting, root-cause clustering, and 0-100 confidence scoring via CLI or AI Agent.
4
- version: 2.0.0
3
+ description: Taint-aware static AST security scanning, cross-file interprocedural dataflow, 88 rules across 22 families, line-shift invariant fingerprinting, and 7-signal calibrated confidence scoring via CLI or AI Agent.
4
+ version: 2.1.1
5
5
  workflow: .torusguard/workflows/audit.md
6
6
  tools: Read, Grep, Glob, Write, run_command
7
7
  scripts-binding:
8
- - internal/scanner/scanner.go
9
- - cmd/torusguard/main.go
8
+ - .torusguard/scripts/audit_runner.py
9
+ - .torusguard/scripts/finding_scorer.py
10
+ - .torusguard/core/taint_graph.py
11
+ - .torusguard/core/cross_file_taint.py
12
+ - .torusguard/core/confidence.py
13
+ - .torusguard/core/parser.py
14
+ - .torusguard/core/incremental.py
15
+
10
16
  ---
11
17
 
12
- # TorusGuard Audit — Static Code Security Analysis
18
+ # TorusGuard Audit — Deep Taint-Aware Static Code Security Analysis
13
19
 
14
20
  ## Objective
15
- Execute static AST analysis across polyglot project files, evaluate code against 74 canonical security rules across 18 families, assign stable line-shift invariant fingerprints, cluster architectural root causes, synchronize findings with `security_report.md`, and score findings with auditable 0–100 confidence ratings.
21
+ Execute deep static analysis combining **Tree-sitter polyglot AST parsing**, **source-to-sink taint tracking**, and **interprocedural call-graph analysis** across Python, JavaScript/TypeScript, Go, Rust, Java, Ruby, PHP, and C#. Evaluates code against 88 canonical rules across 22 architectural families, generates line-shift invariant fingerprints, clusters systemic root causes, and computes 7-signal evidence-chain confidence ratings.
16
22
 
17
23
  ---
18
24
 
19
- ## Tri-Mode Execution
25
+ ## Tri-Mode Execution Parity
20
26
 
21
- ### Mode A: Automated CLI Execution
27
+ ### Mode A: Automated Terminal CLI
22
28
  Run the static security audit from your terminal:
23
29
  ```bash
24
- # Scan current repository
30
+ # Full codebase audit with taint dataflow analysis
25
31
  torusguard audit
26
32
 
27
- # Scan specific directory or example app
33
+ # Incremental scan (sub-second diff on changed files only)
34
+ torusguard audit --incremental
35
+
36
+ # Continuous watch mode (re-scan debounced on file save)
37
+ torusguard audit --watch
38
+
39
+ # Audit specific directory or microservice
28
40
  torusguard audit ./examples/vulnerable-react-express
29
41
 
30
- # Include test fixtures and spec directories
31
- torusguard audit --include-tests
42
+ # Export findings to OASIS SARIF v2.1.0 format
43
+ torusguard audit --sarif --sarif-out ./report.sarif
32
44
 
33
- # Output machine-readable JSON
45
+ # Machine-readable JSON output
34
46
  torusguard audit --json
35
47
  ```
36
- **Under the Hood:** Executes compiled Go static analysis engine (`internal/scanner`).
37
- - Auto-detects repository stack and skips build/cache directories (`node_modules`, `.git`, `.venv`, `dist`, `build`).
38
- - Evaluates files across 18 canonical security families:
39
- - `TG-SEC-*`: Hardcoded credentials, private keys, JWT secrets, client env leaks.
40
- - `TG-INPUT-*`: SQL injection, command injection, path traversal, unsafe HTML rendering.
41
- - `TG-DB-*`: Missing tenant isolation, service role keys in client code.
42
- - `TG-AUTH-*`: Plaintext passwords, missing cookie security flags (httpOnly, secure, sameSite).
43
- - `TG-PLATFORM-*`: Permissive wildcard CORS with credentials, missing security headers.
44
- - `TG-DIFF-*`: Disabled TLS verification (`verify=False`, `InsecureSkipVerify: true`).
45
- - `TG-NPE-*`: Null-pointer exceptions, unchecked nil error dereferences.
46
- - `TG-CONC-*`: Concurrency hazards, goroutine loop variable capture.
47
- - Writes findings directly to `security_report.md` at workspace root.
48
- - Displays standardized 75-column terminal cards.
49
-
50
- ### Mode B: In-Session AI Chat Agent Scan
51
- When auditing files directly in AI chat:
52
- 1. **Discover Sinks:** Use `grep_search` and `view_file` to search for dangerous patterns across server and client code.
53
- 2. **Cluster Root Causes:** Group findings by causal architecture (e.g. `cluster-tenant-isolation`, `cluster-credentials-exposure`, `cluster-injection`).
54
- 3. **Audit Evidence Sufficiency:** Ensure that user-controlled input reaches the vulnerable sink without prior sanitization or schema validation.
55
- 4. **Context Minimization (1/9th Token Strategy):** Inspect only bounded AST context windows ($\pm 3$ lines) via `scanner.ExtractContext` rather than ingesting entire files.
56
- 5. **Present Actionable Findings:** Display finding cards with severity, rule ID, file, line, and remediation recommendation.
57
- 6. **Prompt Next Phase:** Guide the operator to `/torusguard harden` or `torusguard harden`.
58
-
59
- ### Mode C: Native MCP Tool Execution
60
- For autonomous AI coding agents (Antigravity, Cursor, Windsurf, Claude Code):
61
- - **Tool Invocation:** Call `torusguard_audit` with target arguments:
62
- ```json
63
- {
64
- "target": ".",
65
- "include_ocr": true,
66
- "max_image_mb": 10
67
- }
68
- ```
69
- - **Programmatic Return:** Receives formatted finding summaries, active rule counts, and confirmation that `security_report.md` is updated on disk.
70
- - **Resource Companion:** Inspect the living report via resource `torusguard://security_report` or rules catalog via `torusguard://rules_catalog`.
48
+
49
+ ### Mode B: In-Session AI Chat Slash Command (`/torusguard audit`)
50
+ When executing audits directly in AI chat:
51
+ 1. **Trace Dataflow (Sources → Sinks):** Track user inputs (`request.GET`, `req.body`, `r.URL.Query()`) through assignments and helper functions to dangerous sinks (`execute()`, `innerHTML`, `open()`).
52
+ 2. **Verify Sanitizer Absence:** Confirm that input is not cleansed by `int()`, `escape()`, `shlex.quote()`, or `zod.safeParse()`.
53
+ 3. **Cross-File Correlation:** Trace calls across module boundaries up to 5 interprocedural hops using `CrossFileTaintAnalyzer`.
54
+ 4. **Cluster Root Causes:** Group findings into architectural failure patterns (e.g. `cluster-prompt-injection`, `cluster-tenant-isolation`, `cluster-supply-chain`).
55
+ 5. **Calibrate Confidence:** Score findings via the 7-signal evidence chain model.
56
+ 6. **Synchronize Ground Truth:** Record active findings in `security_report.md` at workspace root.
57
+
58
+ ### Mode C: Native MCP Tool Calling
59
+ MCP agents invoke `torusguard_audit(target_root, incremental, use_taint)` via JSON-RPC 2.0 stdio to receive structured findings with verified taint paths and confidence scores.
71
60
 
72
61
  ---
73
62
 
74
- ## Canonical Rule Families
75
- | Family | Scope | Example Violations |
63
+ ## Architectural Rule Taxonomy (86 Rules Across 22 Families)
64
+
65
+ | Family Code | Security Domain | Core Invariant Enforced |
76
66
  | :--- | :--- | :--- |
77
- | **TG-SEC** | Secrets & Credentials | Hardcoded JWT secret, API key strings, token logging |
78
- | **TG-INPUT** | Injection & Input Validation | Raw SQL interpolation, DOM `innerHTML`, `path.join` traversal |
79
- | **TG-DB** | Database & Tenant Scoping | Unscoped `.objects.get(id=...)`, Prisma missing `tenantId` |
80
- | **TG-AUTH** | Authentication & Cookies | Insecure cookies (missing httpOnly/secure/sameSite) |
81
- | **TG-PLATFORM** | Server & Platform Config | Wildcard CORS (`origin: '*'`) with credentials |
82
- | **TG-DIFF** | Security Bypasses | Disabled TLS verification (`verify=False`, `# nosec`) |
83
- | **TG-NPE** | Null Dereference / NPE | Unchecked optional chaining, unhandled nil error returns |
84
- | **TG-CONC** | Concurrency & Thread-Safety | Goroutine loop variable capture, unmutexed map mutations |
67
+ | **`TG-SEC`** | Secrets & Credentials | Zero hardcoded API keys, private certificates, or JWT secrets. |
68
+ | **`TG-AUTH`** | Authentication & Session | Enforce timing-safe compares, strong password hashing, algorithm verification. |
69
+ | **`TG-DB`** | Database & Tenancy | Parameterized SQL queries and tenant partition scoping across all lookups. |
70
+ | **`TG-INPUT`** | Input & Sanitization | Strict path sanitization, command argument escaping, safe template rendering. |
71
+ | **`TG-RATE`** | Rate Limiting | Rate-limiting middleware on auth endpoints and payload size bounds. |
72
+ | **`TG-AGENT`** | AI Agents & Prompts | Structural prompt isolation, inert XML delimiters, MCP tool schema validation. |
73
+ | **`TG-SSRF`** | Outbound Net & SSRF | Hostname whitelisting, private IP blocklist (127.0.0.1, 169.254.169.254). |
74
+ | **`TG-WEBHOOK`**| Webhook Verification | Cryptographic HMAC-SHA256 signature verification and replay prevention. |
75
+ | **`TG-WS`** | WebSockets | Origin verification, handshake authentication, inbound frame size limits. |
76
+ | **`TG-CSRF`** | CSRF Protection | SameSite cookie attributes and anti-CSRF token verification on state mutations. |
77
+ | **`TG-GQL`** | GraphQL Safety | Query depth limiting (max depth 6) and production schema introspection suppression. |
78
+ | **`TG-SUPPLY`** | Supply Chain & CI/CD | Immutable commit SHA pinning in GitHub Actions, lockfile integrity audits. |
79
+ | **`TG-BIZ`** | Business Logic | Non-negative quantity asserts, transaction locks, server-side discount bounds. |
80
+ | **`TG-CACHE`** | Cache Poisoning | Cache-Control headers on sensitive responses, unkeyed header sanitization. |
81
+ | **`TG-CLIENT`** | Client Bundle Secrets | Zero private environment variables (`process.env.SUPABASE_SERVICE_ROLE`) in client. |
82
+ | **`TG-PLATFORM`**| Platform Hardening | Helmet security headers, debug mode suppression, cookie secure flags. |
83
+ | **`TG-DIFF`** | Security Bypasses | Block `# nosec`, `InsecureSkipVerify`, and enforce Ponytail line budgets. |
84
+ | **`TG-EDGE`** | Edge & Serverless | Subrequest fan-out limits and serverless execution timeouts. |
85
+ | **`TG-CONT`** | Container Safety | Enforce non-root execution, zero docker socket mounts, no privileged mode. |
86
+ | **`TG-GIT`** | Git History Secrets | Zero historical committed credentials, no tokens in remote URLs. |
87
+ | **`TG-REDOS`** | ReDoS Prevention | Zero nested quantifiers `(a+)+` or catastrophic backtracking regular expressions. |
88
+ | **`TG-RAG`** | RAG & Vector DB | Mandatory tenant scoping on vector similarity search and inert ingestion. |
85
89
 
86
90
  ---
87
91
 
88
- ## Output Card Format
89
- ```markdown
90
- ### 🛡️ TorusGuard Static Security Audit Completed
91
- - **Run ID:** `run-20260910-121618-audit`
92
- - **Scope:** 7 files evaluated across 18 canonical families
93
- - **Status:** ✖ CRITICAL FINDINGS DETECTED
94
- - **Findings:** 2 Critical, 2 High, 2 Medium/Low (6 total)
95
- - **Clusters:** 3 architectural root causes identified
96
- - **Artifacts:** `security_report.md`
97
- - **Next Action:** Run `torusguard harden` or `/torusguard harden`
92
+ ## Evidence-Chain Confidence Scoring (0–100)
93
+
94
+ Findings are evaluated against 7 empirical signals:
95
+
96
+ ```
97
+ Final Score = Σ weighted signals:
98
+ - rule_severity_base (0.20): Critical=90, High=75, Medium=50, Low=25
99
+ - taint_path_confirmed (0.25): 100 if source→sink reachability is confirmed, 0 otherwise
100
+ - taint_depth (0.10): direct=100, 1-hop=80, 2-hop=60, 3+=40
101
+ - sanitizer_absence (0.15): 100 if no known sanitizer present, 0 if sanitized
102
+ - framework_context_match (0.10): 100 if sink matches detected stack, 50 default
103
+ - evidence_snippet_quality (0.10): 100 for multi-line AST context, 50 for single line
104
+ - test_fixture_penalty (-0.10): -50 penalty if located in test suite or fixtures
105
+ - memory_boost: -30 (false positive class) to +15 (regression watch)
106
+
107
+ Classification Bands:
108
+ - 90–100: Confirmed (High priority for automated Ponytail hardening)
109
+ - 70–89: High Confidence (Requires review & remediation)
110
+ - 50–69: Medium Confidence (Context verification needed)
111
+ - 0–49: Needs Review (Suppressed or test fixture)
98
112
  ```
99
113
 
100
114
  ---
101
115
 
102
- ## 🏛️ OpenCodeReview Precision & Context Minimization
103
- - **1/9th Token Minimization:** Use `scanner.ExtractContext` to extract only the bounded $\pm 3$ lines context window instead of ingesting entire files.
104
- - **Line-Level Pinning:** Every finding is reported with exact 1-indexed line numbers, line content, severity, and suggested remediation.
116
+ ## 🏛️ Context Minimization & Performance Invariants
117
+ - **1/9th Token Strategy:** Never ingest entire files into agent context. Always inspect bounded AST context windows ($\pm 3$ lines) via `scanner.ExtractContext`.
118
+ - **Incremental Cache:** Uses cryptographic content hashes in `.torusguard/cache/ast_cache.json` to complete repeat audits in under 1 second.
119
+ - **Fail-Closed Safety:** Incomplete parses or syntax anomalies in non-standard files gracefully degrade to fallback token parsing without aborting the audit.
105
120
 
106
121
  ---
107
122
 
@@ -109,28 +124,38 @@ For autonomous AI coding agents (Antigravity, Cursor, Windsurf, Claude Code):
109
124
 
110
125
  | Pattern | What AI Does Wrong | What Is Actually Correct |
111
126
  | :--- | :--- | :--- |
112
- | **Unbounded File Reading** | Reads entire 800+ line files to diagnose a 1-line vulnerability. | Read only the bounded context ($\pm 3$ lines) around the finding's line number. |
113
- | **Ignoring NPE / Concurrency** | Focuses only on secrets and misses thread-safety and null-pointer hazards. | Enforce `TG-NPE-001` and `TG-CONC-001` checks during audit review. |
114
- | **False Positive Escalation** | Flags documentation strings or mock test fixtures as production vulnerabilities. | Skip test files (`*_test.go`, `.test.ts`) and verify sink exploitability before reporting. |
115
- | **Missing Sync to Ground Truth** | Produces analysis in chat without checking or updating `security_report.md`. | Always reconcile against `security_report.md` at workspace root. |
127
+ | **Grepping Without Taint** | Flags `db.execute(query)` even when `query` is hardcoded or parameterized. | Verify user input reaches the sink via `core.taint_graph` before reporting. |
128
+ | **Ignoring Sanitizers** | Reports injection even though `int(user_id)` or `shlex.quote()` cleans the input. | Check if any node in the dataflow path acts as a registered sanitizer. |
129
+ | **Single-File Blindness** | Misses vulnerabilities when input enters `utils.py` and reaches a sink in `views.py`. | Trace interprocedural call chains using `CrossFileTaintAnalyzer` (up to 5 hops). |
130
+ | **Unbounded File Dumps** | Reads entire 1,000-line source files into chat context. | Read only the bounded context ($\pm 3$ lines) around the finding's line number. |
131
+ | **Missing Ground Truth Sync** | Produces analysis in chat without synchronizing `security_report.md`. | Always update `security_report.md` with active findings and run IDs. |
116
132
 
117
133
  ---
118
134
 
119
135
  ## ✅ Pre-Flight Self-Audit
120
136
 
121
- Before completing an audit pass, verify:
122
- - [ ] Did I run `torusguard audit` or inspect `security_report.md` first?
123
- - [ ] Are all reported findings pinned to precise line numbers?
124
- - [ ] Did I verify user input reaches the sink without prior validation?
125
- - [ ] Did I extract only the minimal AST context window ($\pm 3$ lines) to conserve tokens?
126
- - [ ] Did I evaluate against all 18 families including NPE and concurrency rules?
137
+ Before finishing an audit pass, confirm:
138
+ - [ ] Did I run `torusguard audit` or inspect `security_report.md`?
139
+ - [ ] Are all reported findings backed by confirmed taint paths or verified regex patterns?
140
+ - [ ] Did I verify user input reaches the sink without prior sanitization?
141
+ - [ ] Are findings pinned to exact 1-indexed line numbers with stable region hashes?
142
+ - [ ] Did I synchronize discovering state into `security_report.md`?
127
143
 
128
144
  ---
129
145
 
130
146
  ## 🔁 VBC Protocol (Verify → Build → Confirm)
131
147
 
132
148
  ```
133
- VERIFY: Scan source code and image assets using torusguard audit or torusguard_audit MCP tool.
134
- BUILD: Synthesize findings clustered by root cause with exact line numbers and bounded AST snippets.
135
- CONFIRM: Synchronize living findings into security_report.md and guide operator to /torusguard harden.
149
+ VERIFY: Scan source code with polyglot AST parser and trace dataflow reachability from sources to sinks.
150
+ BUILD: Group findings by root cause, compute 7-signal calibrated confidence scores, and format 75-column terminal cards.
151
+ CONFIRM: Synchronize all findings to security_report.md at workspace root and guide operator to /torusguard harden.
136
152
  ```
153
+
154
+ ---
155
+
156
+ ## 🔄 Rollback Defaults
157
+
158
+ If audit data becomes corrupted or a run needs to be reverted:
159
+ 1. Historical runs are preserved immutably in `.torusguard/runs/<run_id>/`.
160
+ 2. AST cache can be cleared anytime by deleting `.torusguard/cache/ast_cache.json`.
161
+ 3. Pre-apply code snapshots remain intact in `.torusguard/snapshots/`.
@@ -0,0 +1,27 @@
1
+ ---
2
+ name: torusguard-review
3
+ description: Differential PR and Git diff incremental security review — detects newly introduced vulnerabilities between Git branches or commits with net security score deltas.
4
+ tools: Read, Grep, Glob, Bash, Edit, Write
5
+ version: 2.1.2
6
+ last-updated: 2026-09-30
7
+ skills:
8
+ - torusguard
9
+ - torusguard-audit
10
+ ---
11
+
12
+ # TorusGuard Review — Differential Incremental PR Review
13
+
14
+ TorusGuard Review performs rapid, surgical security review on Git changes (e.g. against `HEAD~1`, `main`, or PR target branch), analyzing only modified lines to prevent introducing security regressions.
15
+
16
+ ## Invocation
17
+
18
+ - **Mode A (Terminal CLI):** `torusguard review [--diff <ref>]`
19
+ - **Mode B (AI Chat Slash Command):** `/torusguard review` or `/torusguard-review`
20
+ - **Mode C (Native MCP Tool):** `torusguard_review(target=".", diff_ref="HEAD~1")`
21
+
22
+ ## Core Capabilities
23
+
24
+ 1. **Sub-300ms Turnaround:** Only parses and checks modified/added lines in Git diffs.
25
+ 2. **Net Security Score Delta:** Tracks introduced vs. resolved vulnerabilities (`+0 Introduced, -2 Resolved`).
26
+ 3. **PR Gate Decision:** Blocks CI/CD pipelines if Critical or High severity flaws are introduced.
27
+ 4. **Alibaba OpenCodeReview Parity:** Outputs ready-to-post pull request inline review comments.
@@ -0,0 +1,27 @@
1
+ ---
2
+ name: torusguard-threatmodel
3
+ description: Synthesizes an architectural STRIDE threat model, interactive Mermaid Data Flow Diagram (DFD), and trust boundary register into SECURITY_THREAT_MODEL.md.
4
+ tools: Read, Grep, Glob, Bash, Edit, Write
5
+ version: 2.1.2
6
+ last-updated: 2026-09-30
7
+ skills:
8
+ - torusguard
9
+ - torusguard-audit
10
+ ---
11
+
12
+ # TorusGuard Threat Model — STRIDE & DFD Synthesizer
13
+
14
+ TorusGuard Threat Model automatically maps routes, trust boundaries, datastores, and egress sinks to construct an architectural Data Flow Diagram and formal STRIDE threat model.
15
+
16
+ ## Invocation
17
+
18
+ - **Mode A (Terminal CLI):** `torusguard threatmodel`
19
+ - **Mode B (AI Chat Slash Command):** `/torusguard threatmodel` or `/torusguard-threatmodel`
20
+ - **Mode C (Native MCP Tool):** `torusguard_threatmodel(target=".")`
21
+
22
+ ## Core Capabilities
23
+
24
+ 1. **Automated DFD Generation:** Renders interactive Mermaid diagram showing Zones 1–4 (Clients, Gateways, Core Services, Datastores).
25
+ 2. **STRIDE Risk Ledger:** Evaluates threats across Spoofing, Tampering, Repudiation, Information Disclosure, Denial of Service, and Elevation of Privilege.
26
+ 3. **Multi-Modal Vision Integration:** Discovers architectural diagrams and verifies that visual trust boundaries match code execution paths.
27
+ 4. **Living Compliance Artifact:** Emits `SECURITY_THREAT_MODEL.md` at workspace root.
@@ -1,5 +1,5 @@
1
1
  ---
2
- description: Static security AST scanning, stable line-shift invariant fingerprinting, root-cause clustering, and 0-100 confidence scoring.
2
+ description: Taint-aware static security AST scanning, cross-file interprocedural dataflow, 86+ rules across 22 families, and 7-signal calibrated confidence scoring.
3
3
  tools: Read, Grep, Glob, Bash, Write, run_command
4
4
  version: 2.0.0
5
5
  agent: auditor
@@ -7,19 +7,21 @@ lifecycle-phase: Phase 2 (Static Audit & Clustering)
7
7
  required-skills:
8
8
  - torusguard-audit
9
9
  scripts-binding:
10
- - internal/scanner/scanner.go
11
- - cmd/torusguard/main.go
12
- - cmd/torusguard/mcp.go
10
+ - .torusguard/scripts/audit_runner.py
11
+ - .torusguard/scripts/finding_scorer.py
12
+ - .torusguard/core/cross_file_taint.py
13
+ - .torusguard/core/confidence.py
14
+ - .torusguard/core/incremental.py
13
15
  ---
14
16
 
15
- # /torusguard audit — Static Security Code Scan & Clustering
17
+ # /torusguard audit — Deep Taint-Aware Static Security Code Scan & Clustering
16
18
 
17
19
  $ARGUMENTS
18
20
 
19
21
  ---
20
22
 
21
23
  ## Objective
22
- Execute static AST analysis across polyglot project files, evaluate code against 74 canonical security rules across 18 families, assign stable line-shift invariant fingerprints, cluster architectural root causes, synchronize findings with `security_report.md`, and score findings with auditable 0–100 confidence ratings.
24
+ Execute deep taint-aware static AST analysis across polyglot project files, evaluate code against 86+ canonical security rules across 22 families, trace interprocedural source-to-sink dataflows, assign stable line-shift invariant fingerprints, cluster architectural root causes, synchronize findings with `security_report.md`, and score findings with auditable 7-signal evidence-chain confidence ratings.
23
25
 
24
26
  ---
25
27
 
@@ -27,7 +29,7 @@ Execute static AST analysis across polyglot project files, evaluate code against
27
29
 
28
30
  | Mode | Command / Tool | Execution Method |
29
31
  | :--- | :--- | :--- |
30
- | **Mode A: Terminal CLI** | `torusguard audit [--include-tests] [--json]` | Shell execution of compiled Go binary. |
32
+ | **Mode A: Terminal CLI** | `torusguard audit [--incremental] [--watch] [--json]` | Shell execution of TorusGuard audit engine. |
31
33
  | **Mode B: AI Chat Slash** | `/torusguard audit` | Conversational guided audit with bounded context inspection. |
32
34
  | **Mode C: Native MCP Tool** | `torusguard_audit` | Autonomous agent tool invocation via stdio JSON-RPC. |
33
35
 
@@ -37,7 +39,7 @@ Execute static AST analysis across polyglot project files, evaluate code against
37
39
 
38
40
  Inspect workspace prerequisites before launching static audit:
39
41
  1. **Init State (`torusguard.json`):** Assert repository is initialized or run `torusguard init`.
40
- 2. **Active Rules (`rules/active/`):** Confirm rule definitions exist.
42
+ 2. **Active Rules (`rules/`):** Confirm rule definitions exist across the 22 architectural families.
41
43
  3. **Exclusions:** Assert `node_modules/`, `.venv/`, `dist/`, `.git/` are skipped.
42
44
  4. **Syntax Check:** Check for syntax errors before parsing ASTs.
43
45
 
@@ -52,13 +54,14 @@ Inspect workspace prerequisites before launching static audit:
52
54
  ## Execution Steps
53
55
 
54
56
  1. **Launch Audit Scan:**
55
- - **Mode A (CLI):** Run `torusguard audit` in terminal.
56
- - **Mode B (Chat):** Parse AST sinks using `grep_search` and bounded `ExtractContext`.
57
+ - **Mode A (CLI):** Run `torusguard audit` (or `python .torusguard/scripts/audit_runner.py`). Use `--incremental` for fast diffs.
58
+ - **Mode B (Chat):** Parse AST sinks and trace sources to sinks using `core.taint_graph` and bounded context windows.
57
59
  - **Mode C (MCP):** Call `torusguard_audit` with `{"target": "."}`.
58
- 2. **Scan Codebase ASTs:** Match 74 canonical rules across 18 families against source trees.
59
- 3. **Compute Stable Fingerprints:** Hash AST context to produce stable line-shift invariant IDs.
60
- 4. **Cluster Root Causes:** Group findings sharing identical sinks or causal architecture.
61
- 5. **Synchronize Ground Truth:** Update `security_report.md` at workspace root.
60
+ 2. **Scan Codebase ASTs:** Match 86+ canonical rules across 22 families against source trees.
61
+ 3. **Trace Taint Dataflows:** Follow untrusted inputs from sources to dangerous sinks across module boundaries.
62
+ 4. **Compute Stable Fingerprints:** Hash AST context to produce stable line-shift invariant IDs (`TG-XXX-<hash12>`).
63
+ 5. **Cluster Root Causes:** Group findings into the 22 canonical architectural root causes.
64
+ 6. **Synchronize Ground Truth:** Update `security_report.md` at workspace root.
62
65
 
63
66
  ---
64
67
 
@@ -66,10 +69,11 @@ Inspect workspace prerequisites before launching static audit:
66
69
 
67
70
  ```markdown
68
71
  ### 🔎 TorusGuard Static Audit Results
69
- - **Files Scanned:** [Count] source files
72
+ - **Files Scanned:** [Count] source files ([Changed] changed)
73
+ - **Taint Paths Confirmed:** [Count] verified dataflows
70
74
  - **Total Findings:** [Count] ([Critical] Critical, [High] High)
71
75
  - **Root Cause Clusters:** [Count] architectural issues
72
- - **Confidence:** [Score]/100
76
+ - **Confidence:** [Score]/100 (Evidence-Chain Calibrated)
73
77
  - **Living Report:** `security_report.md`
74
78
  ```
75
79
 
@@ -78,4 +82,4 @@ Inspect workspace prerequisites before launching static audit:
78
82
  ## Next Steps
79
83
 
80
84
  1. Run `/torusguard verify` to validate exploitability and review evidence.
81
- 2. Run `/torusguard harden` to generate minimal surgical remediation plans.
85
+ 2. Run `/torusguard harden` to generate minimal surgical remediation plans conforming to the Ponytail protocol (<= 35 additions, <= 25 deletions).
@@ -0,0 +1,11 @@
1
+ # TorusGuard Review Workflow (`/torusguard review`)
2
+
3
+ ## Purpose
4
+ Executes differential incremental Git diff security review against a reference branch or commit (default: `HEAD~1`).
5
+ Inspects changed lines for new security regressions, evaluates net security score deltas, and outputs pull request review comments.
6
+
7
+ ## Steps
8
+ 1. Determine reference branch or commit (e.g. `HEAD~1` or `origin/main`).
9
+ 2. Run differential scan via `torusguard review [--diff <ref>]` or native MCP tool `torusguard_review`.
10
+ 3. Check PR Gate Decision (`PASSED` or `BLOCKED`).
11
+ 4. If blocked, review newly introduced violations and formulate surgical Ponytail fixes via `/torusguard harden`.
@@ -0,0 +1,9 @@
1
+ # TorusGuard Threat Model Workflow (`/torusguard threatmodel`)
2
+
3
+ ## Purpose
4
+ Synthesizes architectural STRIDE threat models, interactive Mermaid Data Flow Diagrams (DFDs), and trust boundary registers directly into `SECURITY_THREAT_MODEL.md`.
5
+
6
+ ## Steps
7
+ 1. Execute `torusguard threatmodel` or native MCP tool `torusguard_threatmodel`.
8
+ 2. Inspect `SECURITY_THREAT_MODEL.md` to verify identified entrypoints, data stores, and trust boundaries.
9
+ 3. Align STRIDE risk controls against corresponding TorusGuard security invariants (`TG-AUTH`, `TG-DB`, `TG-RATE`, `TG-SSRF`).