adaptive-memory-multi-model-router 2.14.16 → 2.14.18
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/.a3m-vault.json +23 -0
- package/.github/workflows/ci.yml +253 -5
- package/.publish-tick +1 -1
- package/AGENT_COUNCIL_FINDINGS.md +142 -0
- package/LAUNCH_CHECKLIST.md +141 -0
- package/README.md +15 -17
- package/README.md.bak +836 -0
- package/articles/CHINESE_SUBMISSIONS_READY.md +322 -0
- package/articles/DEVTO_READY.md +255 -0
- package/articles/HN_POST_READY.md +137 -0
- package/articles/INDIEHACKERS_READY.md +120 -0
- package/articles/NEWSLETTER_SEND_NOW.md +259 -0
- package/articles/PRODUCTHUNT_READY.md +106 -0
- package/articles/REDDIT_SUBMISSION_READY.md +348 -0
- package/articles/TWEET_STORM_READY.md +165 -0
- package/benchmark-results.json +45 -43
- package/council-votes/architecture-vote.md +121 -0
- package/council-votes/coverage-vote.md +93 -0
- package/dist/cost/costTracker.d.ts +109 -44
- package/dist/cost/costTracker.js +321 -98
- package/dist/cost/costTracker.js.map +1 -1
- package/dist/ensemble.d.ts +21 -0
- package/dist/ensemble.js +85 -0
- package/dist/index.d.ts +9 -5
- package/dist/index.js +12 -4
- package/dist/routing/advancedRouter.d.ts +38 -43
- package/dist/routing/advancedRouter.js +394 -408
- package/dist/routing/advancedRouter.js.map +1 -1
- package/dist/routing/providers/providerConfig.d.ts +49 -0
- package/dist/routing/providers/providerConfig.js +883 -0
- package/dist/routing/routing/advancedRouter.d.ts +62 -0
- package/dist/routing/routing/advancedRouter.js +447 -0
- package/dist/routing/utils/tokenUtils.d.ts +52 -0
- package/dist/routing/utils/tokenUtils.js +129 -0
- package/dist/server/proxyServer.d.ts +1 -1
- package/dist/tui/dashboard.js +66 -2
- package/dist/tui/dashboard.js.map +1 -1
- package/dist/utils/tokenUtils.d.ts +48 -1
- package/dist/utils/tokenUtils.js +117 -4
- package/dist/utils/tokenUtils.js.map +1 -1
- package/docs/CITATIONS.md +2 -2
- package/docs/GEO_STATUS.md +43 -157
- package/docs/ai-plugin.json +4 -4
- package/docs/llms.txt +21 -27
- package/docs/sitemap.xml +14 -20
- package/package.json +2 -2
- package/research-log.md +49 -0
- package/sitemap.xml +57 -0
- package/src/cost/costTracker.ts +576 -0
- package/src/ensemble.ts +103 -0
- package/src/index.ts +13 -3
- package/src/routing/advancedRouter.ts +536 -0
- package/src/tui/dashboard.ts +76 -3
- package/src/utils/tokenUtils.ts +142 -4
- package/test-council/1-structure-tests.test.js +353 -0
- package/test-council/1-structure-tests.test.ts +353 -0
- package/test-council/2-edge-case-tests.test.ts +361 -0
- package/test-council/3-performance-tests.test.ts +669 -0
- package/test-council/4-integration-tests.test.ts +391 -0
- package/test-council/5-agent-council-eval.test.ts +413 -0
- package/test-council/AGENT_COUNCIL_ARCHITECTURE.md +349 -0
- package/test-council/TEST_COUNCIL_REPORT.md +201 -0
- package/test-council/agents/edge-case-agent.ts +363 -0
- package/test-council/agents/performance-agent.ts +426 -0
- package/test-council/agents/structure-agent.ts +227 -0
- package/test-council/council.md +183 -0
- package/tests/security/guardrailEngine.test.ts +700 -0
- package/docs/.well-known/ai-plugin.json +0 -16
- package/research/PUBLISH_LOG.md +0 -3
|
@@ -0,0 +1,426 @@
|
|
|
1
|
+
#!/usr/bin/env node
|
|
2
|
+
/**
|
|
3
|
+
* Performance Agent - Benchmarks and identifies performance issues
|
|
4
|
+
*
|
|
5
|
+
* This agent measures:
|
|
6
|
+
* - Token counting accuracy and speed
|
|
7
|
+
* - Cost estimation precision
|
|
8
|
+
* - Response time distributions
|
|
9
|
+
* - Memory usage patterns
|
|
10
|
+
* - Throughput under load
|
|
11
|
+
*/
|
|
12
|
+
|
|
13
|
+
import * as fs from 'fs';
|
|
14
|
+
import * as path from 'path';
|
|
15
|
+
|
|
16
|
+
interface BenchmarkResult {
|
|
17
|
+
name: string;
|
|
18
|
+
operations: number;
|
|
19
|
+
totalTimeMs: number;
|
|
20
|
+
avgMs: number;
|
|
21
|
+
minMs: number;
|
|
22
|
+
maxMs: number;
|
|
23
|
+
p50Ms: number;
|
|
24
|
+
p95Ms: number;
|
|
25
|
+
p99Ms: number;
|
|
26
|
+
opsPerSecond: number;
|
|
27
|
+
}
|
|
28
|
+
|
|
29
|
+
interface PerformanceReport {
|
|
30
|
+
benchmarks: BenchmarkResult[];
|
|
31
|
+
issues: PerformanceIssue[];
|
|
32
|
+
recommendations: string[];
|
|
33
|
+
}
|
|
34
|
+
|
|
35
|
+
interface PerformanceIssue {
|
|
36
|
+
name: string;
|
|
37
|
+
description: string;
|
|
38
|
+
severity: 'critical' | 'high' | 'medium' | 'low';
|
|
39
|
+
current: string;
|
|
40
|
+
expected: string;
|
|
41
|
+
}
|
|
42
|
+
|
|
43
|
+
interface PerformanceConfig {
|
|
44
|
+
iterations: number;
|
|
45
|
+
warmupRuns: number;
|
|
46
|
+
maxTimeMs: number;
|
|
47
|
+
}
|
|
48
|
+
|
|
49
|
+
// Default configuration
|
|
50
|
+
const DEFAULT_CONFIG: PerformanceConfig = {
|
|
51
|
+
iterations: 1000,
|
|
52
|
+
warmupRuns: 10,
|
|
53
|
+
maxTimeMs: 30000
|
|
54
|
+
};
|
|
55
|
+
|
|
56
|
+
// Import the modules to benchmark
|
|
57
|
+
function loadModules() {
|
|
58
|
+
try {
|
|
59
|
+
const index = require('../../dist/index.js');
|
|
60
|
+
const tokenUtils = require('../../dist/utils/tokenUtils.js');
|
|
61
|
+
return { index, tokenUtils };
|
|
62
|
+
} catch (e) {
|
|
63
|
+
// Try source if dist not available
|
|
64
|
+
try {
|
|
65
|
+
const tokenUtils = require('../../src/utils/tokenUtils.js');
|
|
66
|
+
return { index: {}, tokenUtils };
|
|
67
|
+
} catch (e2) {
|
|
68
|
+
return null;
|
|
69
|
+
}
|
|
70
|
+
}
|
|
71
|
+
}
|
|
72
|
+
|
|
73
|
+
// Run a benchmark
|
|
74
|
+
function runBenchmark(name: string, fn: () => void, iterations: number): BenchmarkResult {
|
|
75
|
+
const times: number[] = [];
|
|
76
|
+
|
|
77
|
+
// Warmup
|
|
78
|
+
for (let i = 0; i < 10; i++) {
|
|
79
|
+
fn();
|
|
80
|
+
}
|
|
81
|
+
|
|
82
|
+
// Actual benchmark
|
|
83
|
+
for (let i = 0; i < iterations; i++) {
|
|
84
|
+
const start = performance.now();
|
|
85
|
+
fn();
|
|
86
|
+
const end = performance.now();
|
|
87
|
+
times.push(end - start);
|
|
88
|
+
}
|
|
89
|
+
|
|
90
|
+
times.sort((a, b) => a - b);
|
|
91
|
+
|
|
92
|
+
const totalTimeMs = times.reduce((a, b) => a + b, 0);
|
|
93
|
+
|
|
94
|
+
return {
|
|
95
|
+
name,
|
|
96
|
+
operations: iterations,
|
|
97
|
+
totalTimeMs,
|
|
98
|
+
avgMs: totalTimeMs / iterations,
|
|
99
|
+
minMs: times[0],
|
|
100
|
+
maxMs: times[times.length - 1],
|
|
101
|
+
p50Ms: times[Math.floor(iterations * 0.50)],
|
|
102
|
+
p95Ms: times[Math.floor(iterations * 0.95)],
|
|
103
|
+
p99Ms: times[Math.floor(iterations * 0.99)],
|
|
104
|
+
opsPerSecond: (iterations / totalTimeMs) * 1000
|
|
105
|
+
};
|
|
106
|
+
}
|
|
107
|
+
|
|
108
|
+
// Token counting benchmarks
|
|
109
|
+
function benchmarkTokenCounting(iterations: number): BenchmarkResult[] {
|
|
110
|
+
const results: BenchmarkResult[] = [];
|
|
111
|
+
const modules = loadModules();
|
|
112
|
+
|
|
113
|
+
if (!modules || !modules.tokenUtils) {
|
|
114
|
+
console.log('Warning: Could not load tokenUtils for benchmarking');
|
|
115
|
+
return results;
|
|
116
|
+
}
|
|
117
|
+
|
|
118
|
+
const { countTokens, estimateTokens } = modules.tokenUtils;
|
|
119
|
+
|
|
120
|
+
// Short text
|
|
121
|
+
results.push(runBenchmark(
|
|
122
|
+
'countTokens (short text ~10 chars)',
|
|
123
|
+
() => countTokens('Hello world'),
|
|
124
|
+
iterations
|
|
125
|
+
));
|
|
126
|
+
|
|
127
|
+
// Medium text
|
|
128
|
+
results.push(runBenchmark(
|
|
129
|
+
'countTokens (medium text ~100 chars)',
|
|
130
|
+
() => countTokens('This is a moderately long sentence that contains a decent amount of text for benchmarking token counting functions in JavaScript.'),
|
|
131
|
+
iterations
|
|
132
|
+
));
|
|
133
|
+
|
|
134
|
+
// Long text
|
|
135
|
+
const longText = 'The quick brown fox jumps over the lazy dog. '.repeat(50);
|
|
136
|
+
results.push(runBenchmark(
|
|
137
|
+
'countTokens (long text ~2000 chars)',
|
|
138
|
+
() => countTokens(longText),
|
|
139
|
+
iterations
|
|
140
|
+
));
|
|
141
|
+
|
|
142
|
+
// Code text
|
|
143
|
+
const codeText = `function helloWorld() {
|
|
144
|
+
console.log("Hello, World!");
|
|
145
|
+
return 42;
|
|
146
|
+
}
|
|
147
|
+
const x = [1, 2, 3, 4, 5];
|
|
148
|
+
const obj = { a: 1, b: 2 };`;
|
|
149
|
+
results.push(runBenchmark(
|
|
150
|
+
'countTokens (code text)',
|
|
151
|
+
() => countTokens(codeText),
|
|
152
|
+
iterations
|
|
153
|
+
));
|
|
154
|
+
|
|
155
|
+
// Unicode text
|
|
156
|
+
results.push(runBenchmark(
|
|
157
|
+
'countTokens (unicode - Japanese)',
|
|
158
|
+
() => countTokens('こんにちは世界!これはテストです。'),
|
|
159
|
+
iterations
|
|
160
|
+
));
|
|
161
|
+
|
|
162
|
+
// estimateTokens if different from countTokens
|
|
163
|
+
if (estimateTokens !== countTokens) {
|
|
164
|
+
results.push(runBenchmark(
|
|
165
|
+
'estimateTokens (short text)',
|
|
166
|
+
() => estimateTokens('Hello world'),
|
|
167
|
+
iterations
|
|
168
|
+
));
|
|
169
|
+
}
|
|
170
|
+
|
|
171
|
+
return results;
|
|
172
|
+
}
|
|
173
|
+
|
|
174
|
+
// Cost estimation benchmarks
|
|
175
|
+
function benchmarkCostEstimation(iterations: number): BenchmarkResult[] {
|
|
176
|
+
const results: BenchmarkResult[] = [];
|
|
177
|
+
const modules = loadModules();
|
|
178
|
+
|
|
179
|
+
if (!modules || !modules.tokenUtils) {
|
|
180
|
+
console.log('Warning: Could not load tokenUtils for benchmarking');
|
|
181
|
+
return results;
|
|
182
|
+
}
|
|
183
|
+
|
|
184
|
+
const { estimateCost } = modules.tokenUtils;
|
|
185
|
+
|
|
186
|
+
if (typeof estimateCost === 'function') {
|
|
187
|
+
// Small request
|
|
188
|
+
results.push(runBenchmark(
|
|
189
|
+
'estimateCost (100 in, 50 out)',
|
|
190
|
+
() => estimateCost(100, 50, 'gpt-4o'),
|
|
191
|
+
iterations
|
|
192
|
+
));
|
|
193
|
+
|
|
194
|
+
// Large request
|
|
195
|
+
results.push(runBenchmark(
|
|
196
|
+
'estimateCost (10K in, 5K out)',
|
|
197
|
+
() => estimateCost(10000, 5000, 'gpt-4o'),
|
|
198
|
+
iterations
|
|
199
|
+
));
|
|
200
|
+
|
|
201
|
+
// Different models
|
|
202
|
+
for (const model of ['gpt-4o', 'gpt-3.5-turbo', 'claude-3-sonnet']) {
|
|
203
|
+
results.push(runBenchmark(
|
|
204
|
+
`estimateCost (${model})`,
|
|
205
|
+
() => estimateCost(1000, 500, model),
|
|
206
|
+
iterations
|
|
207
|
+
));
|
|
208
|
+
}
|
|
209
|
+
}
|
|
210
|
+
|
|
211
|
+
return results;
|
|
212
|
+
}
|
|
213
|
+
|
|
214
|
+
// Memory benchmarks
|
|
215
|
+
function benchmarkMemory(iterations: number): BenchmarkResult[] {
|
|
216
|
+
const results: BenchmarkResult[] = [];
|
|
217
|
+
const modules = loadModules();
|
|
218
|
+
|
|
219
|
+
if (!modules) {
|
|
220
|
+
return results;
|
|
221
|
+
}
|
|
222
|
+
|
|
223
|
+
// Test MemoryTree operations
|
|
224
|
+
if (modules.index.MemoryTree) {
|
|
225
|
+
const MemoryTree = modules.index.MemoryTree;
|
|
226
|
+
|
|
227
|
+
// Add operation
|
|
228
|
+
const memory = new MemoryTree({ maxSize: 1000 });
|
|
229
|
+
results.push(runBenchmark(
|
|
230
|
+
'MemoryTree.add (single entry)',
|
|
231
|
+
() => {
|
|
232
|
+
memory.add('test entry ' + Math.random(), { tags: ['test'] });
|
|
233
|
+
},
|
|
234
|
+
Math.min(iterations, 100) // Fewer iterations for memory tests
|
|
235
|
+
));
|
|
236
|
+
|
|
237
|
+
// Search operation
|
|
238
|
+
for (let i = 0; i < 100; i++) {
|
|
239
|
+
memory.add('test entry ' + i, { tags: ['test'] });
|
|
240
|
+
}
|
|
241
|
+
results.push(runBenchmark(
|
|
242
|
+
'MemoryTree.search',
|
|
243
|
+
() => memory.search('test'),
|
|
244
|
+
Math.min(iterations, 100)
|
|
245
|
+
));
|
|
246
|
+
}
|
|
247
|
+
|
|
248
|
+
return results;
|
|
249
|
+
}
|
|
250
|
+
|
|
251
|
+
// Identify performance issues
|
|
252
|
+
function identifyIssues(benchmarks: BenchmarkResult[]): PerformanceIssue[] {
|
|
253
|
+
const issues: PerformanceIssue[] = [];
|
|
254
|
+
|
|
255
|
+
for (const bench of benchmarks) {
|
|
256
|
+
// Check if avg is too high
|
|
257
|
+
if (bench.avgMs > 100) {
|
|
258
|
+
issues.push({
|
|
259
|
+
name: `Slow operation: ${bench.name}`,
|
|
260
|
+
description: `Average time per operation exceeds 100ms`,
|
|
261
|
+
severity: bench.avgMs > 1000 ? 'critical' : 'high',
|
|
262
|
+
current: `${bench.avgMs.toFixed(2)}ms avg`,
|
|
263
|
+
expected: '< 100ms avg'
|
|
264
|
+
});
|
|
265
|
+
}
|
|
266
|
+
|
|
267
|
+
// Check p99 vs avg ratio (high variance)
|
|
268
|
+
if (bench.avgMs > 0 && bench.p99Ms / bench.avgMs > 10) {
|
|
269
|
+
issues.push({
|
|
270
|
+
name: `High variance: ${bench.name}`,
|
|
271
|
+
description: `p99 is ${(bench.p99Ms / bench.avgMs).toFixed(1)}x higher than average`,
|
|
272
|
+
severity: 'medium',
|
|
273
|
+
current: `p99: ${bench.p99Ms.toFixed(2)}ms, avg: ${bench.avgMs.toFixed(2)}ms`,
|
|
274
|
+
expected: `p99 should be < 5x avg`
|
|
275
|
+
});
|
|
276
|
+
}
|
|
277
|
+
|
|
278
|
+
// Check operations per second
|
|
279
|
+
if (bench.opsPerSecond < 10 && bench.avgMs > 1) {
|
|
280
|
+
issues.push({
|
|
281
|
+
name: `Low throughput: ${bench.name}`,
|
|
282
|
+
description: `Operations per second is very low`,
|
|
283
|
+
severity: 'medium',
|
|
284
|
+
current: `${bench.opsPerSecond.toFixed(2)} ops/sec`,
|
|
285
|
+
expected: '> 100 ops/sec'
|
|
286
|
+
});
|
|
287
|
+
}
|
|
288
|
+
}
|
|
289
|
+
|
|
290
|
+
return issues;
|
|
291
|
+
}
|
|
292
|
+
|
|
293
|
+
// Generate performance tests
|
|
294
|
+
function generatePerformanceTests(benchmarks: BenchmarkResult[]): string {
|
|
295
|
+
const tests: string[] = [];
|
|
296
|
+
|
|
297
|
+
tests.push(`// Auto-generated performance tests
|
|
298
|
+
// Generated: ${new Date().toISOString()}
|
|
299
|
+
|
|
300
|
+
describe('Performance Benchmarks', () => {`);
|
|
301
|
+
|
|
302
|
+
for (const bench of benchmarks) {
|
|
303
|
+
tests.push(`
|
|
304
|
+
describe('${bench.name}', () => {
|
|
305
|
+
it('completes within acceptable time', () => {
|
|
306
|
+
expect(${bench.avgMs.toFixed(2)}).toBeLessThan(100);
|
|
307
|
+
});
|
|
308
|
+
|
|
309
|
+
it('has reasonable p99 latency', () => {
|
|
310
|
+
expect(${bench.p99Ms.toFixed(2)}).toBeLessThan(500);
|
|
311
|
+
});
|
|
312
|
+
|
|
313
|
+
it('achieves minimum throughput', () => {
|
|
314
|
+
expect(${bench.opsPerSecond.toFixed(2)}).toBeGreaterThan(10);
|
|
315
|
+
});
|
|
316
|
+
});`);
|
|
317
|
+
}
|
|
318
|
+
|
|
319
|
+
tests.push('});');
|
|
320
|
+
|
|
321
|
+
return tests.join('\n');
|
|
322
|
+
}
|
|
323
|
+
|
|
324
|
+
// Main analysis function
|
|
325
|
+
function analyzePerformance(config: Partial<PerformanceConfig> = {}): PerformanceReport {
|
|
326
|
+
const cfg = { ...DEFAULT_CONFIG, ...config };
|
|
327
|
+
|
|
328
|
+
console.log(`\nRunning performance benchmarks (${cfg.iterations} iterations)...`);
|
|
329
|
+
|
|
330
|
+
const benchmarks: BenchmarkResult[] = [];
|
|
331
|
+
|
|
332
|
+
// Token counting benchmarks
|
|
333
|
+
console.log(' - Token counting...');
|
|
334
|
+
benchmarks.push(...benchmarkTokenCounting(cfg.iterations));
|
|
335
|
+
|
|
336
|
+
// Cost estimation benchmarks
|
|
337
|
+
console.log(' - Cost estimation...');
|
|
338
|
+
benchmarks.push(...benchmarkCostEstimation(cfg.iterations));
|
|
339
|
+
|
|
340
|
+
// Memory benchmarks
|
|
341
|
+
console.log(' - Memory operations...');
|
|
342
|
+
benchmarks.push(...benchmarkMemory(cfg.iterations));
|
|
343
|
+
|
|
344
|
+
// Identify issues
|
|
345
|
+
const issues = identifyIssues(benchmarks);
|
|
346
|
+
|
|
347
|
+
// Generate recommendations
|
|
348
|
+
const recommendations: string[] = [];
|
|
349
|
+
|
|
350
|
+
if (issues.some(i => i.severity === 'critical')) {
|
|
351
|
+
recommendations.push('CRITICAL: Address critical performance issues before production');
|
|
352
|
+
}
|
|
353
|
+
|
|
354
|
+
if (issues.some(i => i.severity === 'high')) {
|
|
355
|
+
recommendations.push('HIGH: Optimize high-severity performance bottlenecks');
|
|
356
|
+
}
|
|
357
|
+
|
|
358
|
+
const slowBenchmarks = benchmarks.filter(b => b.avgMs > 50);
|
|
359
|
+
if (slowBenchmarks.length > 0) {
|
|
360
|
+
recommendations.push(`Consider caching for ${slowBenchmarks.length} slow operations`);
|
|
361
|
+
}
|
|
362
|
+
|
|
363
|
+
return { benchmarks, issues, recommendations };
|
|
364
|
+
}
|
|
365
|
+
|
|
366
|
+
// Print report
|
|
367
|
+
function printReport(report: PerformanceReport): void {
|
|
368
|
+
console.log('\n========================================');
|
|
369
|
+
console.log('PERFORMANCE AGENT - Benchmark Report');
|
|
370
|
+
console.log('========================================\n');
|
|
371
|
+
|
|
372
|
+
console.log('Benchmarks:');
|
|
373
|
+
for (const bench of report.benchmarks) {
|
|
374
|
+
console.log(`\n ${bench.name}:`);
|
|
375
|
+
console.log(` Avg: ${bench.avgMs.toFixed(4)}ms`);
|
|
376
|
+
console.log(` Min: ${bench.minMs.toFixed(4)}ms`);
|
|
377
|
+
console.log(` Max: ${bench.maxMs.toFixed(4)}ms`);
|
|
378
|
+
console.log(` p50: ${bench.p50Ms.toFixed(4)}ms`);
|
|
379
|
+
console.log(` p95: ${bench.p95Ms.toFixed(4)}ms`);
|
|
380
|
+
console.log(` p99: ${bench.p99Ms.toFixed(4)}ms`);
|
|
381
|
+
console.log(` Ops/sec: ${bench.opsPerSecond.toFixed(2)}`);
|
|
382
|
+
}
|
|
383
|
+
|
|
384
|
+
if (report.issues.length > 0) {
|
|
385
|
+
console.log('\nIssues Found:');
|
|
386
|
+
for (const issue of report.issues) {
|
|
387
|
+
console.log(` [${issue.severity.toUpperCase()}] ${issue.name}`);
|
|
388
|
+
console.log(` ${issue.description}`);
|
|
389
|
+
console.log(` Current: ${issue.current}`);
|
|
390
|
+
console.log(` Expected: ${issue.expected}`);
|
|
391
|
+
}
|
|
392
|
+
} else {
|
|
393
|
+
console.log('\nNo performance issues detected!');
|
|
394
|
+
}
|
|
395
|
+
|
|
396
|
+
if (report.recommendations.length > 0) {
|
|
397
|
+
console.log('\nRecommendations:');
|
|
398
|
+
for (const rec of report.recommendations) {
|
|
399
|
+
console.log(` - ${rec}`);
|
|
400
|
+
}
|
|
401
|
+
}
|
|
402
|
+
}
|
|
403
|
+
|
|
404
|
+
// Run agent if executed directly
|
|
405
|
+
if (require.main === module) {
|
|
406
|
+
const report = analyzePerformance();
|
|
407
|
+
printReport(report);
|
|
408
|
+
|
|
409
|
+
// Generate test code
|
|
410
|
+
const generated = generatePerformanceTests(report.benchmarks);
|
|
411
|
+
const outputPath = path.join(__dirname, '../3-performance-tests-generated.ts');
|
|
412
|
+
fs.writeFileSync(outputPath, generated);
|
|
413
|
+
console.log(`\nGenerated tests written to: ${outputPath}`);
|
|
414
|
+
}
|
|
415
|
+
|
|
416
|
+
export {
|
|
417
|
+
analyzePerformance,
|
|
418
|
+
runBenchmark,
|
|
419
|
+
benchmarkTokenCounting,
|
|
420
|
+
benchmarkCostEstimation,
|
|
421
|
+
benchmarkMemory,
|
|
422
|
+
generatePerformanceTests,
|
|
423
|
+
PerformanceReport,
|
|
424
|
+
BenchmarkResult,
|
|
425
|
+
PerformanceIssue
|
|
426
|
+
};
|
|
@@ -0,0 +1,227 @@
|
|
|
1
|
+
#!/usr/bin/env node
|
|
2
|
+
/**
|
|
3
|
+
* Structure Agent - Analyzes code structure and exports for test coverage
|
|
4
|
+
*
|
|
5
|
+
* This agent identifies:
|
|
6
|
+
* - Exported functions without tests
|
|
7
|
+
* - Type definitions needing validation
|
|
8
|
+
* - Interface implementations
|
|
9
|
+
* - Public API coverage
|
|
10
|
+
*/
|
|
11
|
+
|
|
12
|
+
import * as fs from 'fs';
|
|
13
|
+
import * as path from 'path';
|
|
14
|
+
|
|
15
|
+
interface ExportInfo {
|
|
16
|
+
name: string;
|
|
17
|
+
type: 'function' | 'class' | 'interface' | 'type' | 'const' | 'unknown';
|
|
18
|
+
file: string;
|
|
19
|
+
line: number;
|
|
20
|
+
tested: boolean;
|
|
21
|
+
}
|
|
22
|
+
|
|
23
|
+
interface CoverageReport {
|
|
24
|
+
total: number;
|
|
25
|
+
tested: number;
|
|
26
|
+
untested: ExportInfo[];
|
|
27
|
+
coverage: number;
|
|
28
|
+
}
|
|
29
|
+
|
|
30
|
+
// Main analysis function
|
|
31
|
+
function analyzeStructure(projectRoot: string): CoverageReport {
|
|
32
|
+
const srcDir = path.join(projectRoot, 'src');
|
|
33
|
+
const exports: ExportInfo[] = [];
|
|
34
|
+
|
|
35
|
+
// Find all TypeScript files
|
|
36
|
+
const files = findTsFiles(srcDir);
|
|
37
|
+
|
|
38
|
+
for (const file of files) {
|
|
39
|
+
const content = fs.readFileSync(file, 'utf-8');
|
|
40
|
+
const fileExports = extractExports(content, file);
|
|
41
|
+
exports.push(...fileExports);
|
|
42
|
+
}
|
|
43
|
+
|
|
44
|
+
// Check which exports are tested
|
|
45
|
+
const testFiles = findTestFiles(projectRoot);
|
|
46
|
+
const testedNames = new Set<string>();
|
|
47
|
+
|
|
48
|
+
for (const testFile of testFiles) {
|
|
49
|
+
const content = fs.readFileSync(testFile, 'utf-8');
|
|
50
|
+
// Extract describe/it blocks to find tested names
|
|
51
|
+
const matches = content.matchAll(/(?:describe|it)\s*\(\s*['"]([^'"]+)/g);
|
|
52
|
+
for (const match of matches) {
|
|
53
|
+
testedNames.add(match[1]);
|
|
54
|
+
}
|
|
55
|
+
}
|
|
56
|
+
|
|
57
|
+
// Mark tested exports
|
|
58
|
+
for (const exp of exports) {
|
|
59
|
+
exp.tested = testedNames.has(exp.name) ||
|
|
60
|
+
testedNames.has(exp.name.toLowerCase()) ||
|
|
61
|
+
testedNames.has(sanitizeTestName(exp.name));
|
|
62
|
+
}
|
|
63
|
+
|
|
64
|
+
const tested = exports.filter(e => e.tested);
|
|
65
|
+
const untested = exports.filter(e => !e.tested);
|
|
66
|
+
|
|
67
|
+
return {
|
|
68
|
+
total: exports.length,
|
|
69
|
+
tested: tested.length,
|
|
70
|
+
untested,
|
|
71
|
+
coverage: exports.length > 0 ? (tested.length / exports.length) * 100 : 0
|
|
72
|
+
};
|
|
73
|
+
}
|
|
74
|
+
|
|
75
|
+
function findTsFiles(dir: string): string[] {
|
|
76
|
+
const files: string[] = [];
|
|
77
|
+
|
|
78
|
+
function walk(d: string) {
|
|
79
|
+
const entries = fs.readdirSync(d, { withFileTypes: true });
|
|
80
|
+
for (const entry of entries) {
|
|
81
|
+
if (entry.name === 'node_modules' || entry.name === 'dist' || entry.name === '__pycache__') continue;
|
|
82
|
+
const fullPath = path.join(d, entry.name);
|
|
83
|
+
if (entry.isDirectory()) {
|
|
84
|
+
walk(fullPath);
|
|
85
|
+
} else if (entry.name.endsWith('.ts') && !entry.name.endsWith('.d.ts')) {
|
|
86
|
+
files.push(fullPath);
|
|
87
|
+
}
|
|
88
|
+
}
|
|
89
|
+
}
|
|
90
|
+
|
|
91
|
+
walk(dir);
|
|
92
|
+
return files;
|
|
93
|
+
}
|
|
94
|
+
|
|
95
|
+
function findTestFiles(projectRoot: string): string[] {
|
|
96
|
+
const testDirs = ['test', 'tests', 'test-council'];
|
|
97
|
+
const files: string[] = [];
|
|
98
|
+
|
|
99
|
+
for (const dir of testDirs) {
|
|
100
|
+
const testDir = path.join(projectRoot, dir);
|
|
101
|
+
if (fs.existsSync(testDir)) {
|
|
102
|
+
files.push(...findTestFilesInDir(testDir));
|
|
103
|
+
}
|
|
104
|
+
}
|
|
105
|
+
|
|
106
|
+
return files;
|
|
107
|
+
}
|
|
108
|
+
|
|
109
|
+
function findTestFilesInDir(dir: string): string[] {
|
|
110
|
+
const files: string[] = [];
|
|
111
|
+
|
|
112
|
+
function walk(d: string) {
|
|
113
|
+
const entries = fs.readdirSync(d, { withFileTypes: true });
|
|
114
|
+
for (const entry of entries) {
|
|
115
|
+
if (entry.name === 'node_modules' || entry.name === 'dist') continue;
|
|
116
|
+
const fullPath = path.join(d, entry.name);
|
|
117
|
+
if (entry.isDirectory()) {
|
|
118
|
+
walk(fullPath);
|
|
119
|
+
} else if (entry.name.match(/\.(test|spec)\.(ts|js)$/)) {
|
|
120
|
+
files.push(fullPath);
|
|
121
|
+
}
|
|
122
|
+
}
|
|
123
|
+
}
|
|
124
|
+
|
|
125
|
+
walk(dir);
|
|
126
|
+
return files;
|
|
127
|
+
}
|
|
128
|
+
|
|
129
|
+
function extractExports(content: string, file: string): ExportInfo[] {
|
|
130
|
+
const exports: ExportInfo[] = [];
|
|
131
|
+
const lines = content.split('\n');
|
|
132
|
+
|
|
133
|
+
// Track line numbers
|
|
134
|
+
for (let i = 0; i < lines.length; i++) {
|
|
135
|
+
const line = lines[i].trim();
|
|
136
|
+
const lineNum = i + 1;
|
|
137
|
+
|
|
138
|
+
// export function/const/class/interface/type
|
|
139
|
+
const exportMatch = line.match(/^export\s+(?:function|class|const|interface|type|enum)\s+(\w+)/);
|
|
140
|
+
if (exportMatch) {
|
|
141
|
+
const name = exportMatch[1];
|
|
142
|
+
let type: ExportInfo['type'] = 'unknown';
|
|
143
|
+
|
|
144
|
+
if (line.includes('function')) type = 'function';
|
|
145
|
+
else if (line.includes('class')) type = 'class';
|
|
146
|
+
else if (line.includes('interface')) type = 'interface';
|
|
147
|
+
else if (line.includes('type')) type = 'type';
|
|
148
|
+
else if (line.includes('const') || line.includes('enum')) type = 'const';
|
|
149
|
+
|
|
150
|
+
exports.push({ name, type, file: path.basename(file), line: lineNum, tested: false });
|
|
151
|
+
}
|
|
152
|
+
|
|
153
|
+
// export { name } from
|
|
154
|
+
const reExportMatch = line.match(/^export\s*{\s*(\w+)/);
|
|
155
|
+
if (reExportMatch) {
|
|
156
|
+
exports.push({
|
|
157
|
+
name: reExportMatch[1],
|
|
158
|
+
type: 'unknown',
|
|
159
|
+
file: path.basename(file),
|
|
160
|
+
line: lineNum,
|
|
161
|
+
tested: false
|
|
162
|
+
});
|
|
163
|
+
}
|
|
164
|
+
}
|
|
165
|
+
|
|
166
|
+
return exports;
|
|
167
|
+
}
|
|
168
|
+
|
|
169
|
+
function sanitizeTestName(name: string): string {
|
|
170
|
+
// Convert PascalCase to kebab-case or space-separated
|
|
171
|
+
return name
|
|
172
|
+
.replace(/([A-Z])/g, ' $1')
|
|
173
|
+
.toLowerCase()
|
|
174
|
+
.trim();
|
|
175
|
+
}
|
|
176
|
+
|
|
177
|
+
// Generate tests for untested exports
|
|
178
|
+
function generateTestsForUntested(report: CoverageReport): string {
|
|
179
|
+
const tests: string[] = [];
|
|
180
|
+
|
|
181
|
+
for (const exp of report.untested) {
|
|
182
|
+
if (exp.type === 'function') {
|
|
183
|
+
tests.push(`
|
|
184
|
+
// TODO: Add test for untested export: ${exp.name}
|
|
185
|
+
// File: ${exp.file}:${exp.line}
|
|
186
|
+
// Type: ${exp.type}
|
|
187
|
+
it('${exp.name} - structure coverage', () => {
|
|
188
|
+
// Structure test placeholder
|
|
189
|
+
expect(true).toBe(true);
|
|
190
|
+
});`);
|
|
191
|
+
}
|
|
192
|
+
}
|
|
193
|
+
|
|
194
|
+
return tests.join('\n');
|
|
195
|
+
}
|
|
196
|
+
|
|
197
|
+
// Run agent if executed directly
|
|
198
|
+
if (require.main === module) {
|
|
199
|
+
const projectRoot = path.resolve(__dirname, '../..');
|
|
200
|
+
const report = analyzeStructure(projectRoot);
|
|
201
|
+
|
|
202
|
+
console.log('\n========================================');
|
|
203
|
+
console.log('STRUCTURE AGENT - Coverage Report');
|
|
204
|
+
console.log('========================================\n');
|
|
205
|
+
console.log(`Total Exports: ${report.total}`);
|
|
206
|
+
console.log(`Tested: ${report.tested}`);
|
|
207
|
+
console.log(`Untested: ${report.untested.length}`);
|
|
208
|
+
console.log(`Coverage: ${report.coverage.toFixed(1)}%\n`);
|
|
209
|
+
|
|
210
|
+
if (report.untested.length > 0) {
|
|
211
|
+
console.log('Untested Exports:');
|
|
212
|
+
for (const exp of report.untested.slice(0, 20)) {
|
|
213
|
+
console.log(` - ${exp.name} (${exp.type}) @ ${exp.file}:${exp.line}`);
|
|
214
|
+
}
|
|
215
|
+
if (report.untested.length > 20) {
|
|
216
|
+
console.log(` ... and ${report.untested.length - 20} more`);
|
|
217
|
+
}
|
|
218
|
+
}
|
|
219
|
+
|
|
220
|
+
// Generate test code
|
|
221
|
+
const generated = generateTestsForUntested(report);
|
|
222
|
+
const outputPath = path.join(__dirname, '../1-structure-tests-generated.ts');
|
|
223
|
+
fs.writeFileSync(outputPath, `// Auto-generated structure tests\n${generated}\n`);
|
|
224
|
+
console.log(`\nGenerated tests written to: ${outputPath}`);
|
|
225
|
+
}
|
|
226
|
+
|
|
227
|
+
export { analyzeStructure, CoverageReport, ExportInfo, generateTestsForUntested };
|