@uipath/flow-tool 1.201.0-preview.131 → 1.201.0-preview.133
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/dist/{authoring-rezk32bm.js → authoring-7r6g06c7.js} +2 -2
- package/dist/{conversion-W6HMHBWW-6qbqgqks.js → conversion-H2VEEYLB-ng48mbf7.js} +536 -161
- package/dist/{de-K6MLR24N-5N5XGLH6-gbwpvdz4.js → de-5HU6GFWV-4EZK6G7N-yrh21f52.js} +273 -13
- package/dist/{de-K6MLR24N-00fjz82b.js → de-5HU6GFWV-zheyw6d2.js} +271 -10
- package/dist/de-5KVP66U6-f7knj511.js +161 -0
- package/dist/de-5KVP66U6-t37qr521.js +161 -0
- package/dist/{de-B2KIY5TC-81f703v3.js → de-HKTOVV4H-d5yw5j21.js} +3 -3
- package/dist/{de-WMKOGI7B-r1821wpe.js → de-WMKOGI7B-6dc1zdnp.js} +3 -3
- package/dist/en-4KXU7N7B-24v330b5.js +118 -0
- package/dist/en-4KXU7N7B-66fntxa2.js +118 -0
- package/dist/{en-ILA4HPE5-ALHGESZE-zrv4qss8.js → en-MABQK6C4-REWE2WG2-8dp4cqdg.js} +109 -6
- package/dist/{en-ILA4HPE5-fhwvq92r.js → en-MABQK6C4-nkzq2j8t.js} +106 -2
- package/dist/{en-P6PNWNTO-g1f6cjh5.js → en-P4AS26D7-szrzfxz1.js} +3 -5
- package/dist/{en-X4DLPKLJ-dt1wtdm7.js → en-X4DLPKLJ-1h9gj5c1.js} +3 -3
- package/dist/{es-BVPGJUWA-9gk2g9yb.js → es-BVPGJUWA-0k59eqfv.js} +3 -3
- package/dist/{es-HTQSD352-YTAEDCE5-nhsbg581.js → es-LUTYN7U7-IY5QWOUF-k2aefzc0.js} +273 -13
- package/dist/{es-HTQSD352-m9rr3f0g.js → es-LUTYN7U7-ye5zn4sn.js} +271 -10
- package/dist/{es-MX-3SP2SSV6-f20qspbr.js → es-MX-H2EAXLYI-qwnqt48x.js} +3 -3
- package/dist/{es-MX-FOX525CG-87g500gr.js → es-MX-POESXYP2-9538722k.js} +272 -11
- package/dist/{es-MX-FOX525CG-HLM27OEL-whr0t9v3.js → es-MX-POESXYP2-UW6VP5QI-cd3a79g9.js} +274 -14
- package/dist/{es-MX-U6YL5KQX-ynwxdjr8.js → es-MX-U6YL5KQX-k4qvpgd4.js} +3 -3
- package/dist/es-MX-XAVLOQW3-47a405kf.js +161 -0
- package/dist/es-MX-XAVLOQW3-9b3cp9n8.js +161 -0
- package/dist/{es-BBX4AUJX-h99fcr53.js → es-O4PAV2CR-jwfaawwk.js} +3 -3
- package/dist/es-VFPNVMP5-6sftsdxe.js +161 -0
- package/dist/es-VFPNVMP5-thm0z9sr.js +161 -0
- package/dist/{fr-FSTBVVGK-t3y57xrb.js → fr-5YRGHXF5-q6b63yg0.js} +3 -3
- package/dist/fr-E5MDO4EG-exvrs7ge.js +161 -0
- package/dist/fr-E5MDO4EG-f3x1jeqw.js +161 -0
- package/dist/{fr-EKTTDR7Z-bmpwp31r.js → fr-EKTTDR7Z-m2t68j0m.js} +3 -3
- package/dist/{fr-54THJ4QH-IGABNGU6-fpzw2mba.js → fr-LJIK56LW-STG6QR3T-cextvv8q.js} +274 -14
- package/dist/{fr-54THJ4QH-d6ezjbhw.js → fr-LJIK56LW-xf3xf2kh.js} +272 -11
- package/dist/{index-6hp3p3eq.js → index-dhcd9b82.js} +2 -2
- package/dist/{index-zvxfkpkr.js → index-er6vy3jk.js} +71 -5
- package/dist/index.js +22 -16
- package/dist/init.js +19 -13
- package/dist/ja-DPO5PPLO-fmbg7zt2.js +161 -0
- package/dist/ja-DPO5PPLO-s2a7j3nj.js +161 -0
- package/dist/{ja-MHUPLSWO-CS3KE4Y3-mbtkzwd7.js → ja-FGEGLBJ3-MG3BSWXD-vq19zt2p.js} +273 -13
- package/dist/{ja-MHUPLSWO-3qwy4jc0.js → ja-FGEGLBJ3-w9byjhze.js} +271 -10
- package/dist/{ja-FW33BJBW-htfk4d8t.js → ja-SDAPAVCR-943p5gx7.js} +3 -3
- package/dist/{ja-Y2YBTXNN-96cvy5av.js → ja-Y2YBTXNN-5sszk2zv.js} +3 -3
- package/dist/ko-ITOPNIYI-aqy82jk8.js +161 -0
- package/dist/ko-ITOPNIYI-qz39kamf.js +161 -0
- package/dist/{ko-PDO5AXNW-d653t4pc.js → ko-NAGOZLC3-wb3y9yhh.js} +3 -3
- package/dist/{ko-S33XAV26-cxf4dn2t.js → ko-S33XAV26-rfsv9myd.js} +3 -3
- package/dist/{ko-XPJTXIKK-DNSW4KYQ-956vp61f.js → ko-TUAK3G3A-TBEBKOTM-ns78t2rv.js} +277 -17
- package/dist/{ko-XPJTXIKK-w1hs4px7.js → ko-TUAK3G3A-s25w0ckg.js} +275 -14
- package/dist/{node-rnmmv28d.js → node-qyxvhbw6.js} +2 -2
- package/dist/{packager-tool-cqrsre00.js → packager-tool-0567pxpm.js} +121 -12
- package/dist/{packager-tool-3wf3tkrc.js → packager-tool-09xysehb.js} +3 -5
- package/dist/{packager-tool-b0s5d4ek.js → packager-tool-0sz8se22.js} +1 -1
- package/dist/{packager-tool-s5y06ny4.js → packager-tool-2xgc101h.js} +544 -45
- package/dist/{packager-tool-03knb0er.js → packager-tool-3wygxefb.js} +116 -17
- package/dist/packager-tool-6ejqwbw9.js +26 -0
- package/dist/{packager-tool-952kqpax.js → packager-tool-78a3szsp.js} +890 -339
- package/dist/{packager-tool-xzcywm2b.js → packager-tool-c42dp002.js} +1 -1
- package/dist/{packager-tool-v61wap3v.js → packager-tool-djagtpkt.js} +12111 -6271
- package/dist/{packager-tool-1gjx8w7m.js → packager-tool-fgdsmd8k.js} +258 -486
- package/dist/packager-tool-gasyxww9.js +23964 -0
- package/dist/packager-tool-gmrwvnyh.js +1171 -0
- package/dist/packager-tool-jysvpt8p.js +11 -0
- package/dist/packager-tool-m6x5sczh.js +899 -0
- package/dist/packager-tool-mcaf4ewf.js +46 -0
- package/dist/packager-tool-p9j43p8h.js +115 -0
- package/dist/{packager-tool-v2p6g3fd.js → packager-tool-pkq70dec.js} +1 -1
- package/dist/{packager-tool-zh6g3h65.js → packager-tool-sf18719z.js} +5 -5
- package/dist/packager-tool-tq9q9m77.js +115 -0
- package/dist/{packager-tool-1yy9x4vk.js → packager-tool-v74yk918.js} +2 -2
- package/dist/packager-tool-vdd02vns.js +55924 -0
- package/dist/packager-tool-xj0agqpk.js +11 -0
- package/dist/packager-tool.js +15 -9
- package/dist/{pt-JLZYSSXX-7jth254x.js → pt-7TR74ZMB-1g8my56g.js} +3 -3
- package/dist/{pt-BR-TV7NR3ID-PDSSHRHU-nddwzj29.js → pt-BR-7XGDXRDJ-TR5D7JYL-533wvy7p.js} +274 -14
- package/dist/{pt-BR-TV7NR3ID-vgd26pj1.js → pt-BR-7XGDXRDJ-espnkgrj.js} +272 -11
- package/dist/pt-BR-BXURWX4W-dqr4dnf9.js +161 -0
- package/dist/pt-BR-BXURWX4W-xzyndmrx.js +161 -0
- package/dist/{pt-BR-LQH3RSBB-kyb4jmcw.js → pt-BR-DXXSYJDO-bn84ewp7.js} +3 -3
- package/dist/{pt-BR-JO45SCNI-4kbfnzc4.js → pt-BR-JO45SCNI-mt93s33c.js} +3 -3
- package/dist/{pt-4JEPEAST-JO5WN33S-ccdbdab2.js → pt-TMVOOITK-H2ADYHSD-ezpggv33.js} +273 -13
- package/dist/{pt-4JEPEAST-40npwxyw.js → pt-TMVOOITK-myha14p9.js} +271 -10
- package/dist/{pt-W3NA4LIZ-f02gapvs.js → pt-W3NA4LIZ-fwd1q4j0.js} +3 -3
- package/dist/pt-YI4VE4YQ-521681zf.js +161 -0
- package/dist/pt-YI4VE4YQ-v7gjeg82.js +161 -0
- package/dist/{ro-RKJ74FV3-QS3I6XKJ-kb31d85k.js → ro-2VALF3S5-NXBJUEZN-nqw1zs33.js} +277 -17
- package/dist/{ro-RKJ74FV3-qn0ey39v.js → ro-2VALF3S5-cdz5afkr.js} +275 -14
- package/dist/{ro-E2EBPQGT-vn573vcr.js → ro-BJFER5OI-334n28rj.js} +3 -3
- package/dist/ro-RKXZUG7Z-cm3xn3fj.js +161 -0
- package/dist/ro-RKXZUG7Z-pdztdafr.js +161 -0
- package/dist/{ro-WLKX6HF5-fcbe2zz8.js → ro-WLKX6HF5-b9tj8q67.js} +3 -3
- package/dist/ru-GSCT4ARZ-FTVCGQTL-5nh4jf2e.js +14 -0
- package/dist/ru-T73OPRTC-mp5gn7gs.js +10 -0
- package/dist/ru-WBLQ5B2U-8xd4gpv2.js +10 -0
- package/dist/ru-WBLQ5B2U-mswqs9wc.js +10 -0
- package/dist/serialization-GAW7UFOK-9hkdmmd2.js +36541 -0
- package/dist/services/flow-eval-schema/deterministic-evaluators.d.ts +110 -0
- package/dist/services/flow-eval-schema/evaluator-fields.d.ts +66 -0
- package/dist/services/flow-eval-schema/helpers.d.ts +9 -0
- package/dist/services/flow-eval-schema/index.d.ts +4 -0
- package/dist/services/flow-eval-schema/llm-judge-evaluators.d.ts +174 -0
- package/dist/services/flow-eval-schema/prompts.d.ts +2 -0
- package/dist/services/flow-eval-schema/registry.d.ts +375 -0
- package/dist/services/flow-eval-schema/tool-call-evaluators.d.ts +93 -0
- package/dist/services/flow-eval-schema/types.d.ts +78 -0
- package/dist/services/flow-eval-types.d.ts +6 -0
- package/dist/services/flow-eval-utils.d.ts +14 -0
- package/dist/services/flow-validate-service.d.ts +11 -0
- package/dist/services/packaging-utils.d.ts +3 -29
- package/dist/tool.js +22 -16
- package/dist/{tr-3U2W2XCA-zg3zcmq3.js → tr-3U2W2XCA-43ptfc0m.js} +3 -3
- package/dist/tr-AA6O2WSR-bj9n80k8.js +161 -0
- package/dist/tr-AA6O2WSR-ga1xnymp.js +161 -0
- package/dist/{tr-X7R6GQNB-VMEZAWRJ-vdbjqw23.js → tr-RIFCEXFH-BZS7S7QK-q465b62g.js} +279 -19
- package/dist/{tr-X7R6GQNB-shwv16dz.js → tr-RIFCEXFH-hgkb4g8y.js} +277 -16
- package/dist/{tr-ZUUSMHER-pfn5dq1r.js → tr-RQTZ2DMS-bpgd4kma.js} +3 -3
- package/dist/utils/flow-io.d.ts +8 -0
- package/dist/validation.js +19 -13
- package/dist/{zh-CN-FEZ24J4A-q5b5x20w.js → zh-CN-IP3FMXWM-1cfarhcp.js} +271 -10
- package/dist/{zh-CN-FEZ24J4A-LSEMG63K-db8kf6rm.js → zh-CN-IP3FMXWM-H3JI2XUM-ta4yac7h.js} +273 -13
- package/dist/{zh-CN-IPZYQR3P-1h76csc4.js → zh-CN-IPZYQR3P-egjmvyjf.js} +3 -3
- package/dist/zh-CN-R2KLHXQV-wttmznep.js +161 -0
- package/dist/zh-CN-R2KLHXQV-x6gnk2n0.js +161 -0
- package/dist/{zh-CN-YFKQ7HBX-w98q5zqf.js → zh-CN-ZC6LT4IN-vty9mvv8.js} +3 -3
- package/dist/{zh-TW-GKSCT6NW-JLMWTDAF-g2smq1gm.js → zh-TW-AGIDXIPF-R4XZRDJZ-cgxxz5k4.js} +274 -14
- package/dist/{zh-TW-GKSCT6NW-4n1yfre0.js → zh-TW-AGIDXIPF-n6r6w7vy.js} +272 -11
- package/dist/{zh-TW-CP5PCVEJ-wvbxf7c7.js → zh-TW-R7Y5KVMU-q81f4z2p.js} +3 -3
- package/dist/zh-TW-RJGCQANM-gqn0aca0.js +161 -0
- package/dist/zh-TW-RJGCQANM-tt8j5wv2.js +161 -0
- package/dist/{zh-TW-YPVSQVGT-wbybwpwt.js → zh-TW-YPVSQVGT-b378ztph.js} +3 -3
- package/package.json +2 -2
- package/dist/packager-tool-1169tjst.js +0 -87554
- package/dist/packager-tool-c19w8vg9.js +0 -26
- package/dist/packager-tool-gbgfwx8f.js +0 -2
- package/dist/packager-tool-jpcm3xas.js +0 -12
- package/dist/packager-tool-mt2g0zjx.js +0 -36187
- package/dist/packager-tool-pvgqbm48.js +0 -21
- package/dist/ru-D6DDND7Z-e4en56d2.js +0 -10
- package/dist/ru-GSCT4ARZ-3TVTUGWP-bbg6360z.js +0 -15
- package/dist/serialization-GTFWODSH-p7wmsmj7.js +0 -148
- /package/dist/{packager-tool-ny53xmsd.js → packager-tool-cvwr40ax.js} +0 -0
|
@@ -0,0 +1,110 @@
|
|
|
1
|
+
export declare const DETERMINISTIC_EVALUATOR_TYPES: {
|
|
2
|
+
readonly "uipath-contains": {
|
|
3
|
+
readonly name: "Contains Text";
|
|
4
|
+
readonly description: "Check if output contains specific text";
|
|
5
|
+
readonly category: "deterministic";
|
|
6
|
+
readonly shortDescription: "Passes when output contains given text";
|
|
7
|
+
readonly configFields: readonly [{
|
|
8
|
+
readonly key: "targetOutputKey";
|
|
9
|
+
readonly label: "Target Output Key";
|
|
10
|
+
readonly type: "targetOutputKey";
|
|
11
|
+
readonly required: false;
|
|
12
|
+
readonly description: "The key in the output to evaluate";
|
|
13
|
+
}, {
|
|
14
|
+
readonly key: "caseSensitive";
|
|
15
|
+
readonly label: "Case Sensitive";
|
|
16
|
+
readonly type: "boolean";
|
|
17
|
+
readonly required: false;
|
|
18
|
+
readonly default: false;
|
|
19
|
+
readonly description: "If true, the comparison is case-sensitive";
|
|
20
|
+
}, {
|
|
21
|
+
readonly key: "negated";
|
|
22
|
+
readonly type: "boolean";
|
|
23
|
+
readonly required: false;
|
|
24
|
+
readonly default: false;
|
|
25
|
+
readonly description: "If true, the evaluation passes when the match is not found";
|
|
26
|
+
readonly label: "Negated (should NOT contain)";
|
|
27
|
+
}];
|
|
28
|
+
readonly criteriaFields: readonly [{
|
|
29
|
+
readonly key: "searchText";
|
|
30
|
+
readonly label: "Search Text";
|
|
31
|
+
readonly type: "string";
|
|
32
|
+
readonly required: false;
|
|
33
|
+
readonly minLength: 1;
|
|
34
|
+
readonly description: "The text to search for in the output";
|
|
35
|
+
}];
|
|
36
|
+
readonly examples: {
|
|
37
|
+
readonly criteria: {
|
|
38
|
+
readonly searchText: "loan approved";
|
|
39
|
+
};
|
|
40
|
+
};
|
|
41
|
+
};
|
|
42
|
+
readonly "uipath-exact-match": {
|
|
43
|
+
readonly name: "Exact Match";
|
|
44
|
+
readonly description: "Check if output exactly matches expected value";
|
|
45
|
+
readonly category: "deterministic";
|
|
46
|
+
readonly shortDescription: "Passes when output is an exact match";
|
|
47
|
+
readonly configFields: readonly [{
|
|
48
|
+
readonly key: "targetOutputKey";
|
|
49
|
+
readonly label: "Target Output Key";
|
|
50
|
+
readonly type: "targetOutputKey";
|
|
51
|
+
readonly required: false;
|
|
52
|
+
readonly description: "The key in the output to evaluate";
|
|
53
|
+
}, {
|
|
54
|
+
readonly key: "caseSensitive";
|
|
55
|
+
readonly label: "Case Sensitive";
|
|
56
|
+
readonly type: "boolean";
|
|
57
|
+
readonly required: false;
|
|
58
|
+
readonly default: false;
|
|
59
|
+
readonly description: "If true, the comparison is case-sensitive";
|
|
60
|
+
}, {
|
|
61
|
+
readonly key: "negated";
|
|
62
|
+
readonly label: "Negated";
|
|
63
|
+
readonly type: "boolean";
|
|
64
|
+
readonly required: false;
|
|
65
|
+
readonly default: false;
|
|
66
|
+
readonly description: "If true, the evaluation passes when the match is not found";
|
|
67
|
+
}];
|
|
68
|
+
readonly criteriaFields: readonly [{
|
|
69
|
+
readonly key: "expectedOutput";
|
|
70
|
+
readonly label: "Expected Output";
|
|
71
|
+
readonly type: "object";
|
|
72
|
+
readonly required: false;
|
|
73
|
+
readonly description: "The expected output value for this evaluator";
|
|
74
|
+
}];
|
|
75
|
+
readonly examples: {
|
|
76
|
+
readonly criteria: {
|
|
77
|
+
readonly expectedOutput: {
|
|
78
|
+
readonly value: "Application approved";
|
|
79
|
+
};
|
|
80
|
+
};
|
|
81
|
+
};
|
|
82
|
+
};
|
|
83
|
+
readonly "uipath-json-similarity": {
|
|
84
|
+
readonly name: "JSON Similarity";
|
|
85
|
+
readonly description: "Compare JSON output similarity";
|
|
86
|
+
readonly category: "deterministic";
|
|
87
|
+
readonly shortDescription: "Passes when JSON structures are similar";
|
|
88
|
+
readonly configFields: readonly [{
|
|
89
|
+
readonly key: "targetOutputKey";
|
|
90
|
+
readonly label: "Target Output Key";
|
|
91
|
+
readonly type: "targetOutputKey";
|
|
92
|
+
readonly required: false;
|
|
93
|
+
readonly description: "The key in the output to evaluate";
|
|
94
|
+
}];
|
|
95
|
+
readonly criteriaFields: readonly [{
|
|
96
|
+
readonly key: "expectedOutput";
|
|
97
|
+
readonly label: "Expected Output";
|
|
98
|
+
readonly type: "object";
|
|
99
|
+
readonly required: false;
|
|
100
|
+
readonly description: "The expected output value for this evaluator";
|
|
101
|
+
}];
|
|
102
|
+
readonly examples: {
|
|
103
|
+
readonly criteria: {
|
|
104
|
+
readonly expectedOutput: {
|
|
105
|
+
readonly status: "ok";
|
|
106
|
+
};
|
|
107
|
+
};
|
|
108
|
+
};
|
|
109
|
+
};
|
|
110
|
+
};
|
|
@@ -0,0 +1,66 @@
|
|
|
1
|
+
import type { EvaluatorFieldDefinition } from "./types.js";
|
|
2
|
+
export declare const targetOutputKeyConfigField: {
|
|
3
|
+
readonly key: "targetOutputKey";
|
|
4
|
+
readonly label: "Target Output Key";
|
|
5
|
+
readonly type: "targetOutputKey";
|
|
6
|
+
readonly required: false;
|
|
7
|
+
readonly description: "The key in the output to evaluate";
|
|
8
|
+
};
|
|
9
|
+
export declare const caseSensitiveConfigField: {
|
|
10
|
+
readonly key: "caseSensitive";
|
|
11
|
+
readonly label: "Case Sensitive";
|
|
12
|
+
readonly type: "boolean";
|
|
13
|
+
readonly required: false;
|
|
14
|
+
readonly default: false;
|
|
15
|
+
readonly description: "If true, the comparison is case-sensitive";
|
|
16
|
+
};
|
|
17
|
+
export declare const negatedConfigField: {
|
|
18
|
+
readonly key: "negated";
|
|
19
|
+
readonly label: "Negated";
|
|
20
|
+
readonly type: "boolean";
|
|
21
|
+
readonly required: false;
|
|
22
|
+
readonly default: false;
|
|
23
|
+
readonly description: "If true, the evaluation passes when the match is not found";
|
|
24
|
+
};
|
|
25
|
+
export declare const expectedOutputCriteriaField: {
|
|
26
|
+
readonly key: "expectedOutput";
|
|
27
|
+
readonly label: "Expected Output";
|
|
28
|
+
readonly type: "object";
|
|
29
|
+
readonly required: false;
|
|
30
|
+
readonly description: "The expected output value for this evaluator";
|
|
31
|
+
};
|
|
32
|
+
export declare const llmBaseConfigFields: readonly [{
|
|
33
|
+
readonly key: "model";
|
|
34
|
+
readonly label: "Model";
|
|
35
|
+
readonly type: "llmModel";
|
|
36
|
+
readonly required: true;
|
|
37
|
+
readonly description: "The LLM model to use for judging";
|
|
38
|
+
}, {
|
|
39
|
+
readonly key: "prompt";
|
|
40
|
+
readonly label: "Evaluation Prompt";
|
|
41
|
+
readonly type: "string";
|
|
42
|
+
readonly required: true;
|
|
43
|
+
readonly default: "As an expert evaluator, analyze the semantic similarity of these JSON contents to determine a score from 0-100. Focus on comparing the meaning and contextual equivalence of corresponding fields, accounting for alternative valid expressions, synonyms, and reasonable variations in language while maintaining high standards for accuracy and completeness. Provide your score with a justification, explaining briefly and concisely why you gave that score.\n----\nExpectedOutput:\n{{ExpectedOutput}}\n----\nActualOutput:\n{{ActualOutput}}\n";
|
|
44
|
+
readonly description: "Custom prompt template for the LLM judge";
|
|
45
|
+
}, {
|
|
46
|
+
readonly key: "temperature";
|
|
47
|
+
readonly label: "Temperature";
|
|
48
|
+
readonly type: "number";
|
|
49
|
+
readonly required: true;
|
|
50
|
+
readonly default: 0;
|
|
51
|
+
readonly min: 0;
|
|
52
|
+
readonly max: 2;
|
|
53
|
+
readonly step: 0.1;
|
|
54
|
+
readonly description: "Temperature for LLM generation";
|
|
55
|
+
}, {
|
|
56
|
+
readonly key: "maxTokens";
|
|
57
|
+
readonly label: "Max Tokens";
|
|
58
|
+
readonly type: "number";
|
|
59
|
+
readonly required: true;
|
|
60
|
+
readonly default: 8192;
|
|
61
|
+
readonly min: 0;
|
|
62
|
+
readonly max: 16384;
|
|
63
|
+
readonly step: 256;
|
|
64
|
+
readonly description: "Maximum tokens for LLM response";
|
|
65
|
+
}];
|
|
66
|
+
export declare const llmTrajectoryBaseConfigFields: readonly EvaluatorFieldDefinition[];
|
|
@@ -0,0 +1,9 @@
|
|
|
1
|
+
import { type EvaluatorTypeId } from "./registry.js";
|
|
2
|
+
import type { EvaluatorTypeDefinition } from "./types.js";
|
|
3
|
+
export declare function evaluatorTypeAliases(): string[];
|
|
4
|
+
export declare function isEvaluatorTypeId(value: string): value is EvaluatorTypeId;
|
|
5
|
+
export declare function getEvaluatorTypeDefinition(typeId: string): EvaluatorTypeDefinition | null;
|
|
6
|
+
export declare function resolveEvaluatorType(type: string): EvaluatorTypeId;
|
|
7
|
+
export declare function getDefaultConfigForEvaluatorType(typeId: EvaluatorTypeId): Record<string, unknown>;
|
|
8
|
+
export declare function evaluatorTypeHasConfigField(typeId: string, fieldKey: string): boolean;
|
|
9
|
+
export declare function evaluatorCanUseSharedExpectedOutput(typeId: string): boolean;
|
|
@@ -0,0 +1,174 @@
|
|
|
1
|
+
export declare const LLM_JUDGE_EVALUATOR_TYPES: {
|
|
2
|
+
readonly "uipath-llm-judge-output-semantic-similarity": {
|
|
3
|
+
readonly name: "LLM Judge Output";
|
|
4
|
+
readonly description: "Use LLM to judge semantic similarity of outputs";
|
|
5
|
+
readonly category: "llm-judge";
|
|
6
|
+
readonly shortDescription: "AI judges if output semantically matches expected";
|
|
7
|
+
readonly configFields: readonly [{
|
|
8
|
+
readonly key: "model";
|
|
9
|
+
readonly label: "Model";
|
|
10
|
+
readonly type: "llmModel";
|
|
11
|
+
readonly required: true;
|
|
12
|
+
readonly description: "The LLM model to use for judging";
|
|
13
|
+
}, {
|
|
14
|
+
readonly key: "prompt";
|
|
15
|
+
readonly label: "Evaluation Prompt";
|
|
16
|
+
readonly type: "string";
|
|
17
|
+
readonly required: true;
|
|
18
|
+
readonly default: "As an expert evaluator, analyze the semantic similarity of these JSON contents to determine a score from 0-100. Focus on comparing the meaning and contextual equivalence of corresponding fields, accounting for alternative valid expressions, synonyms, and reasonable variations in language while maintaining high standards for accuracy and completeness. Provide your score with a justification, explaining briefly and concisely why you gave that score.\n----\nExpectedOutput:\n{{ExpectedOutput}}\n----\nActualOutput:\n{{ActualOutput}}\n";
|
|
19
|
+
readonly description: "Custom prompt template for the LLM judge";
|
|
20
|
+
}, {
|
|
21
|
+
readonly key: "temperature";
|
|
22
|
+
readonly label: "Temperature";
|
|
23
|
+
readonly type: "number";
|
|
24
|
+
readonly required: true;
|
|
25
|
+
readonly default: 0;
|
|
26
|
+
readonly min: 0;
|
|
27
|
+
readonly max: 2;
|
|
28
|
+
readonly step: 0.1;
|
|
29
|
+
readonly description: "Temperature for LLM generation";
|
|
30
|
+
}, {
|
|
31
|
+
readonly key: "maxTokens";
|
|
32
|
+
readonly label: "Max Tokens";
|
|
33
|
+
readonly type: "number";
|
|
34
|
+
readonly required: true;
|
|
35
|
+
readonly default: 8192;
|
|
36
|
+
readonly min: 0;
|
|
37
|
+
readonly max: 16384;
|
|
38
|
+
readonly step: 256;
|
|
39
|
+
readonly description: "Maximum tokens for LLM response";
|
|
40
|
+
}, {
|
|
41
|
+
readonly key: "targetOutputKey";
|
|
42
|
+
readonly label: "Target Output Key";
|
|
43
|
+
readonly type: "targetOutputKey";
|
|
44
|
+
readonly required: false;
|
|
45
|
+
readonly description: "The key in the output to evaluate";
|
|
46
|
+
}];
|
|
47
|
+
readonly criteriaFields: readonly [{
|
|
48
|
+
readonly key: "expectedOutput";
|
|
49
|
+
readonly label: "Expected Output";
|
|
50
|
+
readonly type: "object";
|
|
51
|
+
readonly required: false;
|
|
52
|
+
readonly description: "The expected output value for this evaluator";
|
|
53
|
+
}];
|
|
54
|
+
readonly examples: {
|
|
55
|
+
readonly config: {
|
|
56
|
+
readonly model: "gpt-4o";
|
|
57
|
+
readonly temperature: 0;
|
|
58
|
+
};
|
|
59
|
+
readonly criteria: {
|
|
60
|
+
readonly expectedOutput: {
|
|
61
|
+
readonly text: "The loan has been approved for the requested amount.";
|
|
62
|
+
};
|
|
63
|
+
};
|
|
64
|
+
};
|
|
65
|
+
};
|
|
66
|
+
readonly "uipath-llm-judge-output-strict-json-similarity": {
|
|
67
|
+
readonly name: "LLM Judge Strict JSON";
|
|
68
|
+
readonly description: "Use LLM to judge strict JSON similarity";
|
|
69
|
+
readonly category: "llm-judge";
|
|
70
|
+
readonly shortDescription: "AI judges strict JSON structural match";
|
|
71
|
+
readonly configFields: readonly [{
|
|
72
|
+
readonly key: "model";
|
|
73
|
+
readonly label: "Model";
|
|
74
|
+
readonly type: "llmModel";
|
|
75
|
+
readonly required: true;
|
|
76
|
+
readonly description: "The LLM model to use for judging";
|
|
77
|
+
}, {
|
|
78
|
+
readonly key: "prompt";
|
|
79
|
+
readonly label: "Evaluation Prompt";
|
|
80
|
+
readonly type: "string";
|
|
81
|
+
readonly required: true;
|
|
82
|
+
readonly default: "As an expert evaluator, analyze the semantic similarity of these JSON contents to determine a score from 0-100. Focus on comparing the meaning and contextual equivalence of corresponding fields, accounting for alternative valid expressions, synonyms, and reasonable variations in language while maintaining high standards for accuracy and completeness. Provide your score with a justification, explaining briefly and concisely why you gave that score.\n----\nExpectedOutput:\n{{ExpectedOutput}}\n----\nActualOutput:\n{{ActualOutput}}\n";
|
|
83
|
+
readonly description: "Custom prompt template for the LLM judge";
|
|
84
|
+
}, {
|
|
85
|
+
readonly key: "temperature";
|
|
86
|
+
readonly label: "Temperature";
|
|
87
|
+
readonly type: "number";
|
|
88
|
+
readonly required: true;
|
|
89
|
+
readonly default: 0;
|
|
90
|
+
readonly min: 0;
|
|
91
|
+
readonly max: 2;
|
|
92
|
+
readonly step: 0.1;
|
|
93
|
+
readonly description: "Temperature for LLM generation";
|
|
94
|
+
}, {
|
|
95
|
+
readonly key: "maxTokens";
|
|
96
|
+
readonly label: "Max Tokens";
|
|
97
|
+
readonly type: "number";
|
|
98
|
+
readonly required: true;
|
|
99
|
+
readonly default: 8192;
|
|
100
|
+
readonly min: 0;
|
|
101
|
+
readonly max: 16384;
|
|
102
|
+
readonly step: 256;
|
|
103
|
+
readonly description: "Maximum tokens for LLM response";
|
|
104
|
+
}, {
|
|
105
|
+
readonly key: "targetOutputKey";
|
|
106
|
+
readonly label: "Target Output Key";
|
|
107
|
+
readonly type: "targetOutputKey";
|
|
108
|
+
readonly required: false;
|
|
109
|
+
readonly description: "The key in the output to evaluate";
|
|
110
|
+
}];
|
|
111
|
+
readonly criteriaFields: readonly [{
|
|
112
|
+
readonly key: "expectedOutput";
|
|
113
|
+
readonly label: "Expected Output";
|
|
114
|
+
readonly type: "object";
|
|
115
|
+
readonly required: false;
|
|
116
|
+
readonly description: "The expected output value for this evaluator";
|
|
117
|
+
}];
|
|
118
|
+
readonly examples: {
|
|
119
|
+
readonly config: {
|
|
120
|
+
readonly model: "gpt-4o";
|
|
121
|
+
readonly temperature: 0;
|
|
122
|
+
};
|
|
123
|
+
readonly criteria: {
|
|
124
|
+
readonly expectedOutput: {
|
|
125
|
+
readonly result: "success";
|
|
126
|
+
};
|
|
127
|
+
};
|
|
128
|
+
};
|
|
129
|
+
};
|
|
130
|
+
readonly "uipath-llm-judge-trajectory-similarity": {
|
|
131
|
+
readonly name: "LLM Judge Trajectory";
|
|
132
|
+
readonly description: "Use LLM to judge agent trajectory and behavior";
|
|
133
|
+
readonly category: "llm-judge";
|
|
134
|
+
readonly shortDescription: "AI judges if agent followed expected behavior";
|
|
135
|
+
readonly configFields: readonly import("./types.js").EvaluatorFieldDefinition[];
|
|
136
|
+
readonly criteriaFields: readonly [{
|
|
137
|
+
readonly key: "expectedAgentBehavior";
|
|
138
|
+
readonly label: "Expected Agent Behavior";
|
|
139
|
+
readonly type: "string";
|
|
140
|
+
readonly required: false;
|
|
141
|
+
readonly description: "Description of expected agent behavior and actions";
|
|
142
|
+
}];
|
|
143
|
+
readonly examples: {
|
|
144
|
+
readonly config: {
|
|
145
|
+
readonly model: "gpt-4o";
|
|
146
|
+
};
|
|
147
|
+
readonly criteria: {
|
|
148
|
+
readonly expectedAgentBehavior: "Agent checks eligibility, then calls approval API";
|
|
149
|
+
};
|
|
150
|
+
};
|
|
151
|
+
};
|
|
152
|
+
readonly "uipath-llm-judge-trajectory-simulation": {
|
|
153
|
+
readonly name: "LLM Judge Simulation";
|
|
154
|
+
readonly description: "Use LLM to judge simulation trajectory";
|
|
155
|
+
readonly category: "llm-judge";
|
|
156
|
+
readonly shortDescription: "AI judges simulated agent execution path";
|
|
157
|
+
readonly configFields: readonly import("./types.js").EvaluatorFieldDefinition[];
|
|
158
|
+
readonly criteriaFields: readonly [{
|
|
159
|
+
readonly key: "expectedAgentBehavior";
|
|
160
|
+
readonly label: "Expected Agent Behavior";
|
|
161
|
+
readonly type: "string";
|
|
162
|
+
readonly required: false;
|
|
163
|
+
readonly description: "Description of expected agent behavior during simulation";
|
|
164
|
+
}];
|
|
165
|
+
readonly examples: {
|
|
166
|
+
readonly config: {
|
|
167
|
+
readonly model: "gpt-4o";
|
|
168
|
+
};
|
|
169
|
+
readonly criteria: {
|
|
170
|
+
readonly expectedAgentBehavior: "Agent simulates the user journey through checkout flow";
|
|
171
|
+
};
|
|
172
|
+
};
|
|
173
|
+
};
|
|
174
|
+
};
|
|
@@ -0,0 +1,2 @@
|
|
|
1
|
+
export declare const DEFAULT_EVALUATOR_PROMPT = "As an expert evaluator, analyze the semantic similarity of these JSON contents to determine a score from 0-100. Focus on comparing the meaning and contextual equivalence of corresponding fields, accounting for alternative valid expressions, synonyms, and reasonable variations in language while maintaining high standards for accuracy and completeness. Provide your score with a justification, explaining briefly and concisely why you gave that score.\n----\nExpectedOutput:\n{{ExpectedOutput}}\n----\nActualOutput:\n{{ActualOutput}}\n";
|
|
2
|
+
export declare const DEFAULT_TRAJECTORY_EVALUATOR_PROMPT = "As an expert evaluator, determine how well the agent did on a scale of 0-100. Focus on if the simulation was successful and if the agent behaved according to the expected output accounting for alternative valid expressions, and reasonable variations in language while maintaining high standards for accuracy and completeness. Provide your score with a justification, explaining briefly and concisely why you gave that score.\n----\nUserOrSyntheticInputGivenToAgent:\n{{UserOrSyntheticInput}}\n----\nSimulationInstructions:\n{{SimulationInstructions}}\n----\nExpectedAgentBehavior:\n{{ExpectedAgentBehavior}}\n----\nAgentRunHistory:\n{{AgentRunHistory}}\n";
|