@oxy.so/contracts 1.0.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
Files changed (147) hide show
  1. package/LICENSE +202 -0
  2. package/NOTICE +16 -0
  3. package/dist/cjs/.tsbuildinfo +1 -0
  4. package/dist/cjs/accountGraph.js +489 -0
  5. package/dist/cjs/agency.js +439 -0
  6. package/dist/cjs/browserHub.js +215 -0
  7. package/dist/cjs/civic.js +163 -0
  8. package/dist/cjs/commonsSignIn.js +59 -0
  9. package/dist/cjs/deviceBoot.js +50 -0
  10. package/dist/cjs/deviceDirectory.js +189 -0
  11. package/dist/cjs/devicePairing.js +138 -0
  12. package/dist/cjs/deviceSession.js +164 -0
  13. package/dist/cjs/emailAgentContext.js +32 -0
  14. package/dist/cjs/followGraph.js +28 -0
  15. package/dist/cjs/identity.js +258 -0
  16. package/dist/cjs/inboxPush.js +24 -0
  17. package/dist/cjs/index.js +618 -0
  18. package/dist/cjs/inference/accountBilling.js +334 -0
  19. package/dist/cjs/inference/aliaModelRelease.js +262 -0
  20. package/dist/cjs/inference/attribution.js +106 -0
  21. package/dist/cjs/inference/catalogue.js +487 -0
  22. package/dist/cjs/inference/entitlement.js +217 -0
  23. package/dist/cjs/inference/errors.js +309 -0
  24. package/dist/cjs/inference/identifiers.js +224 -0
  25. package/dist/cjs/inference/inbox.js +105 -0
  26. package/dist/cjs/inference/modelDocumentation.js +433 -0
  27. package/dist/cjs/inference/money.js +188 -0
  28. package/dist/cjs/inference/priceVersion.js +110 -0
  29. package/dist/cjs/inference/providerConnection.js +455 -0
  30. package/dist/cjs/inference/request.js +477 -0
  31. package/dist/cjs/inference/routingPolicy.js +318 -0
  32. package/dist/cjs/inference/streamEvents.js +258 -0
  33. package/dist/cjs/inference/usage.js +329 -0
  34. package/dist/cjs/inference/version.js +105 -0
  35. package/dist/cjs/keyRecovery.js +91 -0
  36. package/dist/cjs/keyRotation.js +75 -0
  37. package/dist/cjs/links.js +68 -0
  38. package/dist/cjs/moderationReputation.js +298 -0
  39. package/dist/cjs/oauth.js +66 -0
  40. package/dist/cjs/oxyRecordTypes.js +71 -0
  41. package/dist/cjs/protocol.js +53 -0
  42. package/dist/cjs/recommendations.js +168 -0
  43. package/dist/cjs/reputation.js +297 -0
  44. package/dist/cjs/sessionStatus.js +121 -0
  45. package/dist/cjs/transparency.js +89 -0
  46. package/dist/cjs/updates.js +252 -0
  47. package/dist/cjs/userInvalidation.js +89 -0
  48. package/dist/cjs/userResponse.js +245 -0
  49. package/dist/cjs/username.js +290 -0
  50. package/dist/cjs/webauthn.js +71 -0
  51. package/dist/esm/.tsbuildinfo +1 -0
  52. package/dist/esm/accountGraph.js +480 -0
  53. package/dist/esm/agency.js +436 -0
  54. package/dist/esm/browserHub.js +212 -0
  55. package/dist/esm/civic.js +160 -0
  56. package/dist/esm/commonsSignIn.js +56 -0
  57. package/dist/esm/deviceBoot.js +47 -0
  58. package/dist/esm/deviceDirectory.js +186 -0
  59. package/dist/esm/devicePairing.js +135 -0
  60. package/dist/esm/deviceSession.js +161 -0
  61. package/dist/esm/emailAgentContext.js +29 -0
  62. package/dist/esm/followGraph.js +27 -0
  63. package/dist/esm/identity.js +255 -0
  64. package/dist/esm/inboxPush.js +21 -0
  65. package/dist/esm/index.js +172 -0
  66. package/dist/esm/inference/accountBilling.js +331 -0
  67. package/dist/esm/inference/aliaModelRelease.js +259 -0
  68. package/dist/esm/inference/attribution.js +103 -0
  69. package/dist/esm/inference/catalogue.js +484 -0
  70. package/dist/esm/inference/entitlement.js +214 -0
  71. package/dist/esm/inference/errors.js +306 -0
  72. package/dist/esm/inference/identifiers.js +221 -0
  73. package/dist/esm/inference/inbox.js +102 -0
  74. package/dist/esm/inference/modelDocumentation.js +430 -0
  75. package/dist/esm/inference/money.js +185 -0
  76. package/dist/esm/inference/priceVersion.js +107 -0
  77. package/dist/esm/inference/providerConnection.js +452 -0
  78. package/dist/esm/inference/request.js +474 -0
  79. package/dist/esm/inference/routingPolicy.js +315 -0
  80. package/dist/esm/inference/streamEvents.js +255 -0
  81. package/dist/esm/inference/usage.js +326 -0
  82. package/dist/esm/inference/version.js +102 -0
  83. package/dist/esm/keyRecovery.js +88 -0
  84. package/dist/esm/keyRotation.js +72 -0
  85. package/dist/esm/links.js +65 -0
  86. package/dist/esm/moderationReputation.js +295 -0
  87. package/dist/esm/oauth.js +63 -0
  88. package/dist/esm/oxyRecordTypes.js +68 -0
  89. package/dist/esm/protocol.js +50 -0
  90. package/dist/esm/recommendations.js +165 -0
  91. package/dist/esm/reputation.js +293 -0
  92. package/dist/esm/sessionStatus.js +118 -0
  93. package/dist/esm/transparency.js +86 -0
  94. package/dist/esm/updates.js +249 -0
  95. package/dist/esm/userInvalidation.js +85 -0
  96. package/dist/esm/userResponse.js +240 -0
  97. package/dist/esm/username.js +283 -0
  98. package/dist/esm/webauthn.js +68 -0
  99. package/dist/types/.tsbuildinfo +1 -0
  100. package/dist/types/accountGraph.d.ts +378 -0
  101. package/dist/types/agency.d.ts +2162 -0
  102. package/dist/types/browserHub.d.ts +856 -0
  103. package/dist/types/civic.d.ts +338 -0
  104. package/dist/types/commonsSignIn.d.ts +58 -0
  105. package/dist/types/deviceBoot.d.ts +74 -0
  106. package/dist/types/deviceDirectory.d.ts +1317 -0
  107. package/dist/types/devicePairing.d.ts +130 -0
  108. package/dist/types/deviceSession.d.ts +411 -0
  109. package/dist/types/emailAgentContext.d.ts +248 -0
  110. package/dist/types/followGraph.d.ts +150 -0
  111. package/dist/types/identity.d.ts +402 -0
  112. package/dist/types/inboxPush.d.ts +30 -0
  113. package/dist/types/index.d.ts +100 -0
  114. package/dist/types/inference/accountBilling.d.ts +738 -0
  115. package/dist/types/inference/aliaModelRelease.d.ts +609 -0
  116. package/dist/types/inference/attribution.d.ts +176 -0
  117. package/dist/types/inference/catalogue.d.ts +1618 -0
  118. package/dist/types/inference/entitlement.d.ts +519 -0
  119. package/dist/types/inference/errors.d.ts +242 -0
  120. package/dist/types/inference/identifiers.d.ts +182 -0
  121. package/dist/types/inference/inbox.d.ts +374 -0
  122. package/dist/types/inference/modelDocumentation.d.ts +1603 -0
  123. package/dist/types/inference/money.d.ts +185 -0
  124. package/dist/types/inference/priceVersion.d.ts +182 -0
  125. package/dist/types/inference/providerConnection.d.ts +968 -0
  126. package/dist/types/inference/request.d.ts +2800 -0
  127. package/dist/types/inference/routingPolicy.d.ts +616 -0
  128. package/dist/types/inference/streamEvents.d.ts +950 -0
  129. package/dist/types/inference/usage.d.ts +1164 -0
  130. package/dist/types/inference/version.d.ts +102 -0
  131. package/dist/types/keyRecovery.d.ts +138 -0
  132. package/dist/types/keyRotation.d.ts +103 -0
  133. package/dist/types/links.d.ts +96 -0
  134. package/dist/types/moderationReputation.d.ts +487 -0
  135. package/dist/types/oauth.d.ts +86 -0
  136. package/dist/types/oxyRecordTypes.d.ts +62 -0
  137. package/dist/types/protocol.d.ts +86 -0
  138. package/dist/types/recommendations.d.ts +542 -0
  139. package/dist/types/reputation.d.ts +457 -0
  140. package/dist/types/sessionStatus.d.ts +231 -0
  141. package/dist/types/transparency.d.ts +392 -0
  142. package/dist/types/updates.d.ts +545 -0
  143. package/dist/types/userInvalidation.d.ts +94 -0
  144. package/dist/types/userResponse.d.ts +1706 -0
  145. package/dist/types/username.d.ts +265 -0
  146. package/dist/types/webauthn.d.ts +77 -0
  147. package/package.json +87 -0
@@ -0,0 +1,430 @@
1
+ /**
2
+ * Model documentation: what a first-party release must DECLARE, what a
3
+ * downstream developer may READ, and the request that ingests both.
4
+ *
5
+ * Issue #972 §12, the three items under "Future Alia model
6
+ * publication/compliance" that `aliaModelRelease.ts` deliberately left open:
7
+ * accepting the documentation set, publicising the customer-safe half of it, and
8
+ * preserving the metadata an EU AI Act / GPAI documentation workflow needs.
9
+ *
10
+ * ## The section's own scope is what makes the compliance claim falsifiable
11
+ *
12
+ * The issue section is titled "Future Alia model publication/compliance", so the
13
+ * obligations in play are the ones binding a PROVIDER of a general-purpose AI
14
+ * model — Oxy/Alia, for an `alia/*` release it trained or derived. Oxy's position
15
+ * on third-party weights is a different one (it received documentation rather
16
+ * than produced it) with different obligations, and nothing here claims to
17
+ * discharge those. Every field below names the obligation it serves; a field
18
+ * whose obligation could not be named is not here, and two are listed at the
19
+ * bottom as deliberately absent.
20
+ *
21
+ * References are to Regulation (EU) 2024/1689 (the AI Act): Article 50(2)
22
+ * (marking synthetic output), Article 51 (classification as a model with
23
+ * systemic risk), Article 53 (obligations of providers of general-purpose AI
24
+ * models), Article 55 (additional obligations for systemic-risk models), Annex XI
25
+ * (the technical documentation), Annex XII (the information for downstream
26
+ * providers).
27
+ *
28
+ * ## Two shapes, and the Act itself draws the line between them
29
+ *
30
+ * {@link modelGpaiDocumentationSchema} is the whole record. {@link
31
+ * modelDownstreamDocumentationSchema} is the subset served publicly. The split
32
+ * is NOT editorial taste: Annex XI is documentation a provider keeps and
33
+ * provides to the AI Office and national competent authorities on request, while
34
+ * Annex XII is information a provider MAKES AVAILABLE to downstream providers.
35
+ * Training compute, training time, energy consumption and the adversarial-testing
36
+ * report are Annex XI Section 2 and Article 55(1)(a) — the first audience — so
37
+ * they are in the record and not in the public projection, and
38
+ * `db/schema/protectedColumns.ts` says the same thing a second time at the type
39
+ * level.
40
+ *
41
+ * ## The conditionals are the Act's, not a convenience
42
+ *
43
+ * Article 53(2) exempts a model released under a free and open-source licence
44
+ * from 53(1)(a) and 53(1)(b) — the Annex XI and Annex XII sets — UNLESS it is a
45
+ * model with systemic risk. It does not exempt 53(1)(c) or 53(1)(d). So the
46
+ * copyright policy and the training-content summary are required of every
47
+ * release here, while the Annex XI/XII set is required of every release that is
48
+ * not covered by that exemption. Writing it the other way round — everything
49
+ * optional, checked by a human — is what makes a compliance record a field nobody
50
+ * filled in.
51
+ *
52
+ * ## What is deliberately NOT here
53
+ *
54
+ * **The modality and FORMAT of inputs and outputs (Annex XI §1(6), Annex XII
55
+ * §1(b)).** The modality half is already stored, as `inference_models`'
56
+ * `input_modalities` / `output_modalities`. The format half is a property of the
57
+ * Oxy API — one request envelope, one set of endpoints, identical for every model
58
+ * — so a per-model column would record the same value on every row and invite a
59
+ * reader to believe it could differ.
60
+ *
61
+ * **The technical means required for integration (Annex XII §1(c)).** Same
62
+ * reason: for a model served over the Oxy API that is Oxy's own API
63
+ * documentation, not a fact about the weights.
64
+ *
65
+ * **A verification finding for a release signature.** See
66
+ * `aliaModelRelease.ts`: whether a signature checked out is Oxy's finding about
67
+ * the document and not a claim the document makes, and no verifier exists yet
68
+ * because what signs is undecided. The ingestion path stores the signatures and
69
+ * the manifest as received so a verifier that lands later can check them; it
70
+ * records no finding, because there is none.
71
+ *
72
+ * Decided in: docs/adr/0008-catalogue-concept-separation.md, issue #972 §12.
73
+ */
74
+ import { z } from 'zod';
75
+ import { aliaModelReleaseManifestSchema } from './aliaModelRelease.js';
76
+ import { modelCapabilitiesSchema, modelEvaluationResultSchema, modelLicenseSchema, modelProvenanceSchema, modelSafetyMetadataSchema, } from './catalogue.js';
77
+ import { inferenceDateSchema, inferenceHttpsUrlSchema, inferenceTimestampSchema, modelIdSchema, modelReferenceSchema, modelRevisionLabelSchema, sha256DigestSchema, } from './identifiers.js';
78
+ /* -------------------------------------------------------------------------- */
79
+ /* Vocabulary */
80
+ /* -------------------------------------------------------------------------- */
81
+ /**
82
+ * How a release reaches the people who use it — Annex XI §1(4) and Annex XII
83
+ * §1(a), "methods of distribution".
84
+ *
85
+ * TWO members, and both exist today: a release is served through the Oxy API, or
86
+ * its weights are published for download, or both. A third channel is a
87
+ * distribution decision somebody would have to make, and a closed enum gaining a
88
+ * member is a MINOR contract-set change the handshake surfaces (`version.ts`),
89
+ * which is the right amount of ceremony for it.
90
+ *
91
+ * `downloadable_weights` is also what the Article 53(2) free-and-open-source
92
+ * exemption is assessed against — that exemption requires the model to be
93
+ * "released under a free and open-source licence that allows for the access,
94
+ * usage, modification and distribution of the model" — so it is required even
95
+ * where the Annex XI set it belongs to is exempt.
96
+ */
97
+ export const modelDistributionMethodSchema = z.enum(['oxy_api', 'downloadable_weights']);
98
+ /**
99
+ * Whether this is a model with systemic risk, and on what basis — Article 51.
100
+ *
101
+ * Three states, because the two ways a model acquires the classification have
102
+ * different evidence and a record that flattened them could not be checked:
103
+ *
104
+ * - `not_designated` — neither presumed nor designated.
105
+ * - `presumed_by_training_compute` — Article 51(2): the cumulative compute used
106
+ * for training exceeds 10^25 floating point operations, which the Act makes a
107
+ * presumption of high-impact capabilities. The FIGURE is what creates it, so
108
+ * {@link modelGpaiDocumentationSchema} requires the figure alongside this
109
+ * value.
110
+ * - `designated_by_commission` — Article 51(1)(b): a Commission decision, ex
111
+ * officio or following a qualified alert, that the model has capabilities
112
+ * equivalent to the presumption. Not derivable from anything Oxy holds, which
113
+ * is exactly why it is a declared value.
114
+ */
115
+ export const modelSystemicRiskTierSchema = z.enum([
116
+ 'not_designated',
117
+ 'presumed_by_training_compute',
118
+ 'designated_by_commission',
119
+ ]);
120
+ /**
121
+ * Cumulative training compute in floating point operations — Annex XI §2(b).
122
+ *
123
+ * TEXT, in the same spirit as `modelEvaluationResultSchema.score` and for a
124
+ * sharper reason: this is a PUBLISHED figure (`4.2e25`, `2.5e26`), the numbers
125
+ * involved are far outside the exactly-representable integer range, and the
126
+ * value is never arithmetic Oxy performs on a customer's behalf. A JSON number
127
+ * would round it silently and make two records of one published figure compare
128
+ * unequal.
129
+ *
130
+ * The one comparison that IS made — against Article 51(2)'s 10^25 threshold — is
131
+ * a magnitude test, and `Number()` on a string this regex admits is exact enough
132
+ * for a magnitude test by a factor of about 10^9. The refinement that performs
133
+ * it is on {@link modelGpaiDocumentationSchema}.
134
+ */
135
+ export const trainingComputeFlopsSchema = z
136
+ .string()
137
+ .max(40)
138
+ .regex(/^(?:0|[1-9][0-9]*)(?:\.[0-9]+)?(?:e\+?(?:0|[1-9][0-9]?))?$/, 'training compute must be a decimal or scientific figure, e.g. 4.2e25');
139
+ /**
140
+ * Article 51(2)'s presumption threshold, as a number.
141
+ *
142
+ * Named rather than inlined so the refinement that applies it and the enum
143
+ * member that describes it (`presumed_by_training_compute`) cannot come to mean
144
+ * different things.
145
+ */
146
+ export const SYSTEMIC_RISK_COMPUTE_THRESHOLD_FLOPS = 1e25;
147
+ /* -------------------------------------------------------------------------- */
148
+ /* The record */
149
+ /* -------------------------------------------------------------------------- */
150
+ /**
151
+ * The subset of the documentation set that is served to downstream developers —
152
+ * Annex XII, plus the two Article 53(1) items that are public by their own terms.
153
+ *
154
+ * Rides inside {@link modelDocumentationSchema} and inherits its version.
155
+ *
156
+ * Every field here is one a developer integrating the model needs in order to
157
+ * decide whether they may use it and what they must say about it: what it is for,
158
+ * how it is distributed, what it is built out of, where the training-content
159
+ * summary and the copyright policy are, and whether it carries the systemic-risk
160
+ * classification that puts obligations on them too.
161
+ *
162
+ * The optional members are optional for a REASON stated in the parent record's
163
+ * refinement — Article 53(2) — and not because a value may be skipped.
164
+ */
165
+ export const modelDownstreamDocumentationSchema = z
166
+ .object({
167
+ /** Annex XI §1(2), Annex XII §1(a): the tasks the model is intended for. */
168
+ intendedTasks: z.string().min(1).max(2000).optional(),
169
+ /** Annex XI §1(4), Annex XII §1(a). */
170
+ distributionMethods: z.array(modelDistributionMethodSchema).min(1),
171
+ /** Annex XI §1(5), reachable through Annex XII §1(a) ("points 1 to 5"). */
172
+ architecture: z.string().min(1).max(500).optional(),
173
+ /** Annex XI §1(5): the number of parameters. */
174
+ parameterCount: z.number().int().positive().safe().optional(),
175
+ /** Article 53(1)(d): the publicly available summary of training content. */
176
+ trainingDataSummaryUrl: inferenceHttpsUrlSchema,
177
+ /**
178
+ * Article 53(1)(c): the policy for complying with Union copyright law,
179
+ * including the reservation of rights under Article 4(3) of Directive
180
+ * (EU) 2019/790. Required of every release — Article 53(2) does not exempt it.
181
+ */
182
+ copyrightPolicyUrl: inferenceHttpsUrlSchema,
183
+ /** Article 51. */
184
+ systemicRisk: modelSystemicRiskTierSchema,
185
+ /**
186
+ * Whether the release is under a free and open-source licence in the sense
187
+ * of Article 53(2). Distinct from `modelLicenseSchema.commercialUseAllowed`,
188
+ * which answers whether OXY may serve the model — a licence can permit
189
+ * commercial use and still not permit access, modification and
190
+ * redistribution of the weights, and it is the second question the exemption
191
+ * turns on.
192
+ */
193
+ freeAndOpenSourceRelease: z.boolean(),
194
+ })
195
+ .strict();
196
+ /**
197
+ * The whole documentation record for one revision, as ingested.
198
+ *
199
+ * `.strict()`, because this is a compliance record arriving over the wire: a
200
+ * field silently dropped at the parse is a field the record does not contain,
201
+ * and "we accepted your documentation" would then be true of less than was sent.
202
+ *
203
+ * Not versioned on its own — it rides inside
204
+ * {@link modelReleaseIngestionRequestSchema} on the way in and inside
205
+ * {@link modelDocumentationSchema} on the way out, and inherits whichever
206
+ * message carries it.
207
+ */
208
+ export const modelGpaiDocumentationSchema = modelDownstreamDocumentationSchema
209
+ .extend({
210
+ /** Annex XI §2(b): the computational resources used for training. */
211
+ trainingComputeFlops: trainingComputeFlopsSchema.optional(),
212
+ /** Annex XI §2(b): the training time. */
213
+ trainingTimeHours: z.number().positive().safe().optional(),
214
+ /**
215
+ * Annex XI §2(c): the known or ESTIMATED energy consumption. The Act asks
216
+ * for an estimate where the figure is not known, so absence here means the
217
+ * Annex XI set is exempt rather than that the number was hard to obtain.
218
+ */
219
+ energyConsumptionMwh: z.number().nonnegative().safe().optional(),
220
+ /**
221
+ * Article 55(1)(a): the model evaluation, including adversarial testing,
222
+ * performed for a model with systemic risk. A pointer, like every other
223
+ * document reference here — the catalogue holds no report.
224
+ */
225
+ adversarialTestingReportUrl: inferenceHttpsUrlSchema.optional(),
226
+ })
227
+ .strict()
228
+ .superRefine((documentation, ctx) => {
229
+ // Article 53(2): the free-and-open-source exemption from 53(1)(a) and (b)
230
+ // does not apply to a model with systemic risk. So the Annex XI/XII set is
231
+ // required of everything else, and the ONE state that may omit it is a
232
+ // free-and-open-source release that is not designated.
233
+ const exempt = documentation.freeAndOpenSourceRelease && documentation.systemicRisk === 'not_designated';
234
+ if (!exempt) {
235
+ const required = [
236
+ 'intendedTasks',
237
+ 'architecture',
238
+ 'parameterCount',
239
+ 'trainingTimeHours',
240
+ 'energyConsumptionMwh',
241
+ ];
242
+ for (const field of required) {
243
+ if (documentation[field] === undefined) {
244
+ ctx.addIssue({
245
+ code: z.ZodIssueCode.custom,
246
+ path: [field],
247
+ message: 'required by Annex XI unless the Article 53(2) free-and-open-source exemption applies, which it does not for this release',
248
+ });
249
+ }
250
+ }
251
+ }
252
+ // The presumption IS the compute figure (Article 51(2)). Declaring the tier
253
+ // without the figure asserts a threshold was crossed while withholding the
254
+ // only thing that says so.
255
+ if (documentation.systemicRisk === 'presumed_by_training_compute' &&
256
+ documentation.trainingComputeFlops === undefined) {
257
+ ctx.addIssue({
258
+ code: z.ZodIssueCode.custom,
259
+ path: ['trainingComputeFlops'],
260
+ message: 'a systemic-risk presumption under Article 51(2) is the training-compute figure; declare it',
261
+ });
262
+ }
263
+ // The other direction, which is the one that matters: a release whose own
264
+ // declared compute is past the threshold cannot also declare that no
265
+ // classification applies. Without this the field pair would let the record
266
+ // contradict itself and still parse.
267
+ if (documentation.trainingComputeFlops !== undefined &&
268
+ documentation.systemicRisk === 'not_designated' &&
269
+ Number(documentation.trainingComputeFlops) >= SYSTEMIC_RISK_COMPUTE_THRESHOLD_FLOPS) {
270
+ ctx.addIssue({
271
+ code: z.ZodIssueCode.custom,
272
+ path: ['systemicRisk'],
273
+ message: 'training compute at or above 10^25 FLOP is presumed to be a model with systemic risk under Article 51(2)',
274
+ });
275
+ }
276
+ // Article 55(1)(a) applies to every model with systemic risk, however it
277
+ // acquired the classification.
278
+ if (documentation.systemicRisk !== 'not_designated' &&
279
+ documentation.adversarialTestingReportUrl === undefined) {
280
+ ctx.addIssue({
281
+ code: z.ZodIssueCode.custom,
282
+ path: ['adversarialTestingReportUrl'],
283
+ message: 'a model with systemic risk documents its evaluation including adversarial testing (Article 55(1)(a))',
284
+ });
285
+ }
286
+ });
287
+ /* -------------------------------------------------------------------------- */
288
+ /* Ingestion */
289
+ /* -------------------------------------------------------------------------- */
290
+ /**
291
+ * What OXY states about the model line a release belongs to.
292
+ *
293
+ * A signed release manifest carries a revision, a licence, a provenance block,
294
+ * evaluations, safety metadata and an artifact inventory. It carries no
295
+ * CAPABILITY SHEET — no modalities, no `maxContextTokens`, none of the
296
+ * tool/streaming flags — and every one of those is required to create a model
297
+ * line at all.
298
+ *
299
+ * That is not a gap in the manifest. A capability sheet is a statement about what
300
+ * the Oxy API will serve, which is Oxy's to make and not the signer's: the same
301
+ * weights behind a different gateway answer a different set of these questions.
302
+ * So it travels beside the manifest, like the documentation record, and the
303
+ * signature keeps covering exactly the document its signer wrote.
304
+ *
305
+ * Ignored when the model line already exists — a release does not edit a model.
306
+ * The licence and provenance in the MANIFEST are checked against the stored ones
307
+ * instead, because those are claims about somebody's rights rather than Oxy's own
308
+ * editorial choices.
309
+ */
310
+ export const modelLineDeclarationSchema = z
311
+ .object({
312
+ displayName: z.string().min(1).max(200),
313
+ description: z.string().max(4000).optional(),
314
+ capabilities: modelCapabilitiesSchema,
315
+ knowledgeCutoff: inferenceDateSchema.optional(),
316
+ releasedOn: inferenceDateSchema.optional(),
317
+ })
318
+ .strict();
319
+ /**
320
+ * The body of the release-ingestion request.
321
+ *
322
+ * The documentation and the capability sheet travel BESIDE the manifest rather
323
+ * than inside it, and that is the whole reason this wrapper exists.
324
+ * `aliaModelReleaseManifestSchema` is a SIGNED document: adding a field to it
325
+ * would change the bytes a signer covers and the version the data plane and Alia
326
+ * compile against, for records that are Oxy's own rather than the signer's.
327
+ * Keeping them separate means the signature still covers exactly what it covered.
328
+ *
329
+ * A signer that later chooses to cover the documentation too can: it would
330
+ * become a second signed document with its own manifest, which is a contract
331
+ * addition rather than a change to this one.
332
+ */
333
+ export const modelReleaseIngestionRequestSchema = z
334
+ .object({
335
+ /** See `version.ts`: an ingestion payload is a whole message on the wire. */
336
+ schemaVersion: z.literal(1),
337
+ manifest: aliaModelReleaseManifestSchema,
338
+ gpaiDocumentation: modelGpaiDocumentationSchema,
339
+ model: modelLineDeclarationSchema,
340
+ })
341
+ .strict();
342
+ /**
343
+ * What ingestion reports back.
344
+ *
345
+ * COUNTS for the artifacts and signatures rather than echoing them: the caller
346
+ * sent them and the interesting fact is that all of them landed. Echoing a
347
+ * signature would also make this response a place a credential-shaped value gets
348
+ * logged, for no gain.
349
+ *
350
+ * No verification field — see this module's header, and `aliaModelRelease.ts`.
351
+ */
352
+ export const modelReleaseIngestionResultSchema = z
353
+ .object({
354
+ /** See `version.ts`: served on its own, so it is versioned. */
355
+ schemaVersion: z.literal(1),
356
+ releaseId: z.string().min(1).max(128),
357
+ modelId: modelIdSchema,
358
+ revision: modelRevisionLabelSchema,
359
+ reference: modelReferenceSchema,
360
+ /** Whether this request created the release, or found it already ingested. */
361
+ outcome: z.enum(['ingested', 'already_ingested']),
362
+ artifactCount: z.number().int().positive().safe(),
363
+ signatureCount: z.number().int().positive().safe(),
364
+ evaluationCount: z.number().int().nonnegative().safe(),
365
+ ingestedAt: inferenceTimestampSchema,
366
+ })
367
+ .strict();
368
+ /* -------------------------------------------------------------------------- */
369
+ /* The customer-safe documentation view */
370
+ /* -------------------------------------------------------------------------- */
371
+ /**
372
+ * The documentation for ONE revision, as a downstream developer reads it.
373
+ *
374
+ * Revision-scoped, and that is the point of it existing beside
375
+ * `modelCatalogueEntrySchema`. The catalogue entry carries the documentation of
376
+ * whichever revision is CURRENT, so a customer who pinned
377
+ * `<publisher>/<model>@<revision>` — which the catalogue invites, and which the
378
+ * immutability trigger on `inference_model_revisions` exists to make meaningful —
379
+ * had no way to read the model card, evaluations or safety metadata of the
380
+ * revision they are actually calling. A model card that only describes the
381
+ * newest weights is the exact conflation ADR 0008 separates revisions to prevent.
382
+ *
383
+ * `license` and `provenance` are the MODEL's, repeated here rather than linked,
384
+ * for the same reason `modelCatalogueEntrySchema` repeats its fields: a
385
+ * projection that nests the operational descriptors is one accident of nesting
386
+ * away from serving an internal identifier.
387
+ */
388
+ export const modelDocumentationSchema = z
389
+ .object({
390
+ /** See `version.ts`: this is a public response shape. */
391
+ schemaVersion: z.literal(1),
392
+ modelId: modelIdSchema,
393
+ revision: modelRevisionLabelSchema,
394
+ /** The exact string a customer pins. */
395
+ reference: modelReferenceSchema,
396
+ /** Whether a bare `<publisher>/<model>` resolves to this revision today. */
397
+ isCurrentRevision: z.boolean(),
398
+ releasedAt: inferenceTimestampSchema,
399
+ retiredAt: inferenceTimestampSchema.optional(),
400
+ modelCardUrl: inferenceHttpsUrlSchema.optional(),
401
+ /**
402
+ * The digest of the served artifact, where Oxy hosts the weights.
403
+ *
404
+ * Customer-safe, deliberately: it is the one field on this view that lets a
405
+ * developer check that the weights they were handed are the weights the
406
+ * documentation describes, and a digest discloses nothing but the identity of
407
+ * bytes Oxy is already serving them.
408
+ */
409
+ artifactDigest: sha256DigestSchema.optional(),
410
+ license: modelLicenseSchema,
411
+ provenance: modelProvenanceSchema,
412
+ evaluations: z.array(modelEvaluationResultSchema).default([]),
413
+ safety: modelSafetyMetadataSchema.optional(),
414
+ /** Absent for a revision with no documentation record — i.e. every one Oxy did not release. */
415
+ gpai: modelDownstreamDocumentationSchema.optional(),
416
+ })
417
+ .strict()
418
+ .superRefine((documentation, ctx) => {
419
+ // The same check `modelRevisionSchema` makes, and it is load-bearing for a
420
+ // different reason here: this view exists so a customer can read the
421
+ // documentation of the revision they PINNED, so a reference that resolves
422
+ // elsewhere would attach a model card to weights nobody is calling.
423
+ if (documentation.reference !== `${documentation.modelId}@${documentation.revision}`) {
424
+ ctx.addIssue({
425
+ code: z.ZodIssueCode.custom,
426
+ path: ['reference'],
427
+ message: 'reference must be exactly <modelId>@<revision>',
428
+ });
429
+ }
430
+ });
@@ -0,0 +1,185 @@
1
+ /**
2
+ * Money and usage units for the inference contracts.
3
+ *
4
+ * The non-negotiable invariant this file exists to make structural: **customer
5
+ * charges never use floating-point values as the financial source of truth.**
6
+ * A JS `number` cannot represent `0.1 + 0.2` exactly, and an inference ledger
7
+ * adds millions of small amounts, so a float total is wrong by construction
8
+ * rather than by accident.
9
+ *
10
+ * Every amount and every price is one representation: {@link exactDecimalSchema},
11
+ * an exact decimal STRING at the scale ADR 0009 declares. Amounts are not
12
+ * integer minor units, and that is the ADR's decision rather than an oversight:
13
+ * one token costs several orders of magnitude less than one cent, so rounding
14
+ * per request would make a customer's bill depend on how their client chunked
15
+ * its work. Rounding happens ONCE, at the invoice boundary, and is itself a
16
+ * ledger entry.
17
+ *
18
+ * A string is also what the driver hands back — `postgres.js` decodes `NUMERIC`
19
+ * as a string — so keeping it a string on the wire means accidental JS
20
+ * arithmetic fails loudly instead of silently losing precision. Money
21
+ * arithmetic happens in SQL or in a decimal type, never in a JS `number`.
22
+ *
23
+ * Units are carried separately from money in every shape: a receipt says both
24
+ * "204 output tokens" and "0.003060000000 USD", and neither is derived from the
25
+ * other at read time. That separation is what lets a price version change
26
+ * without rewriting settled history.
27
+ *
28
+ * Decided in: docs/adr/0009-usage-reservation-and-settlement.md.
29
+ */
30
+ import { z } from 'zod';
31
+ /* -------------------------------------------------------------------------- */
32
+ /* Money */
33
+ /* -------------------------------------------------------------------------- */
34
+ /** ISO 4217 alpha-3 currency code, e.g. `USD`. */
35
+ export const currencyCodeSchema = z
36
+ .string()
37
+ .regex(/^[A-Z]{3}$/, 'currency must be an ISO 4217 alpha-3 code');
38
+ /**
39
+ * The declared fractional scale of every amount and price in this contract, and
40
+ * of the `NUMERIC` columns the ledger stores them in.
41
+ *
42
+ * Twelve digits is sub-minor-unit precision by a wide margin: at $3 per million
43
+ * input tokens, one token costs `0.000003000000`, which this scale represents
44
+ * exactly. Amounts are compared and summed NUMERICALLY, never as text — `3.0`
45
+ * and `3.000000000000` are one amount written two ways.
46
+ */
47
+ export const INFERENCE_MONEY_SCALE = 12;
48
+ /**
49
+ * An exact non-negative decimal, carried as a STRING so no parse step can turn
50
+ * it into a float on the way past. Up to 18 integer digits and
51
+ * {@link INFERENCE_MONEY_SCALE} fractional digits.
52
+ *
53
+ * Non-negative: direction is carried by the SHAPE — a receipt debits, a refund
54
+ * credits — so a stray sign can never silently invert an entry.
55
+ *
56
+ * No exponent form: `1e-6` and `0.000001` are the same number, but only one of
57
+ * them survives a naive string comparison, a cache key or a log grep intact.
58
+ *
59
+ * Branded, so a bare `string` is not assignable and an amount cannot arrive
60
+ * from string concatenation that was never checked. Producers construct one
61
+ * with `exactDecimalSchema.parse(value)`.
62
+ */
63
+ export const exactDecimalSchema = z
64
+ .string()
65
+ .regex(/^(?:0|[1-9][0-9]{0,17})(?:\.[0-9]{1,12})?$/, 'must be an exact non-negative decimal string without an exponent')
66
+ .brand();
67
+ /**
68
+ * An amount of money: the exact decimal plus the currency it is in.
69
+ *
70
+ * `.strict()` so a payload carrying a convenience float beside the exact value
71
+ * (`{ amount: '18.06', amountFloat: 18.06 }`) is REJECTED rather than stripped.
72
+ * A stripped float is the more dangerous outcome: it disappears silently here
73
+ * and survives in the producer, where it is the value somebody eventually
74
+ * displays.
75
+ */
76
+ export const moneySchema = z
77
+ .object({
78
+ amount: exactDecimalSchema,
79
+ currency: currencyCodeSchema,
80
+ })
81
+ .strict();
82
+ /* -------------------------------------------------------------------------- */
83
+ /* Usage units */
84
+ /* -------------------------------------------------------------------------- */
85
+ /**
86
+ * The closed set of units inference is metered in.
87
+ *
88
+ * **The units PARTITION a request: every unit counts material no other unit
89
+ * counts.** `cached_input_tokens` is not part of `input_tokens`, and
90
+ * `reasoning_tokens` is not part of `output_tokens` — they are siblings, not
91
+ * subsets. A request whose 10 000-token prompt was served 9 000 tokens from
92
+ * cache is reported as `input_tokens: 1000` beside `cached_input_tokens: 9000`,
93
+ * never as `input_tokens: 10000` beside it.
94
+ *
95
+ * That belongs to the definition rather than to a convention somewhere else,
96
+ * because settlement applies a price to EVERY reported unit and sums them
97
+ * (`inferenceLedger.service.ts`'s `computeCharge`). Under the partition rule
98
+ * that sum IS the request's cost, and a cached token can carry its own — lower
99
+ * — price. Under the nested reading the same sum charges the cached and
100
+ * reasoning tokens twice: once inside their parent and once on their own line.
101
+ * It fails silently, because every total still looks plausible and the receipt
102
+ * is still internally consistent, and on a reasoning model the reasoning tokens
103
+ * can dominate the completion, so the error is not marginal.
104
+ *
105
+ * **Every OpenAI-compatible provider reports the other way round**:
106
+ * `prompt_tokens` INCLUDES `prompt_tokens_details.cached_tokens`, and
107
+ * `completion_tokens` INCLUDES `completion_tokens_details.reasoning_tokens`.
108
+ * Normalising is the data plane's job and it is subtraction:
109
+ *
110
+ * ```text
111
+ * input_tokens = prompt_tokens - prompt_tokens_details.cached_tokens
112
+ * output_tokens = completion_tokens - completion_tokens_details.reasoning_tokens
113
+ * ```
114
+ *
115
+ * No refinement in this package can enforce it, and saying so is part of the
116
+ * rule: a nested report and a disjoint one are the same four non-negative
117
+ * integers, so no predicate over a single report can tell them apart. The two
118
+ * structural guards that DO exist — refining `cached <= input` and
119
+ * `reasoning <= output`, or deriving the parents instead of reporting them —
120
+ * both encode the nested reading, which is the one this rule rejects. What IS
121
+ * enforceable is the arithmetic that depends on the rule, and that is where the
122
+ * enforcement lives — `inferenceLedger.service.test.ts` prices a report in
123
+ * which cached and reasoning tokens are both non-zero and asserts the exact
124
+ * total, which the nested reading cannot produce.
125
+ *
126
+ * Where the public surface has to speak a nested dialect, the sum is put back
127
+ * at the boundary rather than the internal reading being bent to it
128
+ * (`routes/inferenceEdge.ts` renders `prompt_tokens` as
129
+ * `input_tokens + cached_input_tokens`).
130
+ *
131
+ * Time is carried in integer MILLISECONDS rather than seconds so that no unit
132
+ * quantity is ever fractional: a 12.5-second transcription is `12500`, exactly,
133
+ * and the "units are integers" rule holds for every modality instead of holding
134
+ * for tokens and being quietly broken by audio.
135
+ */
136
+ export const USAGE_UNITS = [
137
+ 'input_tokens',
138
+ 'cached_input_tokens',
139
+ 'output_tokens',
140
+ 'reasoning_tokens',
141
+ 'requests',
142
+ 'images',
143
+ 'audio_input_milliseconds',
144
+ 'audio_output_milliseconds',
145
+ 'video_milliseconds',
146
+ 'characters',
147
+ 'embeddings',
148
+ ];
149
+ export const usageUnitSchema = z.enum(USAGE_UNITS);
150
+ /**
151
+ * A metered quantity of ONE unit. Never money — a quantity carries no price and
152
+ * no currency, so a consumer cannot mistake a token count for an amount owed.
153
+ */
154
+ export const usageQuantitySchema = z
155
+ .object({
156
+ unit: usageUnitSchema,
157
+ quantity: z.number().int().nonnegative().safe(),
158
+ })
159
+ .strict();
160
+ /**
161
+ * Where a metered quantity came from.
162
+ *
163
+ * Kept explicit because the three are not interchangeable when a charge is
164
+ * disputed: `provider_reported` is the upstream's own count, `oxy_measured` is
165
+ * counted by the platform (streamed bytes, wall-clock milliseconds), and
166
+ * `estimated` is a reconstruction used when a provider returned no usage at
167
+ * all. An estimate that is indistinguishable from a reported number is an
168
+ * estimate nobody can later reconcile or refund against.
169
+ */
170
+ export const USAGE_SOURCES = ['provider_reported', 'oxy_measured', 'estimated'];
171
+ export const usageSourceSchema = z.enum(USAGE_SOURCES);
172
+ /**
173
+ * A price for one unit, as `amount` per `per` units — `per` because a price
174
+ * quoted per single token would need more fractional digits than it is worth
175
+ * ("$3.00 per 1000000 input_tokens" is how every provider quotes it, and how
176
+ * every customer reads it).
177
+ */
178
+ export const unitPriceSchema = z
179
+ .object({
180
+ unit: usageUnitSchema,
181
+ amount: exactDecimalSchema,
182
+ per: z.number().int().positive().safe(),
183
+ currency: currencyCodeSchema,
184
+ })
185
+ .strict();