@adobe/spacecat-shared-data-access 4.20.0 → 4.21.1

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/CHANGELOG.md CHANGED
@@ -1,3 +1,15 @@
1
+ ## [@adobe/spacecat-shared-data-access-v4.21.1](https://github.com/adobe/spacecat-shared/compare/@adobe/spacecat-shared-data-access-v4.21.0...@adobe/spacecat-shared-data-access-v4.21.1) (2026-08-14)
2
+
3
+ ### Bug Fixes
4
+
5
+ * **data-access:** publish brand semrushWorkspaceId mirror removal (SITES-49202) ([#1880](https://github.com/adobe/spacecat-shared/issues/1880)) ([80a296a](https://github.com/adobe/spacecat-shared/commit/80a296af1b4f49ffa66aba80d677a5e6b42f3856)), closes [#1867](https://github.com/adobe/spacecat-shared/issues/1867) [#1867](https://github.com/adobe/spacecat-shared/issues/1867) [#1867](https://github.com/adobe/spacecat-shared/issues/1867)
6
+
7
+ ## [@adobe/spacecat-shared-data-access-v4.21.0](https://github.com/adobe/spacecat-shared/compare/@adobe/spacecat-shared-data-access-v4.20.0...@adobe/spacecat-shared-data-access-v4.21.0) (2026-08-13)
8
+
9
+ ### Features
10
+
11
+ * **data-access:** paginated sites query filtered by entitlement tier/productCode ([#1877](https://github.com/adobe/spacecat-shared/issues/1877)) ([00e5e59](https://github.com/adobe/spacecat-shared/commit/00e5e592f1d4990e1578deb8287e84f3f0d0cf1b))
12
+
1
13
  ## [@adobe/spacecat-shared-data-access-v4.20.0](https://github.com/adobe/spacecat-shared/compare/@adobe/spacecat-shared-data-access-v4.19.0...@adobe/spacecat-shared-data-access-v4.20.0) (2026-08-12)
2
14
 
3
15
  ### Features
package/package.json CHANGED
@@ -1,6 +1,6 @@
1
1
  {
2
2
  "name": "@adobe/spacecat-shared-data-access",
3
- "version": "4.20.0",
3
+ "version": "4.21.1",
4
4
  "description": "Shared modules of the Spacecat Services - Data Access",
5
5
  "type": "module",
6
6
  "engines": {
@@ -15,8 +15,9 @@ import BaseCollection from '../base/base.collection.js';
15
15
  /**
16
16
  * BrandCollection - collection of Brand rows. `findById` (the brand UUID) is
17
17
  * the primary access path for serenity dual-mode resolution;
18
- * `findBySemrushWorkspaceId` (from the addAllIndex below) is used to repair a
19
- * pointer when a sub-workspace is found to have been deleted out-of-band.
18
+ * `findBySemrushSubWorkspaceId` (from the addAllIndex on the schema) is used to
19
+ * repair a pointer when a sub-workspace is found to have been deleted
20
+ * out-of-band.
20
21
  *
21
22
  * @class BrandCollection
22
23
  * @extends BaseCollection
@@ -16,17 +16,24 @@ import BaseModel from '../base/base.model.js';
16
16
  * Brand - an Adobe brand, stored in the `brands` table in mysticat-data-service
17
17
  * and served over PostgREST. Intentionally minimal: it surfaces only the fields
18
18
  * the serenity sub-workspace provisioning flows read/write
19
- * (`semrushWorkspaceId`, `status`, `name`). Brands are created and fully
19
+ * (`semrushSubWorkspaceId`, `status`, `name`). Brands are created and fully
20
20
  * managed elsewhere (Brandalf sync, onboarding); this entity is a read +
21
21
  * targeted-patch surface, not a create surface.
22
22
  *
23
- * `semrushWorkspaceId` is the dual-mode switch: NULL = the brand is not
23
+ * `semrushSubWorkspaceId` is the dual-mode switch: NULL = the brand is not
24
24
  * connected to a Semrush sub-workspace (resolves against the org parent
25
25
  * workspace — "flat" mode); set = the brand has its own Semrush sub-workspace.
26
26
  * Deactivation empties the sub-workspace and clears this pointer (the
27
27
  * sub-workspace itself is never deleted). See serenity-docs
28
28
  * brand-semrush-provisioning-v2-phase1-sync.md §6.
29
29
  *
30
+ * NOTE: there is no brand-level `semrushWorkspaceId` accessor. The deprecated
31
+ * read-only mirror (attribute, index, `findBySemrushWorkspaceId`,
32
+ * `allBySemrushWorkspaceId`, `setSemrushWorkspaceId`) was removed in SITES-49202;
33
+ * `semrushSubWorkspaceId` above is the write-of-record. The identically-named
34
+ * `Organization.semrushWorkspaceId` is a DISTINCT field and stays — do not
35
+ * reintroduce a brand mirror by symbol-name sweep.
36
+ *
30
37
  * @class Brand
31
38
  * @extends BaseModel
32
39
  */
@@ -39,23 +46,6 @@ class Brand extends BaseModel {
39
46
  * `pending`; customer offboard writes `deleted`.
40
47
  */
41
48
  static STATUSES = Object.freeze(['pending', 'active', 'deleted', 'ignored']);
42
-
43
- /**
44
- * Deprecated BC-compat setter. `semrushWorkspaceId` is `readOnly: true` in
45
- * the schema (mirrored by the mysticat-data-service sync trigger), so no
46
- * setter is auto-generated for it — this manual method exists purely so an
47
- * existing external caller of `setSemrushWorkspaceId` does not get a
48
- * semver-breaking runtime error on upgrade. Delegates to the real
49
- * write-of-record attribute. Remove once every direct caller has migrated
50
- * to `setSemrushSubWorkspaceId` (see brand.schema.js).
51
- *
52
- * @deprecated Use setSemrushSubWorkspaceId instead.
53
- * @param {string|null} value
54
- * @returns {Brand}
55
- */
56
- setSemrushWorkspaceId(value) {
57
- return this.setSemrushSubWorkspaceId(value);
58
- }
59
49
  }
60
50
 
61
51
  export default Brand;
@@ -28,37 +28,23 @@ const schema = new SchemaBuilder(Brand, BrandCollection)
28
28
  })
29
29
  // reference_status enum on the brands table. Not `required`: this entity
30
30
  // never creates a brand, and a targeted PATCH (e.g. setting only
31
- // semrushWorkspaceId) must not be forced to also send status. The validator
31
+ // semrushSubWorkspaceId) must not be forced to also send status. The validator
32
32
  // still rejects an out-of-enum value when status IS written
33
33
  // (activate → 'active', deactivate → 'pending').
34
34
  .addAttribute('status', {
35
35
  type: Brand.STATUSES,
36
36
  validate: (value) => value == null || Brand.STATUSES.includes(value),
37
37
  })
38
- // DEPRECATED (serenity-docs brand-semrush-mapping-maintenance.md §10
39
- // rename, write-of-record cutover): read-only mirror of
40
- // semrushSubWorkspaceId below, maintained entirely by the
41
- // mysticat-data-service brands_sync_semrush_workspace_id trigger
42
- // (migration 20260702094229). No schema-generated setter — app code must
43
- // write semrushSubWorkspaceId instead. brand.model.js still defines a
44
- // manual, deprecated setSemrushWorkspaceId() that delegates to
45
- // setSemrushSubWorkspaceId(), so an existing external caller of the old
46
- // setter is not broken (a bare readOnly flip here would be a semver-breaking
47
- // removal for any @adobe/spacecat-shared-data-access consumer). Will be
48
- // retired (attribute, column, and trigger) once every direct external
49
- // reader has migrated.
50
- .addAttribute('semrushWorkspaceId', {
51
- type: 'string',
52
- readOnly: true,
53
- })
54
38
  // Brand → Semrush sub-workspace. Nullable (NULL = no sub-workspace
55
- // connected). Same minimum guard as organizations.semrushWorkspaceId: the
56
- // shared `hasText` rejects the empty string (and non-strings) while letting
39
+ // connected). Same minimum guard the distinct `Organization` entity applies
40
+ // to its own `semrushWorkspaceId` field (the brand has no such field — its
41
+ // deprecated mirror was removed in SITES-49202): the shared `hasText` rejects
42
+ // the empty string (and non-strings) while letting
57
43
  // null/undefined short-circuit. Note hasText does NOT trim, so a
58
44
  // whitespace-only value would pass — acceptable here because this column is
59
45
  // only ever written by the activate flow with a real Semrush workspace UUID,
60
- // never user input. This is now the write-of-record (see semrushWorkspaceId
61
- // above for the deprecated BC mirror).
46
+ // never user input. This is the write-of-record for the brand → Semrush
47
+ // sub-workspace pointer.
62
48
  .addAttribute('semrushSubWorkspaceId', {
63
49
  type: 'string',
64
50
  validate: (value) => value == null || hasText(value),
@@ -72,15 +58,9 @@ const schema = new SchemaBuilder(Brand, BrandCollection)
72
58
  type: 'any',
73
59
  validate: (value) => value == null || (typeof value === 'object' && !Array.isArray(value)),
74
60
  })
75
- // Uniqueness is enforced at the DB level via the UNIQUE constraint on the
76
- // deprecated brands.semrush_workspace_id (mysticat-data-service migration
77
- // 20260615102123), so findBySemrushWorkspaceId returns at most one row.
78
- // Kept for BC lookups against the mirrored column; new code should prefer
79
- // findBySemrushSubWorkspaceId below.
80
- .addAllIndex(['semrushWorkspaceId'])
81
- // Same uniqueness guarantee on the write-of-record column
61
+ // Uniqueness guarantee on the write-of-record column
82
62
  // (brands.semrush_sub_workspace_id, mysticat-data-service migration
83
- // 20260702091920).
63
+ // 20260702091920), so findBySemrushSubWorkspaceId returns at most one row.
84
64
  .addAllIndex(['semrushSubWorkspaceId']);
85
65
 
86
66
  export default schema.build();
@@ -17,24 +17,13 @@ import type {
17
17
  export interface Brand extends BaseModel {
18
18
  getName(): string;
19
19
  getStatus(): string;
20
- // Deprecated BC mirror (brands.semrush_workspace_id), maintained by the
21
- // mysticat-data-service sync trigger. No schema-generated setter; the
22
- // deprecated setSemrushWorkspaceId below is a manual delegate defined in
23
- // brand.model.js, kept only for backward compatibility. Use
24
- // getSemrushSubWorkspaceId/setSemrushSubWorkspaceId instead. See
25
- // brand.schema.js.
26
- getSemrushWorkspaceId(): string | null;
27
20
  getSemrushSubWorkspaceId(): string | null;
28
21
  setName(value: string): Brand;
29
22
  setStatus(value: string): Brand;
30
23
  setSemrushSubWorkspaceId(value: string | null): Brand;
31
- /** @deprecated Use setSemrushSubWorkspaceId instead. */
32
- setSemrushWorkspaceId(value: string | null): Brand;
33
24
  }
34
25
 
35
26
  export interface BrandCollection extends BaseCollection<Brand> {
36
- allBySemrushWorkspaceId(semrushWorkspaceId: string): Promise<Brand[]>;
37
- findBySemrushWorkspaceId(semrushWorkspaceId: string): Promise<Brand | null>;
38
27
  allBySemrushSubWorkspaceId(semrushSubWorkspaceId: string): Promise<Brand[]>;
39
28
  findBySemrushSubWorkspaceId(semrushSubWorkspaceId: string): Promise<Brand | null>;
40
29
  }
@@ -14,6 +14,13 @@ import { hasText, isValidHelixPreviewUrl, isValidUrl } from '@adobe/spacecat-sha
14
14
 
15
15
  import DataAccessError from '../../errors/data-access.error.js';
16
16
  import BaseCollection from '../base/base.collection.js';
17
+ import {
18
+ applyWhere,
19
+ decodeCursor,
20
+ DEFAULT_PAGE_SIZE,
21
+ encodeCursor,
22
+ toDbField,
23
+ } from '../../util/postgrest.utils.js';
17
24
 
18
25
  import Site, { AEM_CS_HOST, getAuthoringType } from './site.model.js';
19
26
 
@@ -206,6 +213,133 @@ class SiteCollection extends BaseCollection {
206
213
  return sites;
207
214
  }
208
215
 
216
+ /**
217
+ * Returns sites filtered by entitlement tier and/or product code, paginated,
218
+ * and composable with a caller-supplied `where` and `orderBy` — all in a
219
+ * SINGLE PostgREST query via a two-level inner-join embed
220
+ * (sites -> site_enrollments -> entitlements).
221
+ *
222
+ * PostgREST resource embedding NESTS matching children under each parent row
223
+ * (it does not flatten into a cartesian product), so `.range()` offset
224
+ * pagination stays correct over DISTINCT parent sites even when a site has
225
+ * multiple matching enrollments. `!inner` at both embed levels turns the
226
+ * embedded filter into an INNER JOIN, excluding sites with no matching
227
+ * enrollment/entitlement (a plain embedded filter is a LEFT JOIN on older
228
+ * PostgREST servers, which would return non-matching sites with an empty
229
+ * embed instead of dropping them).
230
+ *
231
+ * Kept symmetric with `Site.all(..., { returnCursor: true, limit })`: it
232
+ * honors the EXACT `limit` passed (no silent cap, no +1) because the
233
+ * api-service caller does N+1 hasMore detection (passes `limit =
234
+ * effectiveLimit + 1` and slices), and returns `{ data, cursor }` when
235
+ * `returnCursor` is set, else a bare array. This lets the controller call it
236
+ * almost identically to `Site.all`.
237
+ *
238
+ * @param {object} [filter]
239
+ * @param {string} [filter.tier] - Entitlement tier (e.g. 'PAID', 'FREE_TRIAL').
240
+ * @param {string} [filter.productCode] - Entitlement product code (e.g. 'LLMO').
241
+ * At least one of `tier` / `productCode` is required.
242
+ * @param {object} [options]
243
+ * @param {Function} [options.where] - Caller `where` fn `(attrs, op) => expr`
244
+ * applied to the sites table (e.g. baseURL substring / deliveryType / isLive).
245
+ * @param {object} [options.orderBy] - `{ attribute, direction }`; defaults to
246
+ * `updatedAt` desc. Always followed by an `id` tiebreaker for stable paging.
247
+ * @param {number} [options.limit] - Max parent rows to return (exact). A
248
+ * non-positive or non-integer value falls back to DEFAULT_PAGE_SIZE.
249
+ * @param {string} [options.cursor] - Base64 offset cursor (see decodeCursor).
250
+ * @param {boolean} [options.returnCursor] - Return `{ data, cursor }` shape.
251
+ * @returns {Promise<Site[] | { data: Site[], cursor: string|null }>}
252
+ */
253
+ async allByEnrollmentFiltered(
254
+ { tier, productCode } = {},
255
+ {
256
+ where, orderBy, limit, cursor, returnCursor,
257
+ } = {},
258
+ ) {
259
+ if (!hasText(tier) && !hasText(productCode)) {
260
+ throw new DataAccessError('tier or productCode is required', this);
261
+ }
262
+
263
+ // The embed must be selected for PostgREST to filter on it. The nested
264
+ // array PostgREST returns per site is stripped below before hydrating the
265
+ // Site model (the Site model has no such attribute).
266
+ const select = '*, site_enrollments!inner(entitlements!inner(tier, product_code))';
267
+
268
+ let query = this.postgrestService
269
+ .from(this.tableName)
270
+ .select(select);
271
+
272
+ if (hasText(tier)) {
273
+ query = query.eq('site_enrollments.entitlements.tier', tier);
274
+ }
275
+ if (hasText(productCode)) {
276
+ query = query.eq('site_enrollments.entitlements.product_code', productCode);
277
+ }
278
+
279
+ // Caller-supplied where composes on the sites table
280
+ // (base_url / delivery_type / is_live).
281
+ query = applyWhere(query, where, this.fieldMaps.toDbMap);
282
+
283
+ // Ordering mirrors base.collection #queryPage: a primary sort (default
284
+ // updatedAt desc) plus a stable id tiebreaker so page boundaries never
285
+ // straddle equal sort keys. An explicit orderBy is validated the same way
286
+ // Site.all does — a clear error beats an opaque PostgREST 400 (unknown
287
+ // column) or a silently-wrong sort direction.
288
+ const hasOrderBy = hasText(orderBy?.attribute);
289
+ let orderField = 'updated_at';
290
+ let ascending = false;
291
+ if (hasOrderBy) {
292
+ const { toDbMap } = this.fieldMaps;
293
+ if (!Object.prototype.hasOwnProperty.call(toDbMap, orderBy.attribute)) {
294
+ throw new DataAccessError(`unknown orderBy attribute: ${orderBy.attribute}`, this);
295
+ }
296
+ const direction = orderBy.direction === undefined
297
+ ? 'asc'
298
+ : String(orderBy.direction).toLowerCase();
299
+ if (direction !== 'asc' && direction !== 'desc') {
300
+ throw new DataAccessError(`invalid orderBy direction: ${orderBy.direction}`, this);
301
+ }
302
+ orderField = toDbField(orderBy.attribute, toDbMap);
303
+ ascending = direction === 'asc';
304
+ }
305
+ query = query.order(orderField, { ascending });
306
+ const idField = this.fieldMaps.toDbMap[this.idName];
307
+ if (idField !== orderField) {
308
+ query = query.order(idField, { ascending });
309
+ }
310
+
311
+ // Honor the exact (positive) limit (no cap, no +1) so the api-service N+1
312
+ // hasMore detection stays symmetric with Site.all's postgrest path. A
313
+ // non-positive or non-integer limit falls back to DEFAULT_PAGE_SIZE so a
314
+ // caller-supplied 0/negative can't produce an inverted PostgREST range.
315
+ const effectiveLimit = Number.isInteger(limit) && limit > 0 ? limit : DEFAULT_PAGE_SIZE;
316
+ const offset = decodeCursor(cursor);
317
+ query = query.range(offset, offset + effectiveLimit - 1);
318
+
319
+ const { data, error } = await query;
320
+ if (error) {
321
+ this.log.error(`[SiteCollection] Failed to query sites by enrollment filter - ${error.message}`, error);
322
+ throw new DataAccessError('Failed to query sites by enrollment filter', this, error);
323
+ }
324
+
325
+ const instances = (data || []).map((row) => {
326
+ // Drop the embed PostgREST nests on each parent row so it cannot leak
327
+ // onto the hydrated Site record.
328
+ const siteRow = { ...row };
329
+ delete siteRow.site_enrollments;
330
+ return this.createInstanceFromRow(siteRow);
331
+ });
332
+
333
+ if (returnCursor) {
334
+ const nextCursor = instances.length === effectiveLimit
335
+ ? encodeCursor(offset + effectiveLimit)
336
+ : null;
337
+ return { data: instances, cursor: nextCursor };
338
+ }
339
+
340
+ return instances;
341
+ }
342
+
209
343
  async allByOrganizationIdAndProjectName(organizationId, projectName) {
210
344
  if (!hasText(organizationId)) {
211
345
  throw new DataAccessError('organizationId is required', this);