@depup/ai-sdk__google 4.0.50-depup.0 → 4.0.54-depup.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/CHANGELOG.md CHANGED
@@ -1,5 +1,41 @@
1
1
  # @ai-sdk/google
2
2
 
3
+ ## 4.0.54
4
+
5
+ ### Patch Changes
6
+
7
+ - 1f7835c: feat (provider/google, provider/google-vertex): Gemini 3.5 Transcribe support — unary transcription (`gemini-3.5-transcribe`) via generateContent with language detection, speaker diarization, word timestamps, and custom vocabulary, plus streaming transcription (`gemini-3.5-transcribe-live`) over the Live API WebSocket with `mode: 'VERBATIM' | 'SMART'` transcription formatting
8
+
9
+ ## 4.0.53
10
+
11
+ ### Patch Changes
12
+
13
+ - Updated dependencies [3e125ba]
14
+ - @ai-sdk/provider-utils@5.0.32
15
+
16
+ ## 4.0.52
17
+
18
+ ### Patch Changes
19
+
20
+ - 7de3612: Encode provider-returned identifiers before using them in credentialed follow-up request paths.
21
+ - a9782e1: fix: align batch result parsing, request counts, and lifecycle behavior across providers
22
+ - 92e08e6: Preserve recursive tool input schemas without aborting Google model calls.
23
+ - 35841f5: feat: normalize mid-stream provider error events across supported providers into public StreamProviderError instances and preserve provider-owned type, code, status, retry, and raw payload metadata
24
+ - 0246209: Fix Google file uploads failing to type-check with TypeScript 5.9 DOM types.
25
+ - Updated dependencies [a9782e1]
26
+ - Updated dependencies [35841f5]
27
+ - Updated dependencies [d2f3353]
28
+ - @ai-sdk/provider-utils@5.0.31
29
+
30
+ ## 4.0.51
31
+
32
+ ### Patch Changes
33
+
34
+ - e7fc90e: feat(google): support the Gemini Batch API with experimental_startTextBatch
35
+ - Updated dependencies [591d25b]
36
+ - @ai-sdk/provider@4.0.8
37
+ - @ai-sdk/provider-utils@5.0.30
38
+
3
39
  ## 4.0.50
4
40
 
5
41
  ### Patch Changes
package/README.md CHANGED
@@ -13,8 +13,8 @@ npm install @depup/ai-sdk__google
13
13
 
14
14
  | Field | Value |
15
15
  |-------|-------|
16
- | Original | [@ai-sdk/google](https://www.npmjs.com/package/@ai-sdk/google) @ 4.0.50 |
17
- | Processed | 2026-08-22 |
16
+ | Original | [@ai-sdk/google](https://www.npmjs.com/package/@ai-sdk/google) @ 4.0.54 |
17
+ | Processed | 2026-08-27 |
18
18
  | Smoke test | failed |
19
19
  | Deps updated | 0 |
20
20
 
package/changes.json CHANGED
@@ -1,5 +1,5 @@
1
1
  {
2
2
  "bumped": {},
3
- "timestamp": "2026-08-22T08:07:16.694Z",
3
+ "timestamp": "2026-08-27T01:36:36.918Z",
4
4
  "totalUpdated": 0
5
5
  }
package/dist/index.d.ts CHANGED
@@ -1,8 +1,8 @@
1
1
  import * as _ai_sdk_provider_utils from '@ai-sdk/provider-utils';
2
- import { InferSchema, FetchFunction, WebSocketConstructor, WORKFLOW_SERIALIZE, WORKFLOW_DESERIALIZE } from '@ai-sdk/provider-utils';
2
+ import { InferSchema, FetchFunction, WebSocketConstructor, WORKFLOW_SERIALIZE, WORKFLOW_DESERIALIZE, Resolvable } from '@ai-sdk/provider-utils';
3
3
  import { z } from 'zod/v4';
4
4
  import * as _ai_sdk_provider from '@ai-sdk/provider';
5
- import { ProviderV4, LanguageModelV4, ImageModelV4, EmbeddingModelV4, Experimental_VideoModelV4, Experimental_SpeechTranslationModelV4, SpeechModelV4, FilesV4, Experimental_RealtimeFactoryV4, Experimental_RealtimeModelV4, Experimental_RealtimeModelV4ClientSecretOptions, Experimental_RealtimeModelV4ClientSecretResult, Experimental_RealtimeModelV4ServerEvent, Experimental_RealtimeModelV4ClientEvent, Experimental_RealtimeModelV4SessionConfig, Experimental_SpeechTranslationModelV4StreamOptions } from '@ai-sdk/provider';
5
+ import { ProviderV4, Experimental_BatchLanguageModelV4, ImageModelV4, EmbeddingModelV4, Experimental_VideoModelV4, Experimental_SpeechTranslationModelV4, SpeechModelV4, TranscriptionModelV4, FilesV4, LanguageModelV4, Experimental_RealtimeFactoryV4, Experimental_RealtimeModelV4, Experimental_RealtimeModelV4ClientSecretOptions, Experimental_RealtimeModelV4ClientSecretResult, Experimental_RealtimeModelV4ServerEvent, Experimental_RealtimeModelV4ClientEvent, Experimental_RealtimeModelV4SessionConfig, JSONObject, Experimental_TranscriptionModelV4StreamOptions, Experimental_SpeechTranslationModelV4StreamOptions } from '@ai-sdk/provider';
6
6
 
7
7
  declare const googleErrorDataSchema: _ai_sdk_provider_utils.LazySchema<{
8
8
  error: {
@@ -512,6 +512,24 @@ interface GoogleImageSettings {
512
512
  maxImagesPerCall?: number;
513
513
  }
514
514
 
515
+ type GoogleTranscriptionModelId = 'gemini-3.5-transcribe' | 'gemini-3.5-transcribe-live' | (string & {});
516
+ /**
517
+ * Speech recognition options shared by unary (`gemini-3.5-transcribe`) and
518
+ * live (`gemini-3.5-transcribe-live`) transcription. Maps onto Google's
519
+ * `AudioTranscriptionConfig`.
520
+ */
521
+ declare const googleTranscriptionModelOptions: z.ZodObject<{
522
+ languageCodes: z.ZodOptional<z.ZodArray<z.ZodString>>;
523
+ customVocabulary: z.ZodOptional<z.ZodArray<z.ZodString>>;
524
+ wordTimestamp: z.ZodOptional<z.ZodBoolean>;
525
+ diarization: z.ZodOptional<z.ZodBoolean>;
526
+ mode: z.ZodOptional<z.ZodEnum<{
527
+ SMART: "SMART";
528
+ VERBATIM: "VERBATIM";
529
+ }>>;
530
+ }, z.core.$strip>;
531
+ type GoogleTranscriptionModelOptions = z.infer<typeof googleTranscriptionModelOptions>;
532
+
515
533
  type GoogleSpeechTranslationModelId = 'gemini-3.5-live-translate-preview' | (string & {});
516
534
  declare const googleSpeechTranslationModelOptions: _ai_sdk_provider_utils.LazySchema<{
517
535
  echoTargetLanguage?: boolean | undefined;
@@ -519,9 +537,9 @@ declare const googleSpeechTranslationModelOptions: _ai_sdk_provider_utils.LazySc
519
537
  type GoogleSpeechTranslationModelOptions = InferSchema<typeof googleSpeechTranslationModelOptions>;
520
538
 
521
539
  interface GoogleProvider extends ProviderV4 {
522
- (modelId: GoogleModelId): LanguageModelV4;
523
- languageModel(modelId: GoogleModelId): LanguageModelV4;
524
- chat(modelId: GoogleModelId): LanguageModelV4;
540
+ (modelId: GoogleModelId): Experimental_BatchLanguageModelV4;
541
+ languageModel(modelId: GoogleModelId): Experimental_BatchLanguageModelV4;
542
+ chat(modelId: GoogleModelId): Experimental_BatchLanguageModelV4;
525
543
  /**
526
544
  * Creates a model for image generation.
527
545
  */
@@ -529,7 +547,7 @@ interface GoogleProvider extends ProviderV4 {
529
547
  /**
530
548
  * @deprecated Use `chat()` instead.
531
549
  */
532
- generativeAI(modelId: GoogleModelId): LanguageModelV4;
550
+ generativeAI(modelId: GoogleModelId): Experimental_BatchLanguageModelV4;
533
551
  /**
534
552
  * Creates a model for text embeddings.
535
553
  */
@@ -570,6 +588,17 @@ interface GoogleProvider extends ProviderV4 {
570
588
  * Creates a model for speech generation (text-to-speech).
571
589
  */
572
590
  speechModel(modelId: GoogleSpeechModelId): SpeechModelV4;
591
+ /**
592
+ * Creates a model for transcription (speech-to-text). Unary models
593
+ * (e.g. `gemini-3.5-transcribe`) transcribe audio files; live models
594
+ * (e.g. `gemini-3.5-transcribe-live`) stream transcription over the
595
+ * Gemini Live API WebSocket via `experimental_streamTranscribe`.
596
+ */
597
+ transcription(modelId: GoogleTranscriptionModelId): TranscriptionModelV4;
598
+ /**
599
+ * Creates a model for transcription (speech-to-text).
600
+ */
601
+ transcriptionModel(modelId: GoogleTranscriptionModelId): TranscriptionModelV4;
573
602
  files(): FilesV4;
574
603
  /**
575
604
  * Creates a language model targeting the Gemini Interactions API
@@ -679,6 +708,36 @@ type GoogleRealtimeModelOptions = {
679
708
  };
680
709
  };
681
710
 
711
+ interface GoogleTranscriptionModelConfig {
712
+ provider: string;
713
+ baseURL: string;
714
+ headers?: Resolvable<Record<string, string | undefined>>;
715
+ fetch?: FetchFunction;
716
+ webSocket?: WebSocketConstructor;
717
+ _internal?: {
718
+ currentDate?: () => Date;
719
+ finishGraceMs?: number;
720
+ };
721
+ }
722
+ declare class GoogleTranscriptionModel implements TranscriptionModelV4 {
723
+ readonly modelId: GoogleTranscriptionModelId;
724
+ private readonly config;
725
+ readonly specificationVersion = "v4";
726
+ static [WORKFLOW_SERIALIZE](model: GoogleTranscriptionModel): {
727
+ modelId: string;
728
+ config: JSONObject;
729
+ };
730
+ static [WORKFLOW_DESERIALIZE](options: {
731
+ modelId: GoogleTranscriptionModelId;
732
+ config: GoogleTranscriptionModelConfig;
733
+ }): GoogleTranscriptionModel;
734
+ get provider(): string;
735
+ constructor(modelId: GoogleTranscriptionModelId, config: GoogleTranscriptionModelConfig);
736
+ private parseOptions;
737
+ doGenerate(options: Parameters<TranscriptionModelV4['doGenerate']>[0]): Promise<Awaited<ReturnType<TranscriptionModelV4['doGenerate']>>>;
738
+ doStream(options: Experimental_TranscriptionModelV4StreamOptions): Promise<Awaited<ReturnType<NonNullable<TranscriptionModelV4['doStream']>>>>;
739
+ }
740
+
682
741
  type GoogleSpeechTranslationModelConfig = {
683
742
  provider: string;
684
743
  baseURL: string;
@@ -708,4 +767,4 @@ declare class GoogleSpeechTranslationModel implements Experimental_SpeechTransla
708
767
 
709
768
  declare const VERSION: string;
710
769
 
711
- export { GoogleRealtimeModel as Experimental_GoogleRealtimeModel, type GoogleRealtimeModelConfig as Experimental_GoogleRealtimeModelConfig, type GoogleRealtimeModelId as Experimental_GoogleRealtimeModelId, type GoogleRealtimeModelOptions as Experimental_GoogleRealtimeModelOptions, GoogleSpeechTranslationModel as Experimental_GoogleSpeechTranslationModel, type GoogleSpeechTranslationModelConfig as Experimental_GoogleSpeechTranslationModelConfig, type GoogleSpeechTranslationModelId as Experimental_GoogleSpeechTranslationModelId, type GoogleSpeechTranslationModelOptions as Experimental_GoogleSpeechTranslationModelOptions, GoogleSpeechTranslationModel as Experimental_GoogleTranslationModel, type GoogleSpeechTranslationModelConfig as Experimental_GoogleTranslationModelConfig, type GoogleSpeechTranslationModelId as Experimental_GoogleTranslationModelId, type GoogleSpeechTranslationModelOptions as Experimental_GoogleTranslationModelOptions, type GoogleEmbeddingModelOptions, type GoogleErrorData, type GoogleFilesUploadOptions, type GoogleEmbeddingModelOptions as GoogleGenerativeAIEmbeddingProviderOptions, type GoogleImageModelOptions as GoogleGenerativeAIImageProviderOptions, type GoogleProvider as GoogleGenerativeAIProvider, type GoogleProviderMetadata as GoogleGenerativeAIProviderMetadata, type GoogleLanguageModelOptions as GoogleGenerativeAIProviderOptions, type GoogleProviderSettings as GoogleGenerativeAIProviderSettings, type GoogleVideoModelId as GoogleGenerativeAIVideoModelId, type GoogleVideoModelOptions as GoogleGenerativeAIVideoProviderOptions, type GoogleImageModelOptions, type GoogleInteractionsAgentName, type GoogleInteractionsModelId, type GoogleInteractionsProviderMetadata, type GoogleLanguageModelInteractionsOptions, type GoogleLanguageModelOptions, type GoogleProvider, type GoogleProviderMetadata, type GoogleProviderSettings, type GoogleSpeechModelId, type GoogleSpeechModelOptions, type GoogleVideoModelId, type GoogleVideoModelOptions, VERSION, createGoogle, createGoogle as createGoogleGenerativeAI, google };
770
+ export { GoogleRealtimeModel as Experimental_GoogleRealtimeModel, type GoogleRealtimeModelConfig as Experimental_GoogleRealtimeModelConfig, type GoogleRealtimeModelId as Experimental_GoogleRealtimeModelId, type GoogleRealtimeModelOptions as Experimental_GoogleRealtimeModelOptions, GoogleSpeechTranslationModel as Experimental_GoogleSpeechTranslationModel, type GoogleSpeechTranslationModelConfig as Experimental_GoogleSpeechTranslationModelConfig, type GoogleSpeechTranslationModelId as Experimental_GoogleSpeechTranslationModelId, type GoogleSpeechTranslationModelOptions as Experimental_GoogleSpeechTranslationModelOptions, GoogleSpeechTranslationModel as Experimental_GoogleTranslationModel, type GoogleSpeechTranslationModelConfig as Experimental_GoogleTranslationModelConfig, type GoogleSpeechTranslationModelId as Experimental_GoogleTranslationModelId, type GoogleSpeechTranslationModelOptions as Experimental_GoogleTranslationModelOptions, type GoogleEmbeddingModelOptions, type GoogleErrorData, type GoogleFilesUploadOptions, type GoogleEmbeddingModelOptions as GoogleGenerativeAIEmbeddingProviderOptions, type GoogleImageModelOptions as GoogleGenerativeAIImageProviderOptions, type GoogleProvider as GoogleGenerativeAIProvider, type GoogleProviderMetadata as GoogleGenerativeAIProviderMetadata, type GoogleLanguageModelOptions as GoogleGenerativeAIProviderOptions, type GoogleProviderSettings as GoogleGenerativeAIProviderSettings, type GoogleVideoModelId as GoogleGenerativeAIVideoModelId, type GoogleVideoModelOptions as GoogleGenerativeAIVideoProviderOptions, type GoogleImageModelOptions, type GoogleInteractionsAgentName, type GoogleInteractionsModelId, type GoogleInteractionsProviderMetadata, type GoogleLanguageModelInteractionsOptions, type GoogleLanguageModelOptions, type GoogleProvider, type GoogleProviderMetadata, type GoogleProviderSettings, type GoogleSpeechModelId, type GoogleSpeechModelOptions, GoogleTranscriptionModel, type GoogleTranscriptionModelId, type GoogleTranscriptionModelOptions, type GoogleVideoModelId, type GoogleVideoModelOptions, VERSION, createGoogle, createGoogle as createGoogleGenerativeAI, google };