@ondewo/s2t-client-typescript 4.0.0 → 5.3.0

This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
package/README.md CHANGED
@@ -1,20 +1,34 @@
1
- <p align="center">
2
- <a href="https://www.ondewo.com">
3
- <img alt="ONDEWO Logo" src="https://raw.githubusercontent.com/ondewo/ondewo-logos/master/github/ondewo_logo_github_2.png"/>
4
- </a>
1
+ <div align="center">
2
+ <table>
3
+ <tr>
4
+ <td>
5
+ <a href="https://ondewo.com/en/products/natural-language-understanding/">
6
+ <img width="400px" src="https://raw.githubusercontent.com/ondewo/ondewo-logos/master/ondewo_we_automate_your_phone_calls.png"/>
7
+ </a>
8
+ </td>
9
+ </tr>
10
+ <tr>
11
+ <td align="center">
12
+ <a href="https://www.linkedin.com/company/ondewo "><img width="40px" src="https://cdn-icons-png.flaticon.com/512/3536/3536505.png"></a>
13
+ <a href="https://www.facebook.com/ondewo"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/733/733547.png"></a>
14
+ <a href="https://twitter.com/ondewo"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/733/733579.png"> </a>
15
+ <a href="https://www.instagram.com/ondewo.ai/"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/174/174855.png"></a>
16
+ <a href="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript"><img src="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript.svg" alt="npm version" height="32"></a>
17
+ </td>
18
+ </tr>
19
+ </table>
5
20
  <h1 align="center">
6
21
  ONDEWO S2T Client Typescript
7
22
  </h1>
8
- <p align="center">
9
- <a href="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript"><img src="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript.svg" alt="npm version" height="18"></a>
10
- </p>
11
- </p>
23
+ </div>
12
24
 
13
- # @ondewo/s2t-client-typescript
25
+ ## Overview
14
26
 
15
- `@ondewo/s2t-client-typescript` is a compiled version of the [ONDEWO S2T API](https://github.com/ondewo/ondewo-s2t-api). Here you can find the S2T API [documentation](https://ondewo.github.io/ondewo-s2t-api/).
27
+ `@ondewo/s2t-client-typescript` is a compiled version of the [ONDEWO S2T API](https://github.com/ondewo/ondewo-s2t-api) using the [ONDEWO PROTO COMPILER](https://github.com/ondewo/ondewo-proto-compiler). Here you can find the S2T API [documentation](https://ondewo.github.io).
16
28
 
17
- ## Installation
29
+ ONDEWO APIs use [Protocol Buffers](https://github.com/google/protobuf) version 3 (proto3) as their Interface Definition Language (IDL) to define the API interface and the structure of the payload messages. The same interface definition is used for gRPC versions of the API in all languages.
30
+
31
+ ## Setup
18
32
 
19
33
  Using NPM:
20
34
 
@@ -22,6 +36,33 @@ Using NPM:
22
36
  npm i --save @ondewo/s2t-client-typescript
23
37
  ```
24
38
 
25
- ## Documentation
39
+ Using GitHub:
40
+
41
+ ```shell
42
+ git clone https://github.com/ondewo/ondewo-s2t-client-typescript.git ## Clone repository
43
+ cd ondewo-s2t-client-typescript ## Change into repo-directoy
44
+ make setup_developer_environment_locally ## Install dependencies
45
+ ```
46
+
47
+ ## Package structure
48
+
49
+ ```
50
+ npm
51
+ ├── api
52
+ │ ├── google
53
+ │ │ └── protobuf
54
+ │ │ ├── empty_pb.d.ts
55
+ │ │ └── empty_pb.js
56
+ │ └── ondewo
57
+ │ └── s2t
58
+ │ ├── speech-to-text_grpc_web_pb.d.ts
59
+ │ ├── speech-to-text_grpc_web_pb.js
60
+ │ ├── speech-to-text_pb.d.ts
61
+ │ └── speech-to-text_pb.js
62
+ ├── LICENSE
63
+ ├── package.json
64
+ ├── public-api.d.ts
65
+ ├── public-api.js
66
+ └── README.md
67
+ ```
26
68
 
27
- Read the full documentation on [GitHub](https://github.com/ondewo/ondewo-s2t-client-typescript).
@@ -7,8 +7,8 @@ export class TranscribeRequestConfig extends jspb.Message {
7
7
  getS2tPipelineId(): string;
8
8
  setS2tPipelineId(value: string): TranscribeRequestConfig;
9
9
 
10
- getCtcDecoding(): CTCDecoding;
11
- setCtcDecoding(value: CTCDecoding): TranscribeRequestConfig;
10
+ getDecoding(): Decoding;
11
+ setDecoding(value: Decoding): TranscribeRequestConfig;
12
12
 
13
13
  getLanguageModelName(): string;
14
14
  setLanguageModelName(value: string): TranscribeRequestConfig;
@@ -59,7 +59,7 @@ export class TranscribeRequestConfig extends jspb.Message {
59
59
  export namespace TranscribeRequestConfig {
60
60
  export type AsObject = {
61
61
  s2tPipelineId: string,
62
- ctcDecoding: CTCDecoding,
62
+ decoding: Decoding,
63
63
  languageModelName: string,
64
64
  postProcessing?: PostProcessingOptions.AsObject,
65
65
  utteranceDetection?: UtteranceDetectionOptions.AsObject,
@@ -102,11 +102,20 @@ export class TranscriptionReturnOptions extends jspb.Message {
102
102
  getReturnAudio(): boolean;
103
103
  setReturnAudio(value: boolean): TranscriptionReturnOptions;
104
104
 
105
+ getReturnConfidenceScore(): boolean;
106
+ setReturnConfidenceScore(value: boolean): TranscriptionReturnOptions;
107
+
105
108
  getReturnAlternativeTranscriptions(): boolean;
106
109
  setReturnAlternativeTranscriptions(value: boolean): TranscriptionReturnOptions;
107
110
 
108
- getReturnConfidenceScore(): boolean;
109
- setReturnConfidenceScore(value: boolean): TranscriptionReturnOptions;
111
+ getReturnAlternativeTranscriptionsNr(): number;
112
+ setReturnAlternativeTranscriptionsNr(value: number): TranscriptionReturnOptions;
113
+
114
+ getReturnAlternativeWords(): boolean;
115
+ setReturnAlternativeWords(value: boolean): TranscriptionReturnOptions;
116
+
117
+ getReturnAlternativeWordsNr(): number;
118
+ setReturnAlternativeWordsNr(value: number): TranscriptionReturnOptions;
110
119
 
111
120
  getReturnWordTiming(): boolean;
112
121
  setReturnWordTiming(value: boolean): TranscriptionReturnOptions;
@@ -123,8 +132,11 @@ export namespace TranscriptionReturnOptions {
123
132
  export type AsObject = {
124
133
  returnStartOfSpeech: boolean,
125
134
  returnAudio: boolean,
126
- returnAlternativeTranscriptions: boolean,
127
135
  returnConfidenceScore: boolean,
136
+ returnAlternativeTranscriptions: boolean,
137
+ returnAlternativeTranscriptionsNr: number,
138
+ returnAlternativeWords: boolean,
139
+ returnAlternativeWordsNr: number,
128
140
  returnWordTiming: boolean,
129
141
  }
130
142
  }
@@ -668,16 +680,19 @@ export namespace S2TDescription {
668
680
  }
669
681
 
670
682
  export class S2TInference extends jspb.Message {
671
- getCtcAcousticModels(): CtcAcousticModels | undefined;
672
- setCtcAcousticModels(value?: CtcAcousticModels): S2TInference;
673
- hasCtcAcousticModels(): boolean;
674
- clearCtcAcousticModels(): S2TInference;
683
+ getAcousticModels(): AcousticModels | undefined;
684
+ setAcousticModels(value?: AcousticModels): S2TInference;
685
+ hasAcousticModels(): boolean;
686
+ clearAcousticModels(): S2TInference;
675
687
 
676
688
  getLanguageModels(): LanguageModels | undefined;
677
689
  setLanguageModels(value?: LanguageModels): S2TInference;
678
690
  hasLanguageModels(): boolean;
679
691
  clearLanguageModels(): S2TInference;
680
692
 
693
+ getInferenceBackend(): InferenceBackend;
694
+ setInferenceBackend(value: InferenceBackend): S2TInference;
695
+
681
696
  serializeBinary(): Uint8Array;
682
697
  toObject(includeInstance?: boolean): S2TInference.AsObject;
683
698
  static toObject(includeInstance: boolean, msg: S2TInference): S2TInference.AsObject;
@@ -688,50 +703,119 @@ export class S2TInference extends jspb.Message {
688
703
 
689
704
  export namespace S2TInference {
690
705
  export type AsObject = {
691
- ctcAcousticModels?: CtcAcousticModels.AsObject,
706
+ acousticModels?: AcousticModels.AsObject,
692
707
  languageModels?: LanguageModels.AsObject,
708
+ inferenceBackend: InferenceBackend,
693
709
  }
694
710
  }
695
711
 
696
- export class CtcAcousticModels extends jspb.Message {
712
+ export class AcousticModels extends jspb.Message {
697
713
  getType(): string;
698
- setType(value: string): CtcAcousticModels;
714
+ setType(value: string): AcousticModels;
699
715
 
700
716
  getQuartznet(): Quartznet | undefined;
701
- setQuartznet(value?: Quartznet): CtcAcousticModels;
717
+ setQuartznet(value?: Quartznet): AcousticModels;
702
718
  hasQuartznet(): boolean;
703
- clearQuartznet(): CtcAcousticModels;
719
+ clearQuartznet(): AcousticModels;
704
720
 
705
721
  getQuartznetTriton(): QuartznetTriton | undefined;
706
- setQuartznetTriton(value?: QuartznetTriton): CtcAcousticModels;
722
+ setQuartznetTriton(value?: QuartznetTriton): AcousticModels;
707
723
  hasQuartznetTriton(): boolean;
708
- clearQuartznetTriton(): CtcAcousticModels;
724
+ clearQuartznetTriton(): AcousticModels;
709
725
 
710
726
  getWav2vec(): Wav2Vec | undefined;
711
- setWav2vec(value?: Wav2Vec): CtcAcousticModels;
727
+ setWav2vec(value?: Wav2Vec): AcousticModels;
712
728
  hasWav2vec(): boolean;
713
- clearWav2vec(): CtcAcousticModels;
729
+ clearWav2vec(): AcousticModels;
714
730
 
715
731
  getWav2vecTriton(): Wav2VecTriton | undefined;
716
- setWav2vecTriton(value?: Wav2VecTriton): CtcAcousticModels;
732
+ setWav2vecTriton(value?: Wav2VecTriton): AcousticModels;
717
733
  hasWav2vecTriton(): boolean;
718
- clearWav2vecTriton(): CtcAcousticModels;
734
+ clearWav2vecTriton(): AcousticModels;
735
+
736
+ getWhisper(): Whisper | undefined;
737
+ setWhisper(value?: Whisper): AcousticModels;
738
+ hasWhisper(): boolean;
739
+ clearWhisper(): AcousticModels;
740
+
741
+ getWhisperTriton(): WhisperTriton | undefined;
742
+ setWhisperTriton(value?: WhisperTriton): AcousticModels;
743
+ hasWhisperTriton(): boolean;
744
+ clearWhisperTriton(): AcousticModels;
719
745
 
720
746
  serializeBinary(): Uint8Array;
721
- toObject(includeInstance?: boolean): CtcAcousticModels.AsObject;
722
- static toObject(includeInstance: boolean, msg: CtcAcousticModels): CtcAcousticModels.AsObject;
723
- static serializeBinaryToWriter(message: CtcAcousticModels, writer: jspb.BinaryWriter): void;
724
- static deserializeBinary(bytes: Uint8Array): CtcAcousticModels;
725
- static deserializeBinaryFromReader(message: CtcAcousticModels, reader: jspb.BinaryReader): CtcAcousticModels;
747
+ toObject(includeInstance?: boolean): AcousticModels.AsObject;
748
+ static toObject(includeInstance: boolean, msg: AcousticModels): AcousticModels.AsObject;
749
+ static serializeBinaryToWriter(message: AcousticModels, writer: jspb.BinaryWriter): void;
750
+ static deserializeBinary(bytes: Uint8Array): AcousticModels;
751
+ static deserializeBinaryFromReader(message: AcousticModels, reader: jspb.BinaryReader): AcousticModels;
726
752
  }
727
753
 
728
- export namespace CtcAcousticModels {
754
+ export namespace AcousticModels {
729
755
  export type AsObject = {
730
756
  type: string,
731
757
  quartznet?: Quartznet.AsObject,
732
758
  quartznetTriton?: QuartznetTriton.AsObject,
733
759
  wav2vec?: Wav2Vec.AsObject,
734
760
  wav2vecTriton?: Wav2VecTriton.AsObject,
761
+ whisper?: Whisper.AsObject,
762
+ whisperTriton?: WhisperTriton.AsObject,
763
+ }
764
+ }
765
+
766
+ export class Whisper extends jspb.Message {
767
+ getModelPath(): string;
768
+ setModelPath(value: string): Whisper;
769
+
770
+ getUseGpu(): boolean;
771
+ setUseGpu(value: boolean): Whisper;
772
+
773
+ getLanguage(): string;
774
+ setLanguage(value: string): Whisper;
775
+
776
+ serializeBinary(): Uint8Array;
777
+ toObject(includeInstance?: boolean): Whisper.AsObject;
778
+ static toObject(includeInstance: boolean, msg: Whisper): Whisper.AsObject;
779
+ static serializeBinaryToWriter(message: Whisper, writer: jspb.BinaryWriter): void;
780
+ static deserializeBinary(bytes: Uint8Array): Whisper;
781
+ static deserializeBinaryFromReader(message: Whisper, reader: jspb.BinaryReader): Whisper;
782
+ }
783
+
784
+ export namespace Whisper {
785
+ export type AsObject = {
786
+ modelPath: string,
787
+ useGpu: boolean,
788
+ language: string,
789
+ }
790
+ }
791
+
792
+ export class WhisperTriton extends jspb.Message {
793
+ getProcessorPath(): string;
794
+ setProcessorPath(value: string): WhisperTriton;
795
+
796
+ getTritonModelName(): string;
797
+ setTritonModelName(value: string): WhisperTriton;
798
+
799
+ getTritonModelVersion(): string;
800
+ setTritonModelVersion(value: string): WhisperTriton;
801
+
802
+ getCheckStatusTimeout(): number;
803
+ setCheckStatusTimeout(value: number): WhisperTriton;
804
+
805
+ serializeBinary(): Uint8Array;
806
+ toObject(includeInstance?: boolean): WhisperTriton.AsObject;
807
+ static toObject(includeInstance: boolean, msg: WhisperTriton): WhisperTriton.AsObject;
808
+ static serializeBinaryToWriter(message: WhisperTriton, writer: jspb.BinaryWriter): void;
809
+ static deserializeBinary(bytes: Uint8Array): WhisperTriton;
810
+ static deserializeBinaryFromReader(message: WhisperTriton, reader: jspb.BinaryReader): WhisperTriton;
811
+ }
812
+
813
+ export namespace WhisperTriton {
814
+ export type AsObject = {
815
+ processorPath: string,
816
+ tritonModelName: string,
817
+ tritonModelVersion: string,
818
+ checkStatusTimeout: number,
735
819
  }
736
820
  }
737
821
 
@@ -961,8 +1045,8 @@ export class StreamingSpeechRecognition extends jspb.Message {
961
1045
  getTranscribeNotFinal(): boolean;
962
1046
  setTranscribeNotFinal(value: boolean): StreamingSpeechRecognition;
963
1047
 
964
- getCtcDecodingMethod(): string;
965
- setCtcDecodingMethod(value: string): StreamingSpeechRecognition;
1048
+ getDecodingMethod(): string;
1049
+ setDecodingMethod(value: string): StreamingSpeechRecognition;
966
1050
 
967
1051
  getSamplingRate(): number;
968
1052
  setSamplingRate(value: number): StreamingSpeechRecognition;
@@ -990,7 +1074,7 @@ export class StreamingSpeechRecognition extends jspb.Message {
990
1074
  export namespace StreamingSpeechRecognition {
991
1075
  export type AsObject = {
992
1076
  transcribeNotFinal: boolean,
993
- ctcDecodingMethod: string,
1077
+ decodingMethod: string,
994
1078
  samplingRate: number,
995
1079
  minAudioChunkSize: number,
996
1080
  startOfUtteranceThreshold: number,
@@ -1372,8 +1456,14 @@ export namespace TrainUserLanguageModelRequest {
1372
1456
  }
1373
1457
  }
1374
1458
 
1375
- export enum CTCDecoding {
1459
+ export enum InferenceBackend {
1460
+ INFERENCE_BACKEND_UNKNOWN = 0,
1461
+ INFERENCE_BACKEND_PYTORCH = 1,
1462
+ INFERENCE_BACKEND_FLAX = 2,
1463
+ }
1464
+ export enum Decoding {
1376
1465
  DEFAULT = 0,
1377
1466
  GREEDY = 1,
1378
1467
  BEAM_SEARCH_WITH_LM = 2,
1468
+ BEAM_SEARCH = 3,
1379
1469
  }