@ondewo/s2t-client-typescript 4.0.0 → 5.3.0
This diff represents the content of publicly available package versions that have been released to one of the supported registries. The information contained in this diff is provided for informational purposes only and reflects changes between package versions as they appear in their respective public registries.
- package/README.md +54 -13
- package/api/ondewo/s2t/speech-to-text_pb.d.ts +121 -31
- package/api/ondewo/s2t/speech-to-text_pb.js +808 -121
- package/package.json +1 -1
- package/public-api.d.ts +1 -1
- package/public-api.js +2 -2
package/README.md
CHANGED
|
@@ -1,20 +1,34 @@
|
|
|
1
|
-
<
|
|
2
|
-
<
|
|
3
|
-
<
|
|
4
|
-
|
|
1
|
+
<div align="center">
|
|
2
|
+
<table>
|
|
3
|
+
<tr>
|
|
4
|
+
<td>
|
|
5
|
+
<a href="https://ondewo.com/en/products/natural-language-understanding/">
|
|
6
|
+
<img width="400px" src="https://raw.githubusercontent.com/ondewo/ondewo-logos/master/ondewo_we_automate_your_phone_calls.png"/>
|
|
7
|
+
</a>
|
|
8
|
+
</td>
|
|
9
|
+
</tr>
|
|
10
|
+
<tr>
|
|
11
|
+
<td align="center">
|
|
12
|
+
<a href="https://www.linkedin.com/company/ondewo "><img width="40px" src="https://cdn-icons-png.flaticon.com/512/3536/3536505.png"></a>
|
|
13
|
+
<a href="https://www.facebook.com/ondewo"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/733/733547.png"></a>
|
|
14
|
+
<a href="https://twitter.com/ondewo"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/733/733579.png"> </a>
|
|
15
|
+
<a href="https://www.instagram.com/ondewo.ai/"><img width="40px" src="https://cdn-icons-png.flaticon.com/512/174/174855.png"></a>
|
|
16
|
+
<a href="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript"><img src="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript.svg" alt="npm version" height="32"></a>
|
|
17
|
+
</td>
|
|
18
|
+
</tr>
|
|
19
|
+
</table>
|
|
5
20
|
<h1 align="center">
|
|
6
21
|
ONDEWO S2T Client Typescript
|
|
7
22
|
</h1>
|
|
8
|
-
|
|
9
|
-
<a href="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript"><img src="https://badge.fury.io/js/%40ondewo%2Fs2t-client-typescript.svg" alt="npm version" height="18"></a>
|
|
10
|
-
</p>
|
|
11
|
-
</p>
|
|
23
|
+
</div>
|
|
12
24
|
|
|
13
|
-
|
|
25
|
+
## Overview
|
|
14
26
|
|
|
15
|
-
`@ondewo/s2t-client-typescript` is a compiled version of the [ONDEWO S2T API](https://github.com/ondewo/ondewo-s2t-api). Here you can find the S2T API [documentation](https://ondewo.github.io
|
|
27
|
+
`@ondewo/s2t-client-typescript` is a compiled version of the [ONDEWO S2T API](https://github.com/ondewo/ondewo-s2t-api) using the [ONDEWO PROTO COMPILER](https://github.com/ondewo/ondewo-proto-compiler). Here you can find the S2T API [documentation](https://ondewo.github.io).
|
|
16
28
|
|
|
17
|
-
|
|
29
|
+
ONDEWO APIs use [Protocol Buffers](https://github.com/google/protobuf) version 3 (proto3) as their Interface Definition Language (IDL) to define the API interface and the structure of the payload messages. The same interface definition is used for gRPC versions of the API in all languages.
|
|
30
|
+
|
|
31
|
+
## Setup
|
|
18
32
|
|
|
19
33
|
Using NPM:
|
|
20
34
|
|
|
@@ -22,6 +36,33 @@ Using NPM:
|
|
|
22
36
|
npm i --save @ondewo/s2t-client-typescript
|
|
23
37
|
```
|
|
24
38
|
|
|
25
|
-
|
|
39
|
+
Using GitHub:
|
|
40
|
+
|
|
41
|
+
```shell
|
|
42
|
+
git clone https://github.com/ondewo/ondewo-s2t-client-typescript.git ## Clone repository
|
|
43
|
+
cd ondewo-s2t-client-typescript ## Change into repo-directoy
|
|
44
|
+
make setup_developer_environment_locally ## Install dependencies
|
|
45
|
+
```
|
|
46
|
+
|
|
47
|
+
## Package structure
|
|
48
|
+
|
|
49
|
+
```
|
|
50
|
+
npm
|
|
51
|
+
├── api
|
|
52
|
+
│ ├── google
|
|
53
|
+
│ │ └── protobuf
|
|
54
|
+
│ │ ├── empty_pb.d.ts
|
|
55
|
+
│ │ └── empty_pb.js
|
|
56
|
+
│ └── ondewo
|
|
57
|
+
│ └── s2t
|
|
58
|
+
│ ├── speech-to-text_grpc_web_pb.d.ts
|
|
59
|
+
│ ├── speech-to-text_grpc_web_pb.js
|
|
60
|
+
│ ├── speech-to-text_pb.d.ts
|
|
61
|
+
│ └── speech-to-text_pb.js
|
|
62
|
+
├── LICENSE
|
|
63
|
+
├── package.json
|
|
64
|
+
├── public-api.d.ts
|
|
65
|
+
├── public-api.js
|
|
66
|
+
└── README.md
|
|
67
|
+
```
|
|
26
68
|
|
|
27
|
-
Read the full documentation on [GitHub](https://github.com/ondewo/ondewo-s2t-client-typescript).
|
|
@@ -7,8 +7,8 @@ export class TranscribeRequestConfig extends jspb.Message {
|
|
|
7
7
|
getS2tPipelineId(): string;
|
|
8
8
|
setS2tPipelineId(value: string): TranscribeRequestConfig;
|
|
9
9
|
|
|
10
|
-
|
|
11
|
-
|
|
10
|
+
getDecoding(): Decoding;
|
|
11
|
+
setDecoding(value: Decoding): TranscribeRequestConfig;
|
|
12
12
|
|
|
13
13
|
getLanguageModelName(): string;
|
|
14
14
|
setLanguageModelName(value: string): TranscribeRequestConfig;
|
|
@@ -59,7 +59,7 @@ export class TranscribeRequestConfig extends jspb.Message {
|
|
|
59
59
|
export namespace TranscribeRequestConfig {
|
|
60
60
|
export type AsObject = {
|
|
61
61
|
s2tPipelineId: string,
|
|
62
|
-
|
|
62
|
+
decoding: Decoding,
|
|
63
63
|
languageModelName: string,
|
|
64
64
|
postProcessing?: PostProcessingOptions.AsObject,
|
|
65
65
|
utteranceDetection?: UtteranceDetectionOptions.AsObject,
|
|
@@ -102,11 +102,20 @@ export class TranscriptionReturnOptions extends jspb.Message {
|
|
|
102
102
|
getReturnAudio(): boolean;
|
|
103
103
|
setReturnAudio(value: boolean): TranscriptionReturnOptions;
|
|
104
104
|
|
|
105
|
+
getReturnConfidenceScore(): boolean;
|
|
106
|
+
setReturnConfidenceScore(value: boolean): TranscriptionReturnOptions;
|
|
107
|
+
|
|
105
108
|
getReturnAlternativeTranscriptions(): boolean;
|
|
106
109
|
setReturnAlternativeTranscriptions(value: boolean): TranscriptionReturnOptions;
|
|
107
110
|
|
|
108
|
-
|
|
109
|
-
|
|
111
|
+
getReturnAlternativeTranscriptionsNr(): number;
|
|
112
|
+
setReturnAlternativeTranscriptionsNr(value: number): TranscriptionReturnOptions;
|
|
113
|
+
|
|
114
|
+
getReturnAlternativeWords(): boolean;
|
|
115
|
+
setReturnAlternativeWords(value: boolean): TranscriptionReturnOptions;
|
|
116
|
+
|
|
117
|
+
getReturnAlternativeWordsNr(): number;
|
|
118
|
+
setReturnAlternativeWordsNr(value: number): TranscriptionReturnOptions;
|
|
110
119
|
|
|
111
120
|
getReturnWordTiming(): boolean;
|
|
112
121
|
setReturnWordTiming(value: boolean): TranscriptionReturnOptions;
|
|
@@ -123,8 +132,11 @@ export namespace TranscriptionReturnOptions {
|
|
|
123
132
|
export type AsObject = {
|
|
124
133
|
returnStartOfSpeech: boolean,
|
|
125
134
|
returnAudio: boolean,
|
|
126
|
-
returnAlternativeTranscriptions: boolean,
|
|
127
135
|
returnConfidenceScore: boolean,
|
|
136
|
+
returnAlternativeTranscriptions: boolean,
|
|
137
|
+
returnAlternativeTranscriptionsNr: number,
|
|
138
|
+
returnAlternativeWords: boolean,
|
|
139
|
+
returnAlternativeWordsNr: number,
|
|
128
140
|
returnWordTiming: boolean,
|
|
129
141
|
}
|
|
130
142
|
}
|
|
@@ -668,16 +680,19 @@ export namespace S2TDescription {
|
|
|
668
680
|
}
|
|
669
681
|
|
|
670
682
|
export class S2TInference extends jspb.Message {
|
|
671
|
-
|
|
672
|
-
|
|
673
|
-
|
|
674
|
-
|
|
683
|
+
getAcousticModels(): AcousticModels | undefined;
|
|
684
|
+
setAcousticModels(value?: AcousticModels): S2TInference;
|
|
685
|
+
hasAcousticModels(): boolean;
|
|
686
|
+
clearAcousticModels(): S2TInference;
|
|
675
687
|
|
|
676
688
|
getLanguageModels(): LanguageModels | undefined;
|
|
677
689
|
setLanguageModels(value?: LanguageModels): S2TInference;
|
|
678
690
|
hasLanguageModels(): boolean;
|
|
679
691
|
clearLanguageModels(): S2TInference;
|
|
680
692
|
|
|
693
|
+
getInferenceBackend(): InferenceBackend;
|
|
694
|
+
setInferenceBackend(value: InferenceBackend): S2TInference;
|
|
695
|
+
|
|
681
696
|
serializeBinary(): Uint8Array;
|
|
682
697
|
toObject(includeInstance?: boolean): S2TInference.AsObject;
|
|
683
698
|
static toObject(includeInstance: boolean, msg: S2TInference): S2TInference.AsObject;
|
|
@@ -688,50 +703,119 @@ export class S2TInference extends jspb.Message {
|
|
|
688
703
|
|
|
689
704
|
export namespace S2TInference {
|
|
690
705
|
export type AsObject = {
|
|
691
|
-
|
|
706
|
+
acousticModels?: AcousticModels.AsObject,
|
|
692
707
|
languageModels?: LanguageModels.AsObject,
|
|
708
|
+
inferenceBackend: InferenceBackend,
|
|
693
709
|
}
|
|
694
710
|
}
|
|
695
711
|
|
|
696
|
-
export class
|
|
712
|
+
export class AcousticModels extends jspb.Message {
|
|
697
713
|
getType(): string;
|
|
698
|
-
setType(value: string):
|
|
714
|
+
setType(value: string): AcousticModels;
|
|
699
715
|
|
|
700
716
|
getQuartznet(): Quartznet | undefined;
|
|
701
|
-
setQuartznet(value?: Quartznet):
|
|
717
|
+
setQuartznet(value?: Quartznet): AcousticModels;
|
|
702
718
|
hasQuartznet(): boolean;
|
|
703
|
-
clearQuartznet():
|
|
719
|
+
clearQuartznet(): AcousticModels;
|
|
704
720
|
|
|
705
721
|
getQuartznetTriton(): QuartznetTriton | undefined;
|
|
706
|
-
setQuartznetTriton(value?: QuartznetTriton):
|
|
722
|
+
setQuartznetTriton(value?: QuartznetTriton): AcousticModels;
|
|
707
723
|
hasQuartznetTriton(): boolean;
|
|
708
|
-
clearQuartznetTriton():
|
|
724
|
+
clearQuartznetTriton(): AcousticModels;
|
|
709
725
|
|
|
710
726
|
getWav2vec(): Wav2Vec | undefined;
|
|
711
|
-
setWav2vec(value?: Wav2Vec):
|
|
727
|
+
setWav2vec(value?: Wav2Vec): AcousticModels;
|
|
712
728
|
hasWav2vec(): boolean;
|
|
713
|
-
clearWav2vec():
|
|
729
|
+
clearWav2vec(): AcousticModels;
|
|
714
730
|
|
|
715
731
|
getWav2vecTriton(): Wav2VecTriton | undefined;
|
|
716
|
-
setWav2vecTriton(value?: Wav2VecTriton):
|
|
732
|
+
setWav2vecTriton(value?: Wav2VecTriton): AcousticModels;
|
|
717
733
|
hasWav2vecTriton(): boolean;
|
|
718
|
-
clearWav2vecTriton():
|
|
734
|
+
clearWav2vecTriton(): AcousticModels;
|
|
735
|
+
|
|
736
|
+
getWhisper(): Whisper | undefined;
|
|
737
|
+
setWhisper(value?: Whisper): AcousticModels;
|
|
738
|
+
hasWhisper(): boolean;
|
|
739
|
+
clearWhisper(): AcousticModels;
|
|
740
|
+
|
|
741
|
+
getWhisperTriton(): WhisperTriton | undefined;
|
|
742
|
+
setWhisperTriton(value?: WhisperTriton): AcousticModels;
|
|
743
|
+
hasWhisperTriton(): boolean;
|
|
744
|
+
clearWhisperTriton(): AcousticModels;
|
|
719
745
|
|
|
720
746
|
serializeBinary(): Uint8Array;
|
|
721
|
-
toObject(includeInstance?: boolean):
|
|
722
|
-
static toObject(includeInstance: boolean, msg:
|
|
723
|
-
static serializeBinaryToWriter(message:
|
|
724
|
-
static deserializeBinary(bytes: Uint8Array):
|
|
725
|
-
static deserializeBinaryFromReader(message:
|
|
747
|
+
toObject(includeInstance?: boolean): AcousticModels.AsObject;
|
|
748
|
+
static toObject(includeInstance: boolean, msg: AcousticModels): AcousticModels.AsObject;
|
|
749
|
+
static serializeBinaryToWriter(message: AcousticModels, writer: jspb.BinaryWriter): void;
|
|
750
|
+
static deserializeBinary(bytes: Uint8Array): AcousticModels;
|
|
751
|
+
static deserializeBinaryFromReader(message: AcousticModels, reader: jspb.BinaryReader): AcousticModels;
|
|
726
752
|
}
|
|
727
753
|
|
|
728
|
-
export namespace
|
|
754
|
+
export namespace AcousticModels {
|
|
729
755
|
export type AsObject = {
|
|
730
756
|
type: string,
|
|
731
757
|
quartznet?: Quartznet.AsObject,
|
|
732
758
|
quartznetTriton?: QuartznetTriton.AsObject,
|
|
733
759
|
wav2vec?: Wav2Vec.AsObject,
|
|
734
760
|
wav2vecTriton?: Wav2VecTriton.AsObject,
|
|
761
|
+
whisper?: Whisper.AsObject,
|
|
762
|
+
whisperTriton?: WhisperTriton.AsObject,
|
|
763
|
+
}
|
|
764
|
+
}
|
|
765
|
+
|
|
766
|
+
export class Whisper extends jspb.Message {
|
|
767
|
+
getModelPath(): string;
|
|
768
|
+
setModelPath(value: string): Whisper;
|
|
769
|
+
|
|
770
|
+
getUseGpu(): boolean;
|
|
771
|
+
setUseGpu(value: boolean): Whisper;
|
|
772
|
+
|
|
773
|
+
getLanguage(): string;
|
|
774
|
+
setLanguage(value: string): Whisper;
|
|
775
|
+
|
|
776
|
+
serializeBinary(): Uint8Array;
|
|
777
|
+
toObject(includeInstance?: boolean): Whisper.AsObject;
|
|
778
|
+
static toObject(includeInstance: boolean, msg: Whisper): Whisper.AsObject;
|
|
779
|
+
static serializeBinaryToWriter(message: Whisper, writer: jspb.BinaryWriter): void;
|
|
780
|
+
static deserializeBinary(bytes: Uint8Array): Whisper;
|
|
781
|
+
static deserializeBinaryFromReader(message: Whisper, reader: jspb.BinaryReader): Whisper;
|
|
782
|
+
}
|
|
783
|
+
|
|
784
|
+
export namespace Whisper {
|
|
785
|
+
export type AsObject = {
|
|
786
|
+
modelPath: string,
|
|
787
|
+
useGpu: boolean,
|
|
788
|
+
language: string,
|
|
789
|
+
}
|
|
790
|
+
}
|
|
791
|
+
|
|
792
|
+
export class WhisperTriton extends jspb.Message {
|
|
793
|
+
getProcessorPath(): string;
|
|
794
|
+
setProcessorPath(value: string): WhisperTriton;
|
|
795
|
+
|
|
796
|
+
getTritonModelName(): string;
|
|
797
|
+
setTritonModelName(value: string): WhisperTriton;
|
|
798
|
+
|
|
799
|
+
getTritonModelVersion(): string;
|
|
800
|
+
setTritonModelVersion(value: string): WhisperTriton;
|
|
801
|
+
|
|
802
|
+
getCheckStatusTimeout(): number;
|
|
803
|
+
setCheckStatusTimeout(value: number): WhisperTriton;
|
|
804
|
+
|
|
805
|
+
serializeBinary(): Uint8Array;
|
|
806
|
+
toObject(includeInstance?: boolean): WhisperTriton.AsObject;
|
|
807
|
+
static toObject(includeInstance: boolean, msg: WhisperTriton): WhisperTriton.AsObject;
|
|
808
|
+
static serializeBinaryToWriter(message: WhisperTriton, writer: jspb.BinaryWriter): void;
|
|
809
|
+
static deserializeBinary(bytes: Uint8Array): WhisperTriton;
|
|
810
|
+
static deserializeBinaryFromReader(message: WhisperTriton, reader: jspb.BinaryReader): WhisperTriton;
|
|
811
|
+
}
|
|
812
|
+
|
|
813
|
+
export namespace WhisperTriton {
|
|
814
|
+
export type AsObject = {
|
|
815
|
+
processorPath: string,
|
|
816
|
+
tritonModelName: string,
|
|
817
|
+
tritonModelVersion: string,
|
|
818
|
+
checkStatusTimeout: number,
|
|
735
819
|
}
|
|
736
820
|
}
|
|
737
821
|
|
|
@@ -961,8 +1045,8 @@ export class StreamingSpeechRecognition extends jspb.Message {
|
|
|
961
1045
|
getTranscribeNotFinal(): boolean;
|
|
962
1046
|
setTranscribeNotFinal(value: boolean): StreamingSpeechRecognition;
|
|
963
1047
|
|
|
964
|
-
|
|
965
|
-
|
|
1048
|
+
getDecodingMethod(): string;
|
|
1049
|
+
setDecodingMethod(value: string): StreamingSpeechRecognition;
|
|
966
1050
|
|
|
967
1051
|
getSamplingRate(): number;
|
|
968
1052
|
setSamplingRate(value: number): StreamingSpeechRecognition;
|
|
@@ -990,7 +1074,7 @@ export class StreamingSpeechRecognition extends jspb.Message {
|
|
|
990
1074
|
export namespace StreamingSpeechRecognition {
|
|
991
1075
|
export type AsObject = {
|
|
992
1076
|
transcribeNotFinal: boolean,
|
|
993
|
-
|
|
1077
|
+
decodingMethod: string,
|
|
994
1078
|
samplingRate: number,
|
|
995
1079
|
minAudioChunkSize: number,
|
|
996
1080
|
startOfUtteranceThreshold: number,
|
|
@@ -1372,8 +1456,14 @@ export namespace TrainUserLanguageModelRequest {
|
|
|
1372
1456
|
}
|
|
1373
1457
|
}
|
|
1374
1458
|
|
|
1375
|
-
export enum
|
|
1459
|
+
export enum InferenceBackend {
|
|
1460
|
+
INFERENCE_BACKEND_UNKNOWN = 0,
|
|
1461
|
+
INFERENCE_BACKEND_PYTORCH = 1,
|
|
1462
|
+
INFERENCE_BACKEND_FLAX = 2,
|
|
1463
|
+
}
|
|
1464
|
+
export enum Decoding {
|
|
1376
1465
|
DEFAULT = 0,
|
|
1377
1466
|
GREEDY = 1,
|
|
1378
1467
|
BEAM_SEARCH_WITH_LM = 2,
|
|
1468
|
+
BEAM_SEARCH = 3,
|
|
1379
1469
|
}
|