Skip to content

Commit 9dbd3da

Browse files
author
OpenRouter SDK Bot
committed
chore: update OpenAPI spec [sdk-bot]
1 parent 5a8274d commit 9dbd3da

1 file changed

Lines changed: 80 additions & 0 deletions

File tree

‎.speakeasy/in.openapi.yaml‎

Lines changed: 80 additions & 0 deletions
Original file line numberDiff line numberDiff line change
@@ -20303,6 +20303,7 @@ components:
2030320303
- 'top_p'
2030420304
- 'max_tokens'
2030520305
supports_implicit_caching: true
20306+
supports_voice_cloning: false
2030620307
tag: 'openai'
2030720308
throughput_last_30m:
2030820309
p50: 45.2
@@ -20404,6 +20405,10 @@ components:
2040420405
type: 'array'
2040520406
supports_implicit_caching:
2040620407
type: 'boolean'
20408+
supports_voice_cloning:
20409+
default: false
20410+
description: 'Whether this TTS endpoint accepts inline reference audio (`input_references`) for stateless voice cloning. Requests carrying reference audio are only routed to endpoints where this is true.'
20411+
type: 'boolean'
2040720412
tag:
2040820413
type: 'string'
2040920414
throughput_last_30m:
@@ -21836,6 +21841,68 @@ components:
2183621841
oneOf:
2183721842
- $ref: '#/components/schemas/ContainerAutoEnvironment'
2183821843
- $ref: '#/components/schemas/ContainerReferenceEnvironment'
21844+
SpeechInputReference:
21845+
description: 'Reference content part for stateless voice cloning'
21846+
discriminator:
21847+
mapping:
21848+
input_audio: '#/components/schemas/SpeechInputReferenceAudio'
21849+
text: '#/components/schemas/SpeechInputReferenceText'
21850+
propertyName: 'type'
21851+
oneOf:
21852+
- $ref: '#/components/schemas/SpeechInputReferenceAudio'
21853+
- $ref: '#/components/schemas/SpeechInputReferenceText'
21854+
SpeechInputReferenceAudio:
21855+
description: 'Reference audio input for stateless voice cloning'
21856+
example:
21857+
input_audio:
21858+
data: 'data:audio/wav;base64,UklGRuQXDABXQVZF...'
21859+
type: 'input_audio'
21860+
properties:
21861+
input_audio:
21862+
$ref: '#/components/schemas/SpeechInputReferenceAudioInput'
21863+
type:
21864+
enum:
21865+
- 'input_audio'
21866+
type: 'string'
21867+
required:
21868+
- 'type'
21869+
- 'input_audio'
21870+
type: 'object'
21871+
SpeechInputReferenceAudioInput:
21872+
description: 'Reference audio input object'
21873+
properties:
21874+
data:
21875+
description: 'Base64-encoded reference audio (optionally a data URI). Supported audio formats are provider-specific. Limited to 20 MiB of base64 (15 MiB of decoded audio).'
21876+
example: 'data:audio/wav;base64,UklGRuQXDABXQVZF...'
21877+
maxLength: 20971520
21878+
minLength: 1
21879+
type: 'string'
21880+
format:
21881+
description: 'Audio format of the reference audio (e.g., wav, mp3). Optional; most providers detect the format from the audio bytes.'
21882+
example: 'wav'
21883+
type: 'string'
21884+
required:
21885+
- 'data'
21886+
type: 'object'
21887+
SpeechInputReferenceText:
21888+
description: 'Transcript of the accompanying reference audio'
21889+
example:
21890+
text: 'I used to rule the world.'
21891+
type: 'text'
21892+
properties:
21893+
text:
21894+
description: 'Transcript of the accompanying reference audio.'
21895+
example: 'I used to rule the world.'
21896+
maxLength: 10000
21897+
type: 'string'
21898+
type:
21899+
enum:
21900+
- 'text'
21901+
type: 'string'
21902+
required:
21903+
- 'type'
21904+
- 'text'
21905+
type: 'object'
2183921906
SpeechRequest:
2184021907
description: 'Text-to-speech request input'
2184121908
example:
@@ -21849,6 +21916,17 @@ components:
2184921916
description: 'Text to synthesize'
2185021917
example: 'Hello world'
2185121918
type: 'string'
21919+
input_references:
21920+
description: 'Reference content for stateless voice cloning: one `input_audio` part carrying the voice sample, optionally accompanied by one `text` part with its transcript. Only routed to endpoints that support voice cloning.'
21921+
example:
21922+
- input_audio:
21923+
data: 'data:audio/wav;base64,UklGRuQXDABXQVZF...'
21924+
type: 'input_audio'
21925+
- text: 'I used to rule the world.'
21926+
type: 'text'
21927+
items:
21928+
$ref: '#/components/schemas/SpeechInputReference'
21929+
type: 'array'
2185221930
model:
2185321931
description: 'TTS model identifier'
2185421932
example: 'mistralai/voxtral-mini-tts-2603'
@@ -28169,6 +28247,7 @@ paths:
2816928247
- 'top_p'
2817028248
- 'max_tokens'
2817128249
supports_implicit_caching: true
28250+
supports_voice_cloning: false
2817228251
tag: 'openai'
2817328252
throughput_last_30m:
2817428253
p50: 45.2
@@ -28205,6 +28284,7 @@ paths:
2820528284
- 'top_p'
2820628285
- 'max_tokens'
2820728286
supports_implicit_caching: true
28287+
supports_voice_cloning: false
2820828288
tag: 'openai'
2820928289
throughput_last_30m:
2821028290
p50: 45.2

0 commit comments

Comments
 (0)