mirror of
https://github.com/MindWorkAI/AI-Studio.git
synced 2026-09-13 20:33:38 +00:00
Some checks are pending
Build and Release / Determine run mode (push) Waiting to run
Build and Release / Sync Flatpak repo (push) Blocked by required conditions
Build and Release / Collect Flatpak artifacts (push) Blocked by required conditions
Build and Release / Verify (push) Waiting to run
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-aarch64-apple-darwin, osx-arm64, macos-latest, aarch64-apple-darwin, dmg,app,updater, dmg) (push) Blocked by required conditions
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-aarch64-pc-windows-msvc.exe, win-arm64, windows-latest, aarch64-pc-windows-msvc, nsis,updater, nsis) (push) Blocked by required conditions
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-aarch64-unknown-linux-gnu, linux-arm64, ubuntu-22.04-arm, aarch64-unknown-linux-gnu, appimage,updater, appimage) (push) Blocked by required conditions
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-x86_64-apple-darwin, osx-x64, macos-latest, x86_64-apple-darwin, dmg,app,updater, dmg) (push) Blocked by required conditions
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-x86_64-pc-windows-msvc.exe, win-x64, windows-latest, x86_64-pc-windows-msvc, nsis,updater, nsis) (push) Blocked by required conditions
Build and Release / Build app (${{ matrix.dotnet_runtime }}) (-x86_64-unknown-linux-gnu, linux-x64, ubuntu-22.04, x86_64-unknown-linux-gnu, appimage,updater, appimage) (push) Blocked by required conditions
Build and Release / Prepare & create release (push) Blocked by required conditions
Build and Release / Publish release (push) Blocked by required conditions
38 lines
1.6 KiB
C#
38 lines
1.6 KiB
C#
using AIStudio.Provider;
|
|
|
|
namespace AIStudio.Models.Kinds;
|
|
|
|
/// <summary>
|
|
/// The models which speak.
|
|
/// </summary>
|
|
/// <remarks>
|
|
/// Besides the pure text-to-speech models this covers the ones which answer in audio, such as
|
|
/// gpt-audio and gpt-4o-audio-preview. Those do accept a text-only request, but they are made for
|
|
/// spoken conversations, and the providers offering them keep them out of their chat model lists as
|
|
/// well.
|
|
///
|
|
/// All three words are stated as name parts. The markers they replace carried a hyphen on one side
|
|
/// to say the same thing, which caught one name these do not: Coqui's XTTS glues the word to an x.
|
|
/// It is named outright rather than loosening all three into substrings, where "tts" would be three
|
|
/// characters claiming every name that happens to contain them.
|
|
/// </remarks>
|
|
public sealed class SpeechSynthesisModelsFamily : ModelFamily
|
|
{
|
|
/// <inheritdoc />
|
|
public override ModelVendor Vendor => ModelVendor.UNKNOWN;
|
|
|
|
/// <inheritdoc />
|
|
public override ModelSource Source => new("https://huggingface.co/models?pipeline_tag=text-to-speech", new DateOnly(2026, 9, 12), "Ported from the speech synthesis markers of Provider/ModelKindExtensions.cs, where each of the three was written twice to allow for a separator on either side.");
|
|
|
|
/// <inheritdoc />
|
|
protected override void Declare(ModelFamilyBuilder builder)
|
|
{
|
|
builder.Modifier("tts").AsSegment().Kind(ModelKind.SPEECH_SYNTHESIS);
|
|
|
|
builder.Modifier("xtts").AsSegment().Inherits();
|
|
|
|
builder.Modifier("speech").AsSegment().Inherits();
|
|
|
|
builder.Modifier("audio").AsSegment().Inherits();
|
|
}
|
|
} |