| # | uttid | text | ref |
baseline
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
basket_generation_config:
basket_lang: en-us
basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
batch_size: 1
gpus: 1
inference:
condition_sample_rate: 24000
diff_k: 3
diff_steps: 100
diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
duplicate_reference: true
exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
gpt_generate_args:
do_sample: true
min_new_tokens: 20
num_return_sequences: 50
use_cache: true
out_sample_rate: 24000
override_conditioning_features:
bad_text_proba: 0.0
c50: 0.0
dmcs_flatness: 100500.0
dmcs_roll_off_0.995: 100500.0
emo2vec: null
emotion: null
is_animation: 1.0
pitch_std: 100.0
snr: 100.0
year: 2025.0
reranking_options:
mode: MBR
top_k: 1
vocoder: bigvgan
voice_samples_preprocessing: []
num_workers: 1
output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15__2026-01-23_18-31-03
ref_dir: cc_20250825/ref
ticket: QUALITY-41
basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_18-31-03
|
diff_baseline |
diff_ft
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_ft
lang: en-us
meta:
basket_generation_config:
basket_lang: en-us
basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
batch_size: 1
gpus: 1
inference:
condition_sample_rate: 24000
diff_k: 3
diff_steps: 100
diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_ft
duplicate_reference: true
exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
gpt_generate_args:
do_sample: true
min_new_tokens: 20
num_return_sequences: 50
use_cache: true
out_sample_rate: 24000
override_conditioning_features:
bad_text_proba: 0.0
c50: 100.0
dmcs_flatness: 100500.0
dmcs_roll_off_0.995: 100500.0
emo2vec: null
emotion: null
is_animation: 1.0
pitch_std: 100.0
snr: 100.0
year: 2025.0
reranking_options:
mode: MBR
top_k: 1
vocoder: bigvgan
voice_samples_preprocessing: []
num_workers: 1
output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_ft__2026-03-13_11-52-34
ref_dir: cc_20250825/ref
ticket: QUALITY-41
basket_generation_git_hash: a0fff333bc707fa86bab206dec923bbfe5b51d8c
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-03-13_11-52-34
|
diff_clean
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_clean
lang: en-us
meta:
basket_generation_config:
basket_lang: en-us
basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
batch_size: 1
gpus: 1
inference:
condition_sample_rate: 24000
diff_k: 3
diff_steps: 100
diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_clean
duplicate_reference: true
exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
gpt_generate_args:
do_sample: true
min_new_tokens: 20
num_return_sequences: 50
use_cache: true
out_sample_rate: 24000
override_conditioning_features:
bad_text_proba: 0.0
c50: 100.0
dmcs_flatness: 100500.0
dmcs_roll_off_0.995: 100500.0
emo2vec: null
emotion: null
is_animation: 1.0
pitch_std: 100.0
snr: 100.0
year: 2025.0
reranking_options:
mode: MBR
top_k: 1
vocoder: bigvgan
voice_samples_preprocessing: []
num_workers: 1
output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__diffusion_yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest_freq_feats_noref_clean__2026-03-13_12-05-28
ref_dir: cc_20250825/ref
ticket: QUALITY-41
basket_generation_git_hash: a0fff333bc707fa86bab206dec923bbfe5b51d8c
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-03-13_12-05-28
|
|---|---|---|---|---|---|---|---|
|
DF-creative-commons-basket/03VPqxrlyxA_zh/F1__8.234-8.922
|
Every second.
|
||||||
|
DF-creative-commons-basket/b91pBJWJDhQ_ru/M0__9.240-10.650
|
Please forgive me.
|
||||||
|
DF-creative-commons-basket/CDxg_6317fA_ja/M1__4.346-5.698
|
Don't tell anyone.
|
||||||
|
DF-creative-commons-basket/l-l9x4bLVUY_zh/F0__11.480-13.480
|
It's all perfectly normal.
|
||||||
|
DF-creative-commons-basket/5Ges6LpYtI0_it/M0__9.020-10.034
|
Ah, did you quarrel?
|
||||||
|
DF-creative-commons-basket/jdSyxrG6dfM_it/M0__8.400-9.490
|
But I have a helmet.
|
||||||
|
DF-creative-commons-basket/pr5WOuKPrus_hi/F0__11.240-12.820
|
I won't let you go.
|
||||||
|
DF-creative-commons-basket/DGbTKywfSxw_ru/F0__10.360-12.490
|
What's troubling you, son?
|