Prev
# uttid text ref baseline_tgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15__2026-01-19_14-20-36
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-19_14-20-36
indextts2 literategoggles_idxdistill
basket_config_path: /mount/s3/tts-binary-data-nb/dchebakov/metrics/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: /mount/s3/tts-binary-data-nb/dchebakov/metrics/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      duplicate_reference: false
      exp: /mount/s3/tts-binary-data-nb/dchebakov/models/yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo
      gpt_generate_args:
        do_sample: true
        enforce_silent_start: false
        num_return_sequences: 30
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        pitch_std: 100.0
        snr: 100.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825_en-US/yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo__2026-01-10_07-59-19
    ref_dir: cc_20250825_en-US/ref
    ticket: QUALITY-000
  basket_generation_git_hash: 084ccc0c4313e646e3630e4e5a35b7e04d70fdad
model_data_type: tts-cloning
ticket: QUALITY-000
version: 2026-01-10_07-59-19
indexrefencoder_ft_enus
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-23_13-08-47
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_13-08-47
movies2_finetune_closest_anim
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15__2026-01-14_17-25-18
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-14_17-25-18
baseline
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15__2026-01-23_18-31-03
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_18-31-03
indexrefencoder_ft_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-23_17-32-54
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_17-32-54
movies2_finetune_closest_anim_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15__2026-01-23_18-01-37
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_18-01-37
indexrefencoder_ft_movies3_closest_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-26_11-26-27
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-26_11-26-27
indexrefencoder_ft_movies3_20Kit_closest_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-28_13-09-10
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-28_13-09-10
indexrefencoder_noft_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-28_16-42-41
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-28_16-42-41
DF-creative-commons-basket/cyLLcnTzV5w_pt/F1__1.290-6.550
I was super happy, super excited, wanting a little bit more.
DF-creative-commons-basket/SQV7gz091-0_es/M0__0.180-2.485
Am I not here for the divorce papers?
DF-creative-commons-basket/xtiiG-k5ejA_de/F0__5.840-9.360
I'm afraid that dream will become reality.
DF-creative-commons-basket/DGbTKywfSxw_ru/F0__6.320-8.405
Tomorrow you will face a trial.
DF-creative-commons-basket/d_TKW3INVKs_it/M1__6.500-8.130
But how many meters does the child need?
DF-creative-commons-basket/-d9C5FyXvvw_it/M2__10.259-11.942
It's a season ticket in the curve section.
DF-creative-commons-basket/2iK4DdnoL-s_pt/F1__6.521-7.169
I brought them.
DF-creative-commons-basket/uTbCKWoc1YY_zh/F0__18.920-19.913
You're not coming home today?
DF-creative-commons-basket/2BJj_jAbQSw_pt/F1__5.060-6.178
How coward you are..
DF-creative-commons-basket/-d9C5FyXvvw_it/M0__12.929-13.539
He did it.
DF-creative-commons-basket/qJaAnEUiO6E_it/F0__17.780-18.690
I'm sure.
DF-creative-commons-basket/5k25FVLoCyo_ru/F0__9.440-11.890
You know, our mom taught us to respect our father.
DF-creative-commons-basket/QTa0hoQDonY_pt/M0__3.370-5.520
Luciano, I don't feel comfortable in my own home anymore.
DF-creative-commons-basket/OesbsUiwgo0_pt/F0__7.550-8.920
You can get back to work.
DF-creative-commons-basket/P8yVtZpEY7Y_es/F0__6.550-8.390
You're not going to compare that ugly woman to me.
DF-creative-commons-basket/2BJj_jAbQSw_pt/F0__3.050-4.300
What are you talking about?
DF-creative-commons-basket/mQSxiD4Kp6c_es/F0__10.115-14.715
But don't let him fool you because he's not as much of a gentleman as you are, even though he's carrying a lady in his arms.
DF-creative-commons-basket/P6glOSnV2gM_fr/M0__17.640-19.710
From the beginning, you decide everything on your own and then you run away.
DF-creative-commons-basket/jdSyxrG6dfM_it/M0__3.300-8.250
Well yes, actually yes, I wouldn't have anything to offer you and then, well...
DF-creative-commons-basket/2BJj_jAbQSw_pt/F1__7.410-12.034
Your low blow of abandoning my daughter when she needed you most.
Next