# uttid text ref baseline_tgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15__2026-01-19_14-20-36
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-19_14-20-36
indextts2 literategoggles_idxdistill
basket_config_path: /mount/s3/tts-binary-data-nb/dchebakov/metrics/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: /mount/s3/tts-binary-data-nb/dchebakov/metrics/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      duplicate_reference: false
      exp: /mount/s3/tts-binary-data-nb/dchebakov/models/yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo
      gpt_generate_args:
        do_sample: true
        enforce_silent_start: false
        num_return_sequences: 30
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        pitch_std: 100.0
        snr: 100.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825_en-US/yt4_wavtokenizer_16K_lossent0.15_idxdistill_1ref_emo__2026-01-10_07-59-19
    ref_dir: cc_20250825_en-US/ref
    ticket: QUALITY-000
  basket_generation_git_hash: 084ccc0c4313e646e3630e4e5a35b7e04d70fdad
model_data_type: tts-cloning
ticket: QUALITY-000
version: 2026-01-10_07-59-19
indexrefencoder_ft_enus
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-23_13-08-47
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_13-08-47
movies2_finetune_closest_anim
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      target_len_rate: 1.0
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15__2026-01-14_17-25-18
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-14_17-25-18
baseline
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15__yt4_wavtokenizer_16K_lossent0.15__2026-01-23_18-31-03
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_18-31-03
indexrefencoder_ft_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 0.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies2_finetune_full_v2__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-23_17-32-54
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_17-32-54
movies2_finetune_closest_anim_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_movies2_finetune_closest__yt4_wavtokenizer_16K_lossent0.15__2026-01-23_18-01-37
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-23_18-01-37
indexrefencoder_ft_movies3_closest_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-26_11-26-27
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-26_11-26-27
indexrefencoder_ft_movies3_20Kit_closest_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self_movies3_finetune_full_v2_closest__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-28_13-09-10
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-28_13-09-10
indexrefencoder_noft_enus_notgtlen
basket_config_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
data_meta: null
exp_name: yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
lang: en-us
meta:
  basket_generation_config:
    basket_lang: en-us
    basket_path: quality/tts/tortoise-baskets/cc_20250825_en-US.json
    batch_size: 1
    gpus: 1
    inference:
      condition_sample_rate: 24000
      diff_k: 3
      diff_steps: 100
      diffusion_exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      duplicate_reference: true
      exp: /mount/s3/tts-binary-data-nb/eg/exp/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self
      gpt_generate_args:
        do_sample: true
        min_new_tokens: 20
        num_return_sequences: 50
        use_cache: true
      out_sample_rate: 24000
      override_conditioning_features:
        bad_text_proba: 0.0
        c50: 0.0
        dmcs_flatness: 100500.0
        dmcs_roll_off_0.995: 100500.0
        emo2vec: null
        emotion: null
        is_animation: 1.0
        pitch_std: 100.0
        snr: 100.0
        year: 2025.0
      reranking_options:
        mode: MBR
        top_k: 1
      vocoder: bigvgan
      voice_samples_preprocessing: []
    num_workers: 1
    output_dir: cc_20250825/yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__yt4_wavtokenizer_16K_lossent0.15_indexrefencoder_separate_self__2026-01-28_16-42-41
    ref_dir: cc_20250825/ref
    ticket: QUALITY-41
  basket_generation_git_hash: 7ba982d9bb8ddc0cb968d517f583b0227d2624ed
model_data_type: tts-cloning
ticket: QUALITY-41
version: 2026-01-28_16-42-41
DF-creative-commons-basket/_W3R1g9-ByQ_ru/M0__2.750-5.130
Bobby, do something!
DF-creative-commons-basket/mFpCHV8M_kU_ru/F0__10.080-12.450
Is this respect for the audience and the judges?
DF-creative-commons-basket/Q-jR0549ZkA_de/F0__11.010-15.110
How happy they must be to already be thinking of May at Christmas!
DF-creative-commons-basket/31p-0IxN0XU_de/F0__2.680-4.810
No one knew if it should all be wasted on drink or put to better use.
DF-creative-commons-basket/9ekEozd_9QE_ja/M0__0.000-7.503
I'll survive by any means necessary, but you're already... done.
DF-creative-commons-basket/P8yVtZpEY7Y_es/M0__11.990-13.525
Would you like me to give you my humble opinion?
DF-creative-commons-basket/5k25FVLoCyo_ru/F0__0.120-4.410
When I put these glasses on, I think dad wore them.
DF-creative-commons-basket/mtohEFEZLBE_pt/M0__5.110-6.300
Tell me, Carminha!
DF-creative-commons-basket/w35s_xMyHIc_pt/F0__12.090-15.040
They even made me call my mother to tell her everything was fine.
DF-creative-commons-basket/JZB_ti6Ctlc_it/M0__18.890-19.810
The Pinson.
DF-creative-commons-basket/_5MiAdlrqgA_pt/F0__6.240-8.890
The love that is born in a moment like that
DF-creative-commons-basket/x_gaj_ZhDqs_de/M0__16.030-18.420
Stick to what you know best.
DF-creative-commons-basket/SQV7gz091-0_es/M1__3.395-6.245
I have come as an ambassador of love.
DF-creative-commons-basket/lTa1VcbYZEE_fr/F0__6.880-7.950
What did he tell you?
DF-creative-commons-basket/BFvg-l-e5FE_es/F1__15.220-18.730
As far as I know, you were also born in a hill village, missis Isabella.
DF-creative-commons-basket/4mN7zx-qOt4_zh/F0__4.718-6.860
If possible, I would willingly switch to a different position.
DF-creative-commons-basket/KCV3wodv9w4_ja/F0__22.892-24.696
Even though you say you don't want to get married.
DF-creative-commons-basket/-d9C5FyXvvw_it/M0__3.153-4.752
I only love the bar.
DF-creative-commons-basket/Mzp-y0qfgYQ_pt/M0__4.920-9.250
I want us to remember, if that is the case, with nostalgia for all the good that we lived together.
DF-creative-commons-basket/b91pBJWJDhQ_ru/F1__12.640-13.888
And you forgive me.
Next