diff --git "a/test-th-en.comet" "b/test-th-en.comet" new file mode 100644--- /dev/null +++ "b/test-th-en.comet" @@ -0,0 +1,1013 @@ +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 0 score: 0.6948 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1 score: 0.8730 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 2 score: 0.8525 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 3 score: 0.8060 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 4 score: 0.8556 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 5 score: 0.8187 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 6 score: 0.8535 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 7 score: 0.8447 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 8 score: 0.8426 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 9 score: 0.8395 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 10 score: 0.8594 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 11 score: 0.8452 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 12 score: 0.8751 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 13 score: 0.8424 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 14 score: 0.7758 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 15 score: 0.7902 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 16 score: 0.8919 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 17 score: 0.9089 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 18 score: 0.8836 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 19 score: 0.8688 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 20 score: 0.8661 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 21 score: 0.9130 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 22 score: 0.7145 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 23 score: 0.8672 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 24 score: 0.8590 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 25 score: 0.8718 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 26 score: 0.5910 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 27 score: 0.8815 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 28 score: 0.7852 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 29 score: 0.8130 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 30 score: 0.7833 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 31 score: 0.7295 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 32 score: 0.7437 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 33 score: 0.8710 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 34 score: 0.9031 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 35 score: 0.9245 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 36 score: 0.8602 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 37 score: 0.7391 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 38 score: 0.6488 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 39 score: 0.8055 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 40 score: 0.7183 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 41 score: 0.8969 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 42 score: 0.8804 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 43 score: 0.9197 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 44 score: 0.8156 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 45 score: 0.8173 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 46 score: 0.8339 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 47 score: 0.8796 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 48 score: 0.6620 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 49 score: 0.7309 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 50 score: 0.9013 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 51 score: 0.8989 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 52 score: 0.8585 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 53 score: 0.8025 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 54 score: 0.8110 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 55 score: 0.8511 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 56 score: 0.7449 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 57 score: 0.6186 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 58 score: 0.6859 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 59 score: 0.9031 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 60 score: 0.7707 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 61 score: 0.9196 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 62 score: 0.8551 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 63 score: 0.8871 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 64 score: 0.8738 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 65 score: 0.7466 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 66 score: 0.7863 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 67 score: 0.7830 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 68 score: 0.9258 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 69 score: 0.7576 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 70 score: 0.7902 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 71 score: 0.7119 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 72 score: 0.7370 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 73 score: 0.6663 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 74 score: 0.7197 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 75 score: 0.8287 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 76 score: 0.8303 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 77 score: 0.7958 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 78 score: 0.9776 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 79 score: 0.8912 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 80 score: 0.4667 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 81 score: 0.9160 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 82 score: 0.8936 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 83 score: 0.8298 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 84 score: 0.5348 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 85 score: 0.8390 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 86 score: 0.7581 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 87 score: 0.8593 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 88 score: 0.8667 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 89 score: 0.5735 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 90 score: 0.7231 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 91 score: 0.5815 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 92 score: 0.5826 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 93 score: 0.9036 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 94 score: 0.8890 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 95 score: 0.4995 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 96 score: 0.6649 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 97 score: 0.8532 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 98 score: 0.8453 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 99 score: 0.8589 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 100 score: 0.8861 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 101 score: 0.7318 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 102 score: 0.8365 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 103 score: 0.6815 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 104 score: 0.7980 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 105 score: 0.5568 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 106 score: 0.7507 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 107 score: 0.8470 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 108 score: 0.8040 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 109 score: 0.8023 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 110 score: 0.8004 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 111 score: 0.9457 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 112 score: 0.8746 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 113 score: 0.8399 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 114 score: 0.8550 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 115 score: 0.8202 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 116 score: 0.8768 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 117 score: 0.7497 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 118 score: 0.8275 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 119 score: 0.8691 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 120 score: 0.6513 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 121 score: 0.7812 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 122 score: 0.7700 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 123 score: 0.8936 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 124 score: 0.6294 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 125 score: 0.9212 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 126 score: 0.8172 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 127 score: 0.8902 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 128 score: 0.8951 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 129 score: 0.8734 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 130 score: 0.8945 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 131 score: 0.7122 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 132 score: 0.8561 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 133 score: 0.5277 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 134 score: 0.5536 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 135 score: 0.6231 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 136 score: 0.8857 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 137 score: 0.8722 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 138 score: 0.5610 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 139 score: 0.8543 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 140 score: 0.8303 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 141 score: 0.8331 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 142 score: 0.8114 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 143 score: 0.6388 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 144 score: 0.7732 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 145 score: 0.8648 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 146 score: 0.7284 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 147 score: 0.8487 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 148 score: 0.6126 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 149 score: 0.7906 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 150 score: 0.8075 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 151 score: 0.8316 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 152 score: 0.8156 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 153 score: 0.8081 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 154 score: 0.6518 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 155 score: 0.8312 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 156 score: 0.8165 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 157 score: 0.7004 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 158 score: 0.8677 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 159 score: 0.7272 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 160 score: 0.8455 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 161 score: 0.7486 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 162 score: 0.7467 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 163 score: 0.8755 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 164 score: 0.6892 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 165 score: 0.8085 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 166 score: 0.7618 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 167 score: 0.8541 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 168 score: 0.7903 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 169 score: 0.7672 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 170 score: 0.8282 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 171 score: 0.9022 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 172 score: 0.7955 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 173 score: 0.8354 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 174 score: 0.7974 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 175 score: 0.7018 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 176 score: 0.7283 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 177 score: 0.8158 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 178 score: 0.8098 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 179 score: 0.8639 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 180 score: 0.8532 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 181 score: 0.8561 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 182 score: 0.8561 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 183 score: 0.9036 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 184 score: 0.8617 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 185 score: 0.9182 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 186 score: 0.8361 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 187 score: 0.6769 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 188 score: 0.7435 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 189 score: 0.7542 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 190 score: 0.6820 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 191 score: 0.7700 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 192 score: 0.7909 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 193 score: 0.8159 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 194 score: 0.8252 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 195 score: 0.8768 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 196 score: 0.8242 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 197 score: 0.7636 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 198 score: 0.8058 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 199 score: 0.5683 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 200 score: 0.8623 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 201 score: 0.8220 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 202 score: 0.7768 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 203 score: 0.7369 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 204 score: 0.7618 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 205 score: 0.6703 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 206 score: 0.8375 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 207 score: 0.7956 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 208 score: 0.8043 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 209 score: 0.7957 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 210 score: 0.9302 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 211 score: 0.8778 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 212 score: 0.8804 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 213 score: 0.9355 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 214 score: 0.8684 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 215 score: 0.8194 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 216 score: 0.9074 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 217 score: 0.8307 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 218 score: 0.8684 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 219 score: 0.8978 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 220 score: 0.7536 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 221 score: 0.9165 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 222 score: 0.8771 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 223 score: 0.7390 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 224 score: 0.8157 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 225 score: 0.8659 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 226 score: 0.7751 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 227 score: 0.6899 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 228 score: 0.7449 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 229 score: 0.8253 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 230 score: 0.8953 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 231 score: 0.6639 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 232 score: 0.8608 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 233 score: 0.7416 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 234 score: 0.8173 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 235 score: 0.8611 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 236 score: 0.7679 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 237 score: 0.8729 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 238 score: 0.8652 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 239 score: 0.7727 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 240 score: 0.8009 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 241 score: 0.6715 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 242 score: 0.6723 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 243 score: 0.7174 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 244 score: 0.6891 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 245 score: 0.6031 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 246 score: 0.7554 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 247 score: 0.7552 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 248 score: 0.6318 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 249 score: 0.8924 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 250 score: 0.8783 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 251 score: 0.9156 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 252 score: 0.8141 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 253 score: 0.7061 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 254 score: 0.5715 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 255 score: 0.8247 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 256 score: 0.7895 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 257 score: 0.8409 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 258 score: 0.6599 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 259 score: 0.8270 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 260 score: 0.8976 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 261 score: 0.7747 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 262 score: 0.7709 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 263 score: 0.8780 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 264 score: 0.9214 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 265 score: 0.8704 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 266 score: 0.8789 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 267 score: 0.8871 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 268 score: 0.9287 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 269 score: 0.8022 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 270 score: 0.8043 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 271 score: 0.6930 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 272 score: 0.7390 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 273 score: 0.7300 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 274 score: 0.7345 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 275 score: 0.8300 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 276 score: 0.7471 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 277 score: 0.6992 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 278 score: 0.8137 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 279 score: 0.7952 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 280 score: 0.9049 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 281 score: 0.8168 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 282 score: 0.7861 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 283 score: 0.8102 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 284 score: 0.8820 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 285 score: 0.9224 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 286 score: 0.7968 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 287 score: 0.7838 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 288 score: 0.8795 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 289 score: 0.7914 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 290 score: 0.7443 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 291 score: 0.7836 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 292 score: 0.6975 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 293 score: 0.8615 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 294 score: 0.7717 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 295 score: 0.7398 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 296 score: 0.7865 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 297 score: 0.6917 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 298 score: 0.8496 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 299 score: 0.8172 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 300 score: 0.9214 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 301 score: 0.8169 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 302 score: 0.8892 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 303 score: 0.8950 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 304 score: 0.8804 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 305 score: 0.8248 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 306 score: 0.8677 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 307 score: 0.8502 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 308 score: 0.9228 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 309 score: 0.8083 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 310 score: 0.5867 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 311 score: 0.8265 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 312 score: 0.6337 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 313 score: 0.7410 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 314 score: 0.9176 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 315 score: 0.8144 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 316 score: 0.6499 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 317 score: 0.7679 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 318 score: 0.7704 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 319 score: 0.7524 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 320 score: 0.8219 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 321 score: 0.8197 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 322 score: 0.8749 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 323 score: 0.8914 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 324 score: 0.9046 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 325 score: 0.8434 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 326 score: 0.7800 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 327 score: 0.8748 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 328 score: 0.8874 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 329 score: 0.8473 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 330 score: 0.8624 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 331 score: 0.7956 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 332 score: 0.8819 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 333 score: 0.9022 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 334 score: 0.9540 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 335 score: 0.9246 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 336 score: 0.8879 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 337 score: 0.8024 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 338 score: 0.7963 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 339 score: 0.8858 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 340 score: 0.8406 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 341 score: 0.6309 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 342 score: 0.6944 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 343 score: 0.7395 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 344 score: 0.8325 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 345 score: 0.7738 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 346 score: 0.9009 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 347 score: 0.8303 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 348 score: 0.5084 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 349 score: 0.7296 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 350 score: 0.8677 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 351 score: 0.6420 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 352 score: 0.8315 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 353 score: 0.8440 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 354 score: 0.5933 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 355 score: 0.8306 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 356 score: 0.7339 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 357 score: 0.6013 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 358 score: 0.8038 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 359 score: 0.7707 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 360 score: 0.7229 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 361 score: 0.7119 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 362 score: 0.8662 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 363 score: 0.6768 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 364 score: 0.8019 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 365 score: 0.6218 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 366 score: 0.6398 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 367 score: 0.6898 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 368 score: 0.7525 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 369 score: 0.8177 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 370 score: 0.8073 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 371 score: 0.7608 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 372 score: 0.8600 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 373 score: 0.8127 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 374 score: 0.8696 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 375 score: 0.8250 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 376 score: 0.9050 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 377 score: 0.8522 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 378 score: 0.4346 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 379 score: 0.7268 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 380 score: 0.8130 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 381 score: 0.8390 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 382 score: 0.6837 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 383 score: 0.8179 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 384 score: 0.6145 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 385 score: 0.7737 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 386 score: 0.6828 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 387 score: 0.7748 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 388 score: 0.5303 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 389 score: 0.8226 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 390 score: 0.8630 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 391 score: 0.5377 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 392 score: 0.7314 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 393 score: 0.6275 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 394 score: 0.6409 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 395 score: 0.5590 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 396 score: 0.7036 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 397 score: 0.8714 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 398 score: 0.8231 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 399 score: 0.7988 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 400 score: 0.8415 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 401 score: 0.7492 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 402 score: 0.7368 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 403 score: 0.7474 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 404 score: 0.8059 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 405 score: 0.5949 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 406 score: 0.8058 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 407 score: 0.7471 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 408 score: 0.8885 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 409 score: 0.8680 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 410 score: 0.8342 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 411 score: 0.8466 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 412 score: 0.7480 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 413 score: 0.8848 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 414 score: 0.8382 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 415 score: 0.8357 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 416 score: 0.8865 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 417 score: 0.8497 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 418 score: 0.8596 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 419 score: 0.9276 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 420 score: 0.9321 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 421 score: 0.8792 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 422 score: 0.8580 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 423 score: 0.8124 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 424 score: 0.8360 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 425 score: 0.9045 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 426 score: 0.8702 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 427 score: 0.8746 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 428 score: 0.8388 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 429 score: 0.7564 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 430 score: 0.4914 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 431 score: 0.8243 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 432 score: 0.7660 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 433 score: 0.8297 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 434 score: 0.8196 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 435 score: 0.5146 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 436 score: 0.6301 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 437 score: 0.6039 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 438 score: 0.8419 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 439 score: 0.8367 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 440 score: 0.8180 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 441 score: 0.8778 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 442 score: 0.7815 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 443 score: 0.7764 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 444 score: 0.7045 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 445 score: 0.8934 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 446 score: 0.7743 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 447 score: 0.8814 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 448 score: 0.8140 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 449 score: 0.7939 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 450 score: 0.7976 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 451 score: 0.7912 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 452 score: 0.7410 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 453 score: 0.5190 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 454 score: 0.8915 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 455 score: 0.7635 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 456 score: 0.7358 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 457 score: 0.8931 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 458 score: 0.6402 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 459 score: 0.6974 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 460 score: 0.5965 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 461 score: 0.8490 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 462 score: 0.7088 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 463 score: 0.4703 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 464 score: 0.5782 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 465 score: 0.4038 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 466 score: 0.7903 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 467 score: 0.6422 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 468 score: 0.6879 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 469 score: 0.5844 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 470 score: 0.7471 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 471 score: 0.7341 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 472 score: 0.6472 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 473 score: 0.7051 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 474 score: 0.7427 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 475 score: 0.8030 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 476 score: 0.8235 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 477 score: 0.7812 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 478 score: 0.8719 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 479 score: 0.9111 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 480 score: 0.8161 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 481 score: 0.7513 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 482 score: 0.7767 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 483 score: 0.7675 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 484 score: 0.8462 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 485 score: 0.8460 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 486 score: 0.5899 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 487 score: 0.7693 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 488 score: 0.7771 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 489 score: 0.4678 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 490 score: 0.6791 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 491 score: 0.7994 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 492 score: 0.9424 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 493 score: 0.9052 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 494 score: 0.7816 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 495 score: 0.8693 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 496 score: 0.8013 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 497 score: 0.7371 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 498 score: 0.7342 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 499 score: 0.7403 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 500 score: 0.9159 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 501 score: 0.8932 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 502 score: 0.8340 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 503 score: 0.8052 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 504 score: 0.6216 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 505 score: 0.8524 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 506 score: 0.7824 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 507 score: 0.8371 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 508 score: 0.8912 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 509 score: 0.8827 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 510 score: 0.8534 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 511 score: 0.8292 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 512 score: 0.8494 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 513 score: 0.8311 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 514 score: 0.8344 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 515 score: 0.8434 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 516 score: 0.7437 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 517 score: 0.9072 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 518 score: 0.7940 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 519 score: 0.9333 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 520 score: 0.7655 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 521 score: 0.7685 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 522 score: 0.7580 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 523 score: 0.6815 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 524 score: 0.8053 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 525 score: 0.6148 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 526 score: 0.7788 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 527 score: 0.7616 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 528 score: 0.7620 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 529 score: 0.8540 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 530 score: 0.5404 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 531 score: 0.8146 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 532 score: 0.6332 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 533 score: 0.7186 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 534 score: 0.5153 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 535 score: 0.5948 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 536 score: 0.7828 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 537 score: 0.8275 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 538 score: 0.8422 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 539 score: 0.7724 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 540 score: 0.8332 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 541 score: 0.8359 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 542 score: 0.3716 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 543 score: 0.4247 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 544 score: 0.8519 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 545 score: 0.8873 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 546 score: 0.7296 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 547 score: 0.8846 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 548 score: 0.6930 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 549 score: 0.9244 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 550 score: 0.8500 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 551 score: 0.8207 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 552 score: 0.8272 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 553 score: 0.7231 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 554 score: 0.7447 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 555 score: 0.7911 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 556 score: 0.5269 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 557 score: 0.7137 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 558 score: 0.7910 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 559 score: 0.6065 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 560 score: 0.5079 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 561 score: 0.7481 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 562 score: 0.8569 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 563 score: 0.8299 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 564 score: 0.8432 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 565 score: 0.8055 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 566 score: 0.7930 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 567 score: 0.8779 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 568 score: 0.7868 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 569 score: 0.8075 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 570 score: 0.7380 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 571 score: 0.8969 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 572 score: 0.7496 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 573 score: 0.7796 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 574 score: 0.6257 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 575 score: 0.8202 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 576 score: 0.6988 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 577 score: 0.8349 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 578 score: 0.9046 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 579 score: 0.8096 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 580 score: 0.6630 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 581 score: 0.5046 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 582 score: 0.5006 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 583 score: 0.5541 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 584 score: 0.6363 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 585 score: 0.8064 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 586 score: 0.9272 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 587 score: 0.8898 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 588 score: 0.6849 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 589 score: 0.8885 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 590 score: 0.8668 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 591 score: 0.8262 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 592 score: 0.8320 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 593 score: 0.6542 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 594 score: 0.8900 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 595 score: 0.7555 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 596 score: 0.5654 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 597 score: 0.7856 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 598 score: 0.4665 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 599 score: 0.8489 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 600 score: 0.7799 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 601 score: 0.6651 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 602 score: 0.7521 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 603 score: 0.8489 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 604 score: 0.9033 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 605 score: 0.8935 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 606 score: 0.8241 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 607 score: 0.7311 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 608 score: 0.5538 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 609 score: 0.6743 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 610 score: 0.7931 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 611 score: 0.7336 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 612 score: 0.8337 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 613 score: 0.5034 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 614 score: 0.8405 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 615 score: 0.6840 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 616 score: 0.7328 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 617 score: 0.8532 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 618 score: 0.8092 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 619 score: 0.4516 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 620 score: 0.8396 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 621 score: 0.4388 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 622 score: 0.8926 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 623 score: 0.8732 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 624 score: 0.7407 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 625 score: 0.8494 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 626 score: 0.9212 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 627 score: 0.8383 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 628 score: 0.9088 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 629 score: 0.8823 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 630 score: 0.8383 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 631 score: 0.8270 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 632 score: 0.8324 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 633 score: 0.9101 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 634 score: 0.7997 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 635 score: 0.8863 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 636 score: 0.8917 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 637 score: 0.8633 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 638 score: 0.4799 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 639 score: 0.8542 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 640 score: 0.5180 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 641 score: 0.8603 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 642 score: 0.2227 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 643 score: 0.6262 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 644 score: 0.7839 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 645 score: 0.8524 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 646 score: 0.7551 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 647 score: 0.8595 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 648 score: 0.6320 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 649 score: 0.5225 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 650 score: 0.6491 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 651 score: 0.4118 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 652 score: 0.5723 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 653 score: 0.7962 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 654 score: 0.7870 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 655 score: 0.9124 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 656 score: 0.9044 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 657 score: 0.4702 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 658 score: 0.8027 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 659 score: 0.8237 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 660 score: 0.8459 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 661 score: 0.8100 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 662 score: 0.8332 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 663 score: 0.8358 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 664 score: 0.7902 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 665 score: 0.8514 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 666 score: 0.8217 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 667 score: 0.7863 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 668 score: 0.8099 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 669 score: 0.8161 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 670 score: 0.8595 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 671 score: 0.8810 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 672 score: 0.9170 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 673 score: 0.7921 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 674 score: 0.8300 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 675 score: 0.7992 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 676 score: 0.8504 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 677 score: 0.8473 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 678 score: 0.8439 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 679 score: 0.8788 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 680 score: 0.8207 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 681 score: 0.9541 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 682 score: 0.8875 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 683 score: 0.8784 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 684 score: 0.8155 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 685 score: 0.8273 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 686 score: 0.8985 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 687 score: 0.7419 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 688 score: 0.8901 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 689 score: 0.7877 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 690 score: 0.9250 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 691 score: 0.7709 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 692 score: 0.8527 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 693 score: 0.7221 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 694 score: 0.6635 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 695 score: 0.9072 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 696 score: 0.6594 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 697 score: 0.9129 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 698 score: 0.8127 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 699 score: 0.7654 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 700 score: 0.8669 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 701 score: 0.8524 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 702 score: 0.8542 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 703 score: 0.8596 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 704 score: 0.4715 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 705 score: 0.7266 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 706 score: 0.8305 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 707 score: 0.8659 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 708 score: 0.9014 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 709 score: 0.7960 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 710 score: 0.9057 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 711 score: 0.7972 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 712 score: 0.8842 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 713 score: 0.8753 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 714 score: 0.8665 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 715 score: 0.6698 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 716 score: 0.7660 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 717 score: 0.8698 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 718 score: 0.8395 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 719 score: 0.8208 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 720 score: 0.7591 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 721 score: 0.7113 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 722 score: 0.8658 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 723 score: 0.8811 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 724 score: 0.8995 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 725 score: 0.8392 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 726 score: 0.7911 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 727 score: 0.7685 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 728 score: 0.7249 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 729 score: 0.8855 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 730 score: 0.6687 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 731 score: 0.8069 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 732 score: 0.8538 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 733 score: 0.8591 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 734 score: 0.5274 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 735 score: 0.8895 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 736 score: 0.8316 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 737 score: 0.7894 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 738 score: 0.8356 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 739 score: 0.8700 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 740 score: 0.8642 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 741 score: 0.7202 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 742 score: 0.8356 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 743 score: 0.8803 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 744 score: 0.8403 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 745 score: 0.8687 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 746 score: 0.9150 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 747 score: 0.8112 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 748 score: 0.7609 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 749 score: 0.9143 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 750 score: 0.9135 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 751 score: 0.7495 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 752 score: 0.8229 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 753 score: 0.7712 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 754 score: 0.9048 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 755 score: 0.6061 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 756 score: 0.8805 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 757 score: 0.8767 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 758 score: 0.8348 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 759 score: 0.7178 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 760 score: 0.9093 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 761 score: 0.8124 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 762 score: 0.9112 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 763 score: 0.8600 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 764 score: 0.8940 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 765 score: 0.7835 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 766 score: 0.8832 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 767 score: 0.7253 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 768 score: 0.7626 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 769 score: 0.6600 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 770 score: 0.8437 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 771 score: 0.8363 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 772 score: 0.7927 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 773 score: 0.6123 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 774 score: 0.8111 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 775 score: 0.6454 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 776 score: 0.4566 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 777 score: 0.7791 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 778 score: 0.7070 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 779 score: 0.6563 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 780 score: 0.8742 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 781 score: 0.7983 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 782 score: 0.6067 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 783 score: 0.6420 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 784 score: 0.8408 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 785 score: 0.6308 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 786 score: 0.9191 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 787 score: 0.4969 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 788 score: 0.9260 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 789 score: 0.8520 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 790 score: 0.8732 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 791 score: 0.8436 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 792 score: 0.8293 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 793 score: 0.7990 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 794 score: 0.8018 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 795 score: 0.6240 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 796 score: 0.9369 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 797 score: 0.7598 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 798 score: 0.7225 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 799 score: 0.7774 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 800 score: 0.8644 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 801 score: 0.7724 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 802 score: 0.8536 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 803 score: 0.7901 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 804 score: 0.6056 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 805 score: 0.4667 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 806 score: 0.7212 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 807 score: 0.9150 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 808 score: 0.8552 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 809 score: 0.7527 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 810 score: 0.6341 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 811 score: 0.8127 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 812 score: 0.9149 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 813 score: 0.9172 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 814 score: 0.9002 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 815 score: 0.7664 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 816 score: 0.8549 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 817 score: 0.7193 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 818 score: 0.6538 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 819 score: 0.8446 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 820 score: 0.8628 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 821 score: 0.7925 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 822 score: 0.8290 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 823 score: 0.7580 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 824 score: 0.7176 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 825 score: 0.8377 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 826 score: 0.8798 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 827 score: 0.8411 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 828 score: 0.8154 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 829 score: 0.7930 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 830 score: 0.8579 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 831 score: 0.8312 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 832 score: 0.7946 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 833 score: 0.7913 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 834 score: 0.8237 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 835 score: 0.7942 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 836 score: 0.7212 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 837 score: 0.5619 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 838 score: 0.7564 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 839 score: 0.8937 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 840 score: 0.8459 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 841 score: 0.6915 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 842 score: 0.4701 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 843 score: 0.7958 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 844 score: 0.7875 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 845 score: 0.5679 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 846 score: 0.5474 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 847 score: 0.7922 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 848 score: 0.7731 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 849 score: 0.5191 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 850 score: 0.7680 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 851 score: 0.8005 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 852 score: 0.7224 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 853 score: 0.6510 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 854 score: 0.8251 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 855 score: 0.6503 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 856 score: 0.6497 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 857 score: 0.8853 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 858 score: 0.5854 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 859 score: 0.8106 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 860 score: 0.7873 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 861 score: 0.8531 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 862 score: 0.8750 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 863 score: 0.7438 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 864 score: 0.7575 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 865 score: 0.7357 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 866 score: 0.8482 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 867 score: 0.8745 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 868 score: 0.7497 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 869 score: 0.8076 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 870 score: 0.7588 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 871 score: 0.9023 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 872 score: 0.7794 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 873 score: 0.8860 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 874 score: 0.9099 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 875 score: 0.5162 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 876 score: 0.8442 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 877 score: 0.8780 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 878 score: 0.8813 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 879 score: 0.6653 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 880 score: 0.8354 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 881 score: 0.8875 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 882 score: 0.7514 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 883 score: 0.8402 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 884 score: 0.7960 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 885 score: 0.7968 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 886 score: 0.7808 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 887 score: 0.6994 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 888 score: 0.5114 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 889 score: 0.5262 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 890 score: 0.7956 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 891 score: 0.6674 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 892 score: 0.7359 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 893 score: 0.7627 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 894 score: 0.7538 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 895 score: 0.8989 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 896 score: 0.7455 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 897 score: 0.6971 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 898 score: 0.9108 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 899 score: 0.6859 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 900 score: 0.8236 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 901 score: 0.7197 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 902 score: 0.5542 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 903 score: 0.7759 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 904 score: 0.7445 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 905 score: 0.7440 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 906 score: 0.7058 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 907 score: 0.6199 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 908 score: 0.8523 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 909 score: 0.9486 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 910 score: 0.8696 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 911 score: 0.8276 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 912 score: 0.9386 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 913 score: 0.8615 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 914 score: 0.8943 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 915 score: 0.8890 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 916 score: 0.8778 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 917 score: 0.7215 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 918 score: 0.6806 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 919 score: 0.8329 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 920 score: 0.8590 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 921 score: 0.8398 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 922 score: 0.7980 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 923 score: 0.7721 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 924 score: 0.8902 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 925 score: 0.8287 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 926 score: 0.9173 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 927 score: 0.8541 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 928 score: 0.8892 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 929 score: 0.8063 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 930 score: 0.8366 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 931 score: 0.8134 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 932 score: 0.7476 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 933 score: 0.6230 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 934 score: 0.7454 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 935 score: 0.8942 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 936 score: 0.9040 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 937 score: 0.9496 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 938 score: 0.8972 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 939 score: 0.8968 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 940 score: 0.8265 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 941 score: 0.9245 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 942 score: 0.8900 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 943 score: 0.7367 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 944 score: 0.7476 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 945 score: 0.7463 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 946 score: 0.5596 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 947 score: 0.8606 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 948 score: 0.6614 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 949 score: 0.9311 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 950 score: 0.8318 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 951 score: 0.8584 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 952 score: 0.7937 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 953 score: 0.3905 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 954 score: 0.7915 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 955 score: 0.7470 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 956 score: 0.8720 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 957 score: 0.8007 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 958 score: 0.8122 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 959 score: 0.6280 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 960 score: 0.6891 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 961 score: 0.6729 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 962 score: 0.8925 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 963 score: 0.8427 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 964 score: 0.9272 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 965 score: 0.8644 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 966 score: 0.7573 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 967 score: 0.6728 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 968 score: 0.7553 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 969 score: 0.6398 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 970 score: 0.8605 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 971 score: 0.8262 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 972 score: 0.8051 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 973 score: 0.8288 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 974 score: 0.8217 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 975 score: 0.9223 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 976 score: 0.8705 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 977 score: 0.8316 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 978 score: 0.8068 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 979 score: 0.7717 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 980 score: 0.7483 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 981 score: 0.8764 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 982 score: 0.6827 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 983 score: 0.8004 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 984 score: 0.8452 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 985 score: 0.9022 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 986 score: 0.7688 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 987 score: 0.7626 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 988 score: 0.7839 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 989 score: 0.7000 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 990 score: 0.8767 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 991 score: 0.8637 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 992 score: 0.8595 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 993 score: 0.7549 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 994 score: 0.8303 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 995 score: 0.7769 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 996 score: 0.6002 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 997 score: 0.9398 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 998 score: 0.7903 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 999 score: 0.6073 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1000 score: 0.8739 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1001 score: 0.8735 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1002 score: 0.7908 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1003 score: 0.4519 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1004 score: 0.8526 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1005 score: 0.9312 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1006 score: 0.7448 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1007 score: 0.7069 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1008 score: 0.8786 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1009 score: 0.8520 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1010 score: 0.8142 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en Segment 1011 score: 0.8666 +/beacon-scratch/tongzh24/ALMA-checkpoint/exp_16_languages/alma-13b-sft-16-languages-th-max-tokens-512//test-th-en score: 0.7861