hadeel295/liar2_processed_binary
2134
1---2dataset_info:3- config_name: default4 features:5 - name: id6 dtype: int647 - name: label8 dtype: int649 - name: statement10 dtype: string11 - name: date12 dtype: string13 - name: subject14 dtype: string15 - name: speaker16 dtype: string17 - name: speaker_description18 dtype: string19 - name: state_info20 dtype: string21 - name: true_counts22 dtype: int6423 - name: mostly_true_counts24 dtype: int6425 - name: half_true_counts26 dtype: int6427 - name: mostly_false_counts28 dtype: int6429 - name: false_counts30 dtype: int6431 - name: pants_on_fire_counts32 dtype: int6433 - name: context34 dtype: string35 - name: justification36 dtype: string37 - name: binary_label38 dtype: int6439 - name: clean_statement40 dtype: string41 - name: clean_justification42 dtype: string43 - name: input_ids44 list: int3245 - name: token_type_ids46 list: int847 - name: attention_mask48 list: int849 splits:50 - name: train51 num_bytes: 4684192352 num_examples: 1836953 - name: validation54 num_bytes: 584499655 num_examples: 229756 - name: test57 num_bytes: 587266158 num_examples: 229659 download_size: 5597669260 dataset_size: 5855958061- config_name: llama3_balanced_sample62 features:63 - name: id64 dtype: int6465 - name: label66 dtype: int6467 - name: statement68 dtype: string69 - name: date70 dtype: string71 - name: subject72 dtype: string73 - name: speaker74 dtype: string75 - name: speaker_description76 dtype: string77 - name: state_info78 dtype: string79 - name: true_counts80 dtype: int6481 - name: mostly_true_counts82 dtype: int6483 - name: half_true_counts84 dtype: int6485 - name: mostly_false_counts86 dtype: int6487 - name: false_counts88 dtype: int6489 - name: pants_on_fire_counts90 dtype: int6491 - name: context92 dtype: string93 - name: justification94 dtype: string95 - name: binary_label96 dtype: int6497 - name: clean_statement98 dtype: string99 - name: clean_justification100 dtype: string101 - name: input_ids102 list: int32103 - name: token_type_ids104 list: int8105 - name: attention_mask106 list: int8107 - name: computed_label_text108 dtype: string109 - name: llama3_explanation110 dtype: string111 splits:112 - name: train113 num_bytes: 189536114 num_examples: 60115 download_size: 192607116 dataset_size: 189536117- config_name: llama3_sample118 features:119 - name: id120 dtype: int64121 - name: label122 dtype: int64123 - name: statement124 dtype: string125 - name: date126 dtype: string127 - name: subject128 dtype: string129 - name: speaker130 dtype: string131 - name: speaker_description132 dtype: string133 - name: state_info134 dtype: string135 - name: true_counts136 dtype: int64137 - name: mostly_true_counts138 dtype: int64139 - name: half_true_counts140 dtype: int64141 - name: mostly_false_counts142 dtype: int64143 - name: false_counts144 dtype: int64145 - name: pants_on_fire_counts146 dtype: int64147 - name: context148 dtype: string149 - name: justification150 dtype: string151 - name: binary_label152 dtype: int64153 - name: clean_statement154 dtype: string155 - name: clean_justification156 dtype: string157 - name: input_ids158 list: int32159 - name: token_type_ids160 list: int8161 - name: attention_mask162 list: int8163 - name: llama3_explanation164 dtype: string165 splits:166 - name: train167 num_bytes: 182568168 num_examples: 60169 download_size: 191337170 dataset_size: 182568171configs:172- config_name: default173 data_files:174 - split: train175 path: data/train-*176 - split: validation177 path: data/validation-*178 - split: test179 path: data/test-*180- config_name: llama3_balanced_sample181 data_files:182 - split: train183 path: llama3_balanced_sample/train-*184- config_name: llama3_sample185 data_files:186 - split: train187 path: llama3_sample/train-*188---189 