Team Ai
Datasetpublic

hadeel295/liar2_processed_binary

sourceHugging Faceupdated 3h agoView on Hugging Face
2likes134downloads
README.md189 linesDownload Raw Back to root
1---2dataset_info:3- config_name: default4  features:5  - name: id6    dtype: int647  - name: label8    dtype: int649  - name: statement10    dtype: string11  - name: date12    dtype: string13  - name: subject14    dtype: string15  - name: speaker16    dtype: string17  - name: speaker_description18    dtype: string19  - name: state_info20    dtype: string21  - name: true_counts22    dtype: int6423  - name: mostly_true_counts24    dtype: int6425  - name: half_true_counts26    dtype: int6427  - name: mostly_false_counts28    dtype: int6429  - name: false_counts30    dtype: int6431  - name: pants_on_fire_counts32    dtype: int6433  - name: context34    dtype: string35  - name: justification36    dtype: string37  - name: binary_label38    dtype: int6439  - name: clean_statement40    dtype: string41  - name: clean_justification42    dtype: string43  - name: input_ids44    list: int3245  - name: token_type_ids46    list: int847  - name: attention_mask48    list: int849  splits:50  - name: train51    num_bytes: 4684192352    num_examples: 1836953  - name: validation54    num_bytes: 584499655    num_examples: 229756  - name: test57    num_bytes: 587266158    num_examples: 229659  download_size: 5597669260  dataset_size: 5855958061- config_name: llama3_balanced_sample62  features:63  - name: id64    dtype: int6465  - name: label66    dtype: int6467  - name: statement68    dtype: string69  - name: date70    dtype: string71  - name: subject72    dtype: string73  - name: speaker74    dtype: string75  - name: speaker_description76    dtype: string77  - name: state_info78    dtype: string79  - name: true_counts80    dtype: int6481  - name: mostly_true_counts82    dtype: int6483  - name: half_true_counts84    dtype: int6485  - name: mostly_false_counts86    dtype: int6487  - name: false_counts88    dtype: int6489  - name: pants_on_fire_counts90    dtype: int6491  - name: context92    dtype: string93  - name: justification94    dtype: string95  - name: binary_label96    dtype: int6497  - name: clean_statement98    dtype: string99  - name: clean_justification100    dtype: string101  - name: input_ids102    list: int32103  - name: token_type_ids104    list: int8105  - name: attention_mask106    list: int8107  - name: computed_label_text108    dtype: string109  - name: llama3_explanation110    dtype: string111  splits:112  - name: train113    num_bytes: 189536114    num_examples: 60115  download_size: 192607116  dataset_size: 189536117- config_name: llama3_sample118  features:119  - name: id120    dtype: int64121  - name: label122    dtype: int64123  - name: statement124    dtype: string125  - name: date126    dtype: string127  - name: subject128    dtype: string129  - name: speaker130    dtype: string131  - name: speaker_description132    dtype: string133  - name: state_info134    dtype: string135  - name: true_counts136    dtype: int64137  - name: mostly_true_counts138    dtype: int64139  - name: half_true_counts140    dtype: int64141  - name: mostly_false_counts142    dtype: int64143  - name: false_counts144    dtype: int64145  - name: pants_on_fire_counts146    dtype: int64147  - name: context148    dtype: string149  - name: justification150    dtype: string151  - name: binary_label152    dtype: int64153  - name: clean_statement154    dtype: string155  - name: clean_justification156    dtype: string157  - name: input_ids158    list: int32159  - name: token_type_ids160    list: int8161  - name: attention_mask162    list: int8163  - name: llama3_explanation164    dtype: string165  splits:166  - name: train167    num_bytes: 182568168    num_examples: 60169  download_size: 191337170  dataset_size: 182568171configs:172- config_name: default173  data_files:174  - split: train175    path: data/train-*176  - split: validation177    path: data/validation-*178  - split: test179    path: data/test-*180- config_name: llama3_balanced_sample181  data_files:182  - split: train183    path: llama3_balanced_sample/train-*184- config_name: llama3_sample185  data_files:186  - split: train187    path: llama3_sample/train-*188---189