CoolFace
Datasetpublic

giulio98/LongBench

sourceHugging Faceupdated 1y agoView on Hugging Face
0likes1.1kdownloads
README.md1145 linesDownload Raw Back to root
1---2dataset_info:3- config_name: 2wikimqa4  features:5  - name: question6    dtype: string7  - name: context8    dtype: string9  - name: answers10    sequence: string11  - name: length12    dtype: int3213  - name: dataset14    dtype: string15  - name: language16    dtype: string17  - name: all_classes18    dtype: 'null'19  - name: _id20    dtype: string21  - name: _row_id22    dtype: int6423  - name: context_cheatsheet24    dtype: string25  - name: prompt_cheatsheet26    dtype: string27  - name: answer_prefix28    dtype: string29  - name: task30    dtype: string31  - name: max_new_tokens32    dtype: int6433  - name: cheatsheet34    dtype: string35  splits:36  - name: test37    num_bytes: 1300032238    num_examples: 20039  download_size: 772202240  dataset_size: 1300032241- config_name: 2wikimqa_e42  features:43  - name: question44    dtype: string45  - name: context46    dtype: string47  - name: answers48    sequence: string49  - name: length50    dtype: int3251  - name: dataset52    dtype: string53  - name: language54    dtype: string55  - name: all_classes56    dtype: 'null'57  - name: _id58    dtype: string59  - name: answer_prefix60    dtype: string61  - name: task62    dtype: string63  - name: task_description64    dtype: string65  - name: max_new_tokens66    dtype: int6467  splits:68  - name: test69    num_bytes: 1145278270    num_examples: 30071  download_size: 681309972  dataset_size: 1145278273- config_name: gov_report74  features:75  - name: question76    dtype: string77  - name: context78    dtype: string79  - name: answers80    sequence: string81  - name: length82    dtype: int3283  - name: dataset84    dtype: string85  - name: language86    dtype: string87  - name: all_classes88    dtype: 'null'89  - name: _id90    dtype: string91  - name: _row_id92    dtype: int6493  - name: context_cheatsheet94    dtype: string95  - name: prompt_cheatsheet96    dtype: string97  - name: answer_prefix98    dtype: string99  - name: task100    dtype: string101  - name: max_new_tokens102    dtype: int64103  - name: cheatsheet104    dtype: string105  splits:106  - name: test107    num_bytes: 23339233108    num_examples: 200109  download_size: 10981661110  dataset_size: 23339233111- config_name: gov_report_e112  features:113  - name: question114    dtype: string115  - name: context116    dtype: string117  - name: answers118    sequence: string119  - name: length120    dtype: int32121  - name: dataset122    dtype: string123  - name: language124    dtype: string125  - name: all_classes126    dtype: 'null'127  - name: _id128    dtype: string129  - name: answer_prefix130    dtype: string131  - name: task132    dtype: string133  - name: task_description134    dtype: string135  - name: max_new_tokens136    dtype: int64137  splits:138  - name: test139    num_bytes: 14342598140    num_examples: 300141  download_size: 6669359142  dataset_size: 14342598143- config_name: hotpotqa144  features:145  - name: question146    dtype: string147  - name: context148    dtype: string149  - name: answers150    sequence: string151  - name: length152    dtype: int32153  - name: dataset154    dtype: string155  - name: language156    dtype: string157  - name: all_classes158    dtype: 'null'159  - name: _id160    dtype: string161  - name: _row_id162    dtype: int64163  - name: context_cheatsheet164    dtype: string165  - name: prompt_cheatsheet166    dtype: string167  - name: answer_prefix168    dtype: string169  - name: task170    dtype: string171  - name: max_new_tokens172    dtype: int64173  - name: cheatsheet174    dtype: string175  splits:176  - name: test177    num_bytes: 23875480178    num_examples: 200179  download_size: 13773255180  dataset_size: 23875480181- config_name: hotpotqa_e182  features:183  - name: question184    dtype: string185  - name: context186    dtype: string187  - name: answers188    sequence: string189  - name: length190    dtype: int32191  - name: dataset192    dtype: string193  - name: language194    dtype: string195  - name: all_classes196    dtype: 'null'197  - name: _id198    dtype: string199  - name: answer_prefix200    dtype: string201  - name: task202    dtype: string203  - name: task_description204    dtype: string205  - name: max_new_tokens206    dtype: int64207  splits:208  - name: test209    num_bytes: 12445130210    num_examples: 300211  download_size: 7202317212  dataset_size: 12445130213- config_name: lcc214  features:215  - name: question216    dtype: string217  - name: context218    dtype: string219  - name: answers220    sequence: string221  - name: length222    dtype: int32223  - name: dataset224    dtype: string225  - name: language226    dtype: string227  - name: all_classes228    dtype: 'null'229  - name: _id230    dtype: string231  - name: _row_id232    dtype: int64233  - name: context_cheatsheet234    dtype: string235  - name: prompt_cheatsheet236    dtype: string237  - name: answer_prefix238    dtype: string239  - name: task240    dtype: string241  - name: max_new_tokens242    dtype: int64243  - name: cheatsheet244    dtype: string245  splits:246  - name: test247    num_bytes: 15500741248    num_examples: 500249  download_size: 5278164250  dataset_size: 15500741251- config_name: lcc_e252  features:253  - name: question254    dtype: string255  - name: context256    dtype: string257  - name: answers258    sequence: string259  - name: length260    dtype: int32261  - name: dataset262    dtype: string263  - name: language264    dtype: string265  - name: all_classes266    dtype: 'null'267  - name: _id268    dtype: string269  - name: answer_prefix270    dtype: string271  - name: task272    dtype: string273  - name: task_description274    dtype: string275  - name: max_new_tokens276    dtype: int64277  splits:278  - name: test279    num_bytes: 17789705280    num_examples: 300281  download_size: 5523361282  dataset_size: 17789705283- config_name: multi_news284  features:285  - name: question286    dtype: string287  - name: context288    dtype: string289  - name: answers290    sequence: string291  - name: length292    dtype: int32293  - name: dataset294    dtype: string295  - name: language296    dtype: string297  - name: all_classes298    dtype: 'null'299  - name: _id300    dtype: string301  - name: _row_id302    dtype: int64303  - name: context_cheatsheet304    dtype: string305  - name: prompt_cheatsheet306    dtype: string307  - name: answer_prefix308    dtype: string309  - name: task310    dtype: string311  - name: max_new_tokens312    dtype: int64313  - name: cheatsheet314    dtype: string315  splits:316  - name: test317    num_bytes: 5741623318    num_examples: 200319  download_size: 3067554320  dataset_size: 5741623321- config_name: multi_news_e322  features:323  - name: question324    dtype: string325  - name: context326    dtype: string327  - name: answers328    sequence: string329  - name: length330    dtype: int32331  - name: dataset332    dtype: string333  - name: language334    dtype: string335  - name: all_classes336    dtype: 'null'337  - name: _id338    dtype: string339  - name: answer_prefix340    dtype: string341  - name: task342    dtype: string343  - name: task_description344    dtype: string345  - name: max_new_tokens346    dtype: int64347  splits:348  - name: test349    num_bytes: 11379222350    num_examples: 294351  download_size: 5859143352  dataset_size: 11379222353- config_name: multifieldqa_en354  features:355  - name: question356    dtype: string357  - name: context358    dtype: string359  - name: answers360    sequence: string361  - name: length362    dtype: int32363  - name: dataset364    dtype: string365  - name: language366    dtype: string367  - name: all_classes368    dtype: 'null'369  - name: _id370    dtype: string371  - name: _row_id372    dtype: int64373  - name: context_cheatsheet374    dtype: string375  - name: prompt_cheatsheet376    dtype: string377  - name: answer_prefix378    dtype: string379  - name: task380    dtype: string381  - name: max_new_tokens382    dtype: int64383  - name: cheatsheet384    dtype: string385  splits:386  - name: test387    num_bytes: 9391845388    num_examples: 150389  download_size: 3905544390  dataset_size: 9391845391- config_name: multifieldqa_en_e392  features:393  - name: question394    dtype: string395  - name: context396    dtype: string397  - name: answers398    sequence: string399  - name: length400    dtype: int32401  - name: dataset402    dtype: string403  - name: language404    dtype: string405  - name: all_classes406    dtype: 'null'407  - name: _id408    dtype: string409  - name: answer_prefix410    dtype: string411  - name: task412    dtype: string413  - name: task_description414    dtype: string415  - name: max_new_tokens416    dtype: int64417  splits:418  - name: test419    num_bytes: 4466519420    num_examples: 150421  download_size: 1834530422  dataset_size: 4466519423- config_name: musique424  features:425  - name: question426    dtype: string427  - name: context428    dtype: string429  - name: answers430    sequence: string431  - name: length432    dtype: int32433  - name: dataset434    dtype: string435  - name: language436    dtype: string437  - name: all_classes438    dtype: 'null'439  - name: _id440    dtype: string441  - name: _row_id442    dtype: int64443  - name: context_cheatsheet444    dtype: string445  - name: prompt_cheatsheet446    dtype: string447  - name: answer_prefix448    dtype: string449  - name: task450    dtype: string451  - name: max_new_tokens452    dtype: int64453  - name: cheatsheet454    dtype: string455  splits:456  - name: test457    num_bytes: 29222414458    num_examples: 200459  download_size: 16867087460  dataset_size: 29222414461- config_name: narrativeqa462  features:463  - name: question464    dtype: string465  - name: context466    dtype: string467  - name: answers468    sequence: string469  - name: length470    dtype: int32471  - name: dataset472    dtype: string473  - name: language474    dtype: string475  - name: all_classes476    dtype: 'null'477  - name: _id478    dtype: string479  - name: _row_id480    dtype: int64481  - name: context_cheatsheet482    dtype: string483  - name: prompt_cheatsheet484    dtype: string485  - name: answer_prefix486    dtype: string487  - name: task488    dtype: string489  - name: max_new_tokens490    dtype: int64491  - name: cheatsheet492    dtype: string493  splits:494  - name: test495    num_bytes: 44146824496    num_examples: 200497  download_size: 2822691498  dataset_size: 44146824499- config_name: passage_count500  features:501  - name: question502    dtype: string503  - name: context504    dtype: string505  - name: answers506    sequence: string507  - name: length508    dtype: int32509  - name: dataset510    dtype: string511  - name: language512    dtype: string513  - name: all_classes514    dtype: 'null'515  - name: _id516    dtype: string517  - name: _row_id518    dtype: int64519  - name: context_cheatsheet520    dtype: string521  - name: prompt_cheatsheet522    dtype: string523  - name: answer_prefix524    dtype: string525  - name: task526    dtype: string527  - name: max_new_tokens528    dtype: int64529  - name: cheatsheet530    dtype: string531  splits:532  - name: test533    num_bytes: 28032548534    num_examples: 200535  download_size: 10477014536  dataset_size: 28032548537- config_name: passage_count_e538  features:539  - name: question540    dtype: string541  - name: context542    dtype: string543  - name: answers544    sequence: string545  - name: length546    dtype: int32547  - name: dataset548    dtype: string549  - name: language550    dtype: string551  - name: all_classes552    dtype: 'null'553  - name: _id554    dtype: string555  - name: answer_prefix556    dtype: string557  - name: task558    dtype: string559  - name: task_description560    dtype: string561  - name: max_new_tokens562    dtype: int64563  splits:564  - name: test565    num_bytes: 11329954566    num_examples: 300567  download_size: 3924891568  dataset_size: 11329954569- config_name: passage_retrieval_en570  features:571  - name: question572    dtype: string573  - name: context574    dtype: string575  - name: answers576    sequence: string577  - name: length578    dtype: int32579  - name: dataset580    dtype: string581  - name: language582    dtype: string583  - name: all_classes584    dtype: 'null'585  - name: _id586    dtype: string587  - name: _row_id588    dtype: int64589  - name: context_cheatsheet590    dtype: string591  - name: prompt_cheatsheet592    dtype: string593  - name: answer_prefix594    dtype: string595  - name: task596    dtype: string597  - name: max_new_tokens598    dtype: int64599  - name: cheatsheet600    dtype: string601  splits:602  - name: test603    num_bytes: 23808776604    num_examples: 200605  download_size: 14784200606  dataset_size: 23808776607- config_name: passage_retrieval_en_e608  features:609  - name: question610    dtype: string611  - name: context612    dtype: string613  - name: answers614    sequence: string615  - name: length616    dtype: int32617  - name: dataset618    dtype: string619  - name: language620    dtype: string621  - name: all_classes622    dtype: 'null'623  - name: _id624    dtype: string625  - name: answer_prefix626    dtype: string627  - name: task628    dtype: string629  - name: task_description630    dtype: string631  - name: max_new_tokens632    dtype: int64633  splits:634  - name: test635    num_bytes: 11247035636    num_examples: 300637  download_size: 6972223638  dataset_size: 11247035639- config_name: qasper640  features:641  - name: question642    dtype: string643  - name: context644    dtype: string645  - name: answers646    sequence: string647  - name: length648    dtype: int32649  - name: dataset650    dtype: string651  - name: language652    dtype: string653  - name: all_classes654    dtype: 'null'655  - name: _id656    dtype: string657  - name: _row_id658    dtype: int64659  - name: context_cheatsheet660    dtype: string661  - name: prompt_cheatsheet662    dtype: string663  - name: answer_prefix664    dtype: string665  - name: task666    dtype: string667  - name: max_new_tokens668    dtype: int64669  - name: cheatsheet670    dtype: string671  splits:672  - name: test673    num_bytes: 10430705674    num_examples: 200675  download_size: 4029155676  dataset_size: 10430705677- config_name: qasper_e678  features:679  - name: question680    dtype: string681  - name: context682    dtype: string683  - name: answers684    sequence: string685  - name: length686    dtype: int32687  - name: dataset688    dtype: string689  - name: language690    dtype: string691  - name: all_classes692    dtype: 'null'693  - name: _id694    dtype: string695  - name: answer_prefix696    dtype: string697  - name: task698    dtype: string699  - name: task_description700    dtype: string701  - name: max_new_tokens702    dtype: int64703  splits:704  - name: test705    num_bytes: 7098072706    num_examples: 224707  download_size: 2030778708  dataset_size: 7098072709- config_name: qmsum710  features:711  - name: question712    dtype: string713  - name: context714    dtype: string715  - name: answers716    sequence: string717  - name: length718    dtype: int32719  - name: dataset720    dtype: string721  - name: language722    dtype: string723  - name: all_classes724    dtype: 'null'725  - name: _id726    dtype: string727  - name: _row_id728    dtype: int64729  - name: context_cheatsheet730    dtype: string731  - name: prompt_cheatsheet732    dtype: string733  - name: answer_prefix734    dtype: string735  - name: task736    dtype: string737  - name: max_new_tokens738    dtype: int64739  - name: cheatsheet740    dtype: string741  splits:742  - name: test743    num_bytes: 23873021744    num_examples: 200745  download_size: 2060714746  dataset_size: 23873021747- config_name: repobench-p748  features:749  - name: question750    dtype: string751  - name: context752    dtype: string753  - name: answers754    sequence: string755  - name: length756    dtype: int32757  - name: dataset758    dtype: string759  - name: language760    dtype: string761  - name: all_classes762    dtype: 'null'763  - name: _id764    dtype: string765  - name: _row_id766    dtype: int64767  - name: context_cheatsheet768    dtype: string769  - name: prompt_cheatsheet770    dtype: string771  - name: answer_prefix772    dtype: string773  - name: task774    dtype: string775  - name: max_new_tokens776    dtype: int64777  - name: cheatsheet778    dtype: string779  splits:780  - name: test781    num_bytes: 48584558782    num_examples: 500783  download_size: 15485729784  dataset_size: 48584558785- config_name: repobench-p_e786  features:787  - name: question788    dtype: string789  - name: context790    dtype: string791  - name: answers792    sequence: string793  - name: length794    dtype: int32795  - name: dataset796    dtype: string797  - name: language798    dtype: string799  - name: all_classes800    dtype: 'null'801  - name: _id802    dtype: string803  - name: answer_prefix804    dtype: string805  - name: task806    dtype: string807  - name: task_description808    dtype: string809  - name: max_new_tokens810    dtype: int64811  splits:812  - name: test813    num_bytes: 20414779814    num_examples: 300815  download_size: 6635836816  dataset_size: 20414779817- config_name: samsum818  features:819  - name: question820    dtype: string821  - name: context822    dtype: string823  - name: answers824    sequence: string825  - name: length826    dtype: int32827  - name: dataset828    dtype: string829  - name: language830    dtype: string831  - name: all_classes832    dtype: 'null'833  - name: _id834    dtype: string835  - name: _row_id836    dtype: int64837  - name: context_cheatsheet838    dtype: string839  - name: prompt_cheatsheet840    dtype: string841  - name: answer_prefix842    dtype: string843  - name: task844    dtype: string845  - name: max_new_tokens846    dtype: int64847  - name: cheatsheet848    dtype: string849  splits:850  - name: test851    num_bytes: 14344452852    num_examples: 200853  download_size: 8306671854  dataset_size: 14344452855- config_name: samsum_e856  features:857  - name: question858    dtype: string859  - name: context860    dtype: string861  - name: answers862    sequence: string863  - name: length864    dtype: int32865  - name: dataset866    dtype: string867  - name: language868    dtype: string869  - name: all_classes870    dtype: 'null'871  - name: _id872    dtype: string873  - name: answer_prefix874    dtype: string875  - name: task876    dtype: string877  - name: task_description878    dtype: string879  - name: max_new_tokens880    dtype: int64881  splits:882  - name: test883    num_bytes: 10364277884    num_examples: 300885  download_size: 6088405886  dataset_size: 10364277887- config_name: trec888  features:889  - name: question890    dtype: string891  - name: context892    dtype: string893  - name: answers894    sequence: string895  - name: length896    dtype: int32897  - name: dataset898    dtype: string899  - name: language900    dtype: string901  - name: all_classes902    sequence: string903  - name: _id904    dtype: string905  - name: _row_id906    dtype: int64907  - name: context_cheatsheet908    dtype: string909  - name: prompt_cheatsheet910    dtype: string911  - name: answer_prefix912    dtype: string913  - name: task914    dtype: string915  - name: max_new_tokens916    dtype: int64917  - name: cheatsheet918    dtype: string919  splits:920  - name: test921    num_bytes: 13847523922    num_examples: 200923  download_size: 5785070924  dataset_size: 13847523925- config_name: trec_e926  features:927  - name: question928    dtype: string929  - name: context930    dtype: string931  - name: answers932    sequence: string933  - name: length934    dtype: int32935  - name: dataset936    dtype: string937  - name: language938    dtype: string939  - name: all_classes940    sequence: string941  - name: _id942    dtype: string943  - name: answer_prefix944    dtype: string945  - name: task946    dtype: string947  - name: task_description948    dtype: string949  - name: max_new_tokens950    dtype: int64951  splits:952  - name: test953    num_bytes: 11242848954    num_examples: 300955  download_size: 4851146956  dataset_size: 11242848957- config_name: triviaqa958  features:959  - name: question960    dtype: string961  - name: context962    dtype: string963  - name: answers964    sequence: string965  - name: length966    dtype: int32967  - name: dataset968    dtype: string969  - name: language970    dtype: string971  - name: all_classes972    dtype: 'null'973  - name: _id974    dtype: string975  - name: _row_id976    dtype: int64977  - name: context_cheatsheet978    dtype: string979  - name: prompt_cheatsheet980    dtype: string981  - name: answer_prefix982    dtype: string983  - name: task984    dtype: string985  - name: max_new_tokens986    dtype: int64987  - name: cheatsheet988    dtype: string989  splits:990  - name: test991    num_bytes: 20235861992    num_examples: 200993  download_size: 12389980994  dataset_size: 20235861995- config_name: triviaqa_e996  features:997  - name: question998    dtype: string999  - name: context1000    dtype: string1001  - name: answers1002    sequence: string1003  - name: length1004    dtype: int321005  - name: dataset1006    dtype: string1007  - name: language1008    dtype: string1009  - name: all_classes1010    dtype: 'null'1011  - name: _id1012    dtype: string1013  - name: answer_prefix1014    dtype: string1015  - name: task1016    dtype: string1017  - name: task_description1018    dtype: string1019  - name: max_new_tokens1020    dtype: int641021  splits:1022  - name: test1023    num_bytes: 125584571024    num_examples: 3001025  download_size: 75486401026  dataset_size: 125584571027configs:1028- config_name: 2wikimqa1029  data_files:1030  - split: test1031    path: 2wikimqa/test-*1032- config_name: 2wikimqa_e1033  data_files:1034  - split: test1035    path: 2wikimqa_e/test-*1036- config_name: gov_report1037  data_files:1038  - split: test1039    path: gov_report/test-*1040- config_name: gov_report_e1041  data_files:1042  - split: test1043    path: gov_report_e/test-*1044- config_name: hotpotqa1045  data_files:1046  - split: test1047    path: hotpotqa/test-*1048- config_name: hotpotqa_e1049  data_files:1050  - split: test1051    path: hotpotqa_e/test-*1052- config_name: lcc1053  data_files:1054  - split: test1055    path: lcc/test-*1056- config_name: lcc_e1057  data_files:1058  - split: test1059    path: lcc_e/test-*1060- config_name: multi_news1061  data_files:1062  - split: test1063    path: multi_news/test-*1064- config_name: multi_news_e1065  data_files:1066  - split: test1067    path: multi_news_e/test-*1068- config_name: multifieldqa_en1069  data_files:1070  - split: test1071    path: multifieldqa_en/test-*1072- config_name: multifieldqa_en_e1073  data_files:1074  - split: test1075    path: multifieldqa_en_e/test-*1076- config_name: musique1077  data_files:1078  - split: test1079    path: musique/test-*1080- config_name: narrativeqa1081  data_files:1082  - split: test1083    path: narrativeqa/test-*1084- config_name: passage_count1085  data_files:1086  - split: test1087    path: passage_count/test-*1088- config_name: passage_count_e1089  data_files:1090  - split: test1091    path: passage_count_e/test-*1092- config_name: passage_retrieval_en1093  data_files:1094  - split: test1095    path: passage_retrieval_en/test-*1096- config_name: passage_retrieval_en_e1097  data_files:1098  - split: test1099    path: passage_retrieval_en_e/test-*1100- config_name: qasper1101  data_files:1102  - split: test1103    path: qasper/test-*1104- config_name: qasper_e1105  data_files:1106  - split: test1107    path: qasper_e/test-*1108- config_name: qmsum1109  data_files:1110  - split: test1111    path: qmsum/test-*1112- config_name: repobench-p1113  data_files:1114  - split: test1115    path: repobench-p/test-*1116- config_name: repobench-p_e1117  data_files:1118  - split: test1119    path: repobench-p_e/test-*1120- config_name: samsum1121  data_files:1122  - split: test1123    path: samsum/test-*1124- config_name: samsum_e1125  data_files:1126  - split: test1127    path: samsum_e/test-*1128- config_name: trec1129  data_files:1130  - split: test1131    path: trec/test-*1132- config_name: trec_e1133  data_files:1134  - split: test1135    path: trec_e/test-*1136- config_name: triviaqa1137  data_files:1138  - split: test1139    path: triviaqa/test-*1140- config_name: triviaqa_e1141  data_files:1142  - split: test1143    path: triviaqa_e/test-*1144---1145