{"meta":{"page":1,"per_page":50,"max_per_page":100,"total":34,"total_is_capped":false,"direct_labels_cover":0,"predictions_cover":34,"direct_label_status":"direct model label, unvalidated","prediction_status":"machine_predicted_unvalidated (Codex and Gemma teacher distillation)","score_status":"score_only:v0-immature-baseline (scores rank; they never assert a category)","snapshot":{"source":"OpenAlex, pinned release, all 482 partitions","release":"2026-06-24","frame_built":"2026-07-12","author_layer_release":"2026-06-26"},"query_hash":"156f0c5b8dc6","filters":{"venue":"International Conference on Computational Linguistics"}},"results":[{"id":"W2914220664","doi":"","title":"Deep Models for Arabic Dialect Identification on Benchmarked Data","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":61,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of British Columbia","funders":"","keywords":"Benchmark (surveying); Computer science; Artificial intelligence; Task (project management); Deep learning; Arabic; Natural language processing; Binary classification; Deep neural networks; Identification (biology); Machine learning; Artificial neural network; Recurrent neural network; Test data; Binary number; Speech recognition; Support vector machine; Linguistics; Mathematics; Geography","authors":[{"name":"Mohamed Elaraby","is_ca":false},{"name":"Muhammad Abdul-Mageed","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1011538868834525,"gpt":0.3765717506027052,"spread":0.2754178637192528,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003822666,0.003153874,0.001154637,0.002856643,0.001163006,0.001773805,0.00283606,0.002279087,0.006617148],"category_scores_gemma":[0.01373547,0.000517507,0.001343481,0.0026688,0.001068354,0.003046163,0.00257073,0.003556804,0.005859227],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00225225,"about_ca_system_score_gemma":0.001401767,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.03114979,"about_ca_topic_score_gemma":0.04051483,"domain_scores_codex":[0.9975515,0.000904968,0.0002047483,0.0006413599,0.0004413933,0.0002559908],"domain_scores_gemma":[0.993727,0.002597362,0.0002861056,0.001557822,0.001538584,0.0002931368],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002723394,0.002574268,0.01938711,0.00170758,0.0008150364,0.0008286241,0.0006264319,0.2374206,0.01110296,0.00365418,0.2302635,0.4888963],"study_design_scores_gemma":[0.0004096191,0.0005499137,0.01178978,0.0002336095,0.0001603937,0.0003246539,0.0008291264,0.9316401,0.01672951,0.008715757,0.02848861,0.0001288981],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.7793805,0.01180767,0.0618514,0.003383128,0.002188389,0.0006476797,0.09620999,0.02442597,0.02010529],"genre_scores_gemma":[0.6701142,0.001339031,0.08052433,0.0009237159,0.000316679,0.0004937258,0.2342135,0.0006441923,0.01143065],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.03114979,"threshold_uncertainty_score":0.06193691,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W1956103381","doi":"","title":"Automatic Acquisition of Lexical Formality","year":2010,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":56,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Toronto","funders":"","keywords":"Formality; Computer science; Word (group theory); Natural language processing; Artificial intelligence; Word Association; Metric (unit); Similarity (geometry); Association (psychology); Task (project management); Synonym (taxonomy); Linguistics; Speech recognition; Psychology","authors":[{"name":"Julian Brooke","is_ca":true},{"name":"Tong Wang","is_ca":true},{"name":"Graeme Hirst","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0284122819659165,"gpt":0.3393629900576921,"spread":0.3109507080917756,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002189051,0.0009597391,0.001273933,0.007282081,0.0008027698,0.002958264,0.001291727,0.0006121761,0.006069791],"category_scores_gemma":[0.01571402,0.0008203471,0.0007040604,0.002430234,0.0007788594,0.005937049,0.00400361,0.001451868,0.0038987],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0006262375,"about_ca_system_score_gemma":0.001479681,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001192078,"about_ca_topic_score_gemma":0.00214164,"domain_scores_codex":[0.9976675,0.000461404,0.0003026445,0.0008557551,0.0005591194,0.0001535165],"domain_scores_gemma":[0.9901086,0.004382833,0.00105876,0.001461749,0.00267346,0.0003145325],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0002940958,0.0001437604,0.01841269,0.0008904217,0.00010957,0.0005505041,0.001579623,0.001472758,0.1150108,0.0114391,0.009114193,0.8409824],"study_design_scores_gemma":[0.0003437394,0.0007875534,0.1159348,0.0008326119,0.0004863855,0.006344464,0.005297728,0.4188152,0.1892527,0.1408122,0.1205211,0.0005716207],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3154962,0.001825341,0.6265804,0.0007653699,0.0004007737,0.0005765731,0.0050962,0.03236469,0.0168944],"genre_scores_gemma":[0.6633444,0.0005176776,0.3223411,0.0001528582,0.0001473376,0.0002721966,0.008415918,0.001296974,0.003511537],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.007282081,"threshold_uncertainty_score":0.02030545,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2865911429","doi":"","title":"Abstractive Unsupervised Multi-Document Summarization using Paraphrastic Sentence Fusion","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":53,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Lethbridge","funders":"","keywords":"Automatic summarization; Computer science; Sentence; Natural language processing; Artificial intelligence; Set (abstract data type); Machine translation; Word embedding; Multi-document summarization; Word (group theory); Information retrieval; Embedding; Linguistics","authors":[{"name":"Mir Tafseer Nayeem","is_ca":true},{"name":"Tanvir Ahmed Fuad","is_ca":true},{"name":"Yllias Chali","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.09255504555690312,"gpt":0.3506786421700804,"spread":0.2581235966131773,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00128472,0.001469428,0.001367177,0.001983395,0.0004585274,0.001104998,0.001323101,0.0009671329,0.002291838],"category_scores_gemma":[0.003916661,0.0003553146,0.001241448,0.001640879,0.0003558526,0.002375837,0.001027635,0.001295969,0.002481435],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004535647,"about_ca_system_score_gemma":0.0006494264,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001060087,"about_ca_topic_score_gemma":0.001823186,"domain_scores_codex":[0.9988194,0.0003294304,0.0001243749,0.0003543253,0.0003130929,0.00005918457],"domain_scores_gemma":[0.9974045,0.0008127111,0.0003897112,0.0005116814,0.0008006206,0.00008063941],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0003737762,0.000330324,0.001094889,0.0008228497,0.0002926695,0.0003365183,0.0005130786,0.02888026,0.1512977,0.004211293,0.01209559,0.799751],"study_design_scores_gemma":[0.0001010109,0.0009214451,0.003513892,0.00007960608,0.0004468141,0.0007933794,0.0003316175,0.7692926,0.1862004,0.01231637,0.02587905,0.0001237119],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01835506,0.0009101945,0.9712451,0.0002212879,0.0001117891,0.000166832,0.0006225855,0.007247339,0.001119849],"genre_scores_gemma":[0.1897727,0.0007245846,0.7970082,0.0003431475,0.000349434,0.0003072501,0.006074378,0.0005512945,0.004869003],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.002291838,"threshold_uncertainty_score":0.007666945,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2814750830","doi":"","title":"Learning Emotion-enriched Word Representations","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Sentiment Analysis and Opinion Mining","field":"Computer Science","cited_by":51,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"York University","funders":"","keywords":"Similarity (geometry); Word (group theory); Natural language processing; Computer science; Meaning (existential); Affect (linguistics); Representation (politics); Artificial intelligence; Emotion classification; Contrast (vision); Psychology; Cognitive psychology; Linguistics; Communication","authors":[{"name":"Ameeta Agrawal","is_ca":true},{"name":"Aijun An","is_ca":true},{"name":"Manos Papagelis","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.05774088466866577,"gpt":0.3536746629688343,"spread":0.2959337783001685,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0003559147,0.001087455,0.0004566773,0.0009804579,0.0002021738,0.0007337651,0.0008016798,0.0008409607,0.002251875],"category_scores_gemma":[0.002343219,0.0001913078,0.0006801471,0.0009462062,0.0002773841,0.00197978,0.000851773,0.001110612,0.001029467],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004050607,"about_ca_system_score_gemma":0.0003287043,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0007891998,"about_ca_topic_score_gemma":0.001271988,"domain_scores_codex":[0.9996631,0.00007301661,0.00002543163,0.000149539,0.00004599991,0.00004300805],"domain_scores_gemma":[0.9995515,0.0001704393,0.0000561393,0.0000774562,0.0001232963,0.00002111691],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004972965,0.0004099227,0.003740117,0.0003405162,0.0001845087,0.0002501267,0.0004701553,0.05683611,0.05256256,0.01055502,0.01361941,0.8605343],"study_design_scores_gemma":[0.00005847109,0.000204254,0.002228771,0.00004566032,0.0001229497,0.000142651,0.0002250731,0.945376,0.01427963,0.03312405,0.004159208,0.00003319436],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.1988266,0.001042444,0.7910687,0.0005123278,0.0003239597,0.0001537268,0.00141587,0.002979294,0.003677069],"genre_scores_gemma":[0.7955788,0.0007608095,0.1935095,0.0002394317,0.000192072,0.0002780951,0.004787473,0.0001525785,0.004501216],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.002251875,"threshold_uncertainty_score":0.007533252,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2108711363","doi":"","title":"A Strategy of Mapping Polish WordNet onto Princeton WordNet","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":40,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Ottawa","funders":"","keywords":"WordNet; Premise; Computer science; Set (abstract data type); Natural language processing; Focus (optics); Artificial intelligence; Lexical database; Range (aeronautics); Information retrieval; Linguistics; Programming language; Philosophy","authors":[{"name":"Ewa Rudnicka","is_ca":false},{"name":"Marek Maziarz","is_ca":false},{"name":"Maciej Piasecki","is_ca":false},{"name":"Stan Śzpakowicz","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.06179235534177949,"gpt":0.3410952297029126,"spread":0.2793028743611331,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00155462,0.0007445642,0.000419479,0.003804877,0.001896846,0.002192403,0.0009552566,0.0005701733,0.007832087],"category_scores_gemma":[0.006650538,0.000815952,0.0007725426,0.002735085,0.001270356,0.005428575,0.004388485,0.001652013,0.003417654],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007372301,"about_ca_system_score_gemma":0.001706074,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004162129,"about_ca_topic_score_gemma":0.006140627,"domain_scores_codex":[0.9989713,0.0003190777,0.000104873,0.000315042,0.0002188019,0.00007083003],"domain_scores_gemma":[0.9986494,0.0002553635,0.00006167174,0.0006309948,0.0003364067,0.00006613094],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0001356405,0.0001806922,0.002893368,0.0002623186,0.00009163954,0.0004149135,0.002980398,0.004939711,0.01456842,0.5718862,0.01295955,0.3886872],"study_design_scores_gemma":[0.00006288382,0.0002736166,0.003636281,0.0001799987,0.000112112,0.0008693518,0.002502028,0.06601118,0.0472357,0.5755864,0.3033997,0.0001307208],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.02082188,0.00008530088,0.9531497,0.0007366914,0.0001743856,0.0005412287,0.0008688683,0.002520006,0.02110184],"genre_scores_gemma":[0.15013,0.0002834449,0.8307717,0.0003863994,0.00004632485,0.00156565,0.002044608,0.0009836874,0.01378834],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.007832087,"threshold_uncertainty_score":0.02620089,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2250555305","doi":"","title":"On Panini and the Generative Capacity of Contextualized Replacement Systems","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"semigroups and automata theory","field":"Computer Science","cited_by":37,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Toronto","funders":"","keywords":"Generative grammar; Formalism (music); Sanskrit; Computer science; Grammar; Rewriting; Linguistics; Mildly context-sensitive grammar formalism; Adaptive grammar; Emergent grammar; Natural language processing; Programming language; Mathematics; Artificial intelligence; Philosophy; Literature","authors":[{"name":"Gerald Penn","is_ca":true},{"name":"Paul Kiparsky","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0576393092193052,"gpt":0.3040569194395782,"spread":0.246417610220273,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003329965,0.0004363802,0.0005962364,0.001246851,0.001790991,0.002905121,0.000925817,0.001089754,0.006851135],"category_scores_gemma":[0.008463288,0.0005263391,0.0007923295,0.001102692,0.01146732,0.007678122,0.003455856,0.002418465,0.0006540869],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001813336,"about_ca_system_score_gemma":0.0007668199,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002062083,"about_ca_topic_score_gemma":0.001765558,"domain_scores_codex":[0.9981557,0.0009180547,0.00009165303,0.000372549,0.0002672364,0.0001948589],"domain_scores_gemma":[0.9933358,0.005063246,0.0002420302,0.0009512084,0.0002644162,0.0001432254],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.000006116129,0.000002262595,0.00009789485,0.00001273667,0.000001497403,0.00003437573,0.0006199578,0.0009440727,0.0001159211,0.9952149,0.0001450193,0.002805205],"study_design_scores_gemma":[0.000004595683,0.000007798799,0.00007862467,0.00001751004,0.000003297071,0.00005562307,0.0001171713,0.003432595,0.0002509713,0.9885653,0.007458364,0.000008149916],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1810188,0.009895286,0.418433,0.01186073,0.000305062,0.000100055,0.0002888181,0.001163764,0.3769344],"genre_scores_gemma":[0.957053,0.001627039,0.03091555,0.0005105642,0.000221708,0.00006525363,0.0001187373,0.0002462782,0.009241823],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.006851135,"threshold_uncertainty_score":0.02291936,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2574986170","doi":"","title":"plWordNet 3.0 - a Comprehensive Lexical-Semantic Resource.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":32,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Ottawa","funders":"","keywords":"WordNet; Computer science; Set (abstract data type); Natural language processing; Resource (disambiguation); Word (group theory); Artificial intelligence; Lexical database; Information retrieval; Word list; Linguistics; Class (philosophy); Programming language","authors":[{"name":"Marek Maziarz","is_ca":false},{"name":"Maciej Piasecki","is_ca":false},{"name":"Ewa Rudnicka","is_ca":false},{"name":"Stan Śzpakowicz","is_ca":true},{"name":"Paweł Kędzia","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.04651737362330788,"gpt":0.3294580097768378,"spread":0.2829406361535299,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001874589,0.001950221,0.001068054,0.006645518,0.0014938,0.003640317,0.001695172,0.001408508,0.03720091],"category_scores_gemma":[0.007872437,0.00133755,0.0008754915,0.006012694,0.0007313044,0.01328407,0.004859651,0.002190015,0.05351034],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008579882,"about_ca_system_score_gemma":0.003199798,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005015084,"about_ca_topic_score_gemma":0.006344342,"domain_scores_codex":[0.9986473,0.0002420782,0.0002689118,0.0003519871,0.0003898349,0.00009986013],"domain_scores_gemma":[0.9979882,0.0004550447,0.0001628143,0.0004296777,0.0007878982,0.0001762025],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0003858987,0.0001096331,0.001858166,0.003226863,0.0001379163,0.0006397309,0.001483224,0.0018609,0.01103819,0.07049628,0.6297829,0.2789803],"study_design_scores_gemma":[0.00002996917,0.00002540564,0.0007516993,0.0002950037,0.00003743861,0.0003481063,0.0002850393,0.00218752,0.004871569,0.02260832,0.968511,0.00004888283],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"dataset","genre_gemma":"software","genre_scores_codex":[0.01812159,0.004323892,0.3729278,0.003194607,0.001555383,0.001356673,0.395836,0.09180233,0.1108818],"genre_scores_gemma":[0.04027784,0.003144956,0.1970324,0.001238861,0.0002634539,0.002117562,0.6900859,0.02334528,0.04249373],"genre_candidate":"software","genre_consensus":null,"teacher_disagreement_score":0.03720091,"threshold_uncertainty_score":0.1244494,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W1815301076","doi":"","title":"Measuring the Non-compositionality of Multiword Expressions","year":2010,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":30,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Waterloo","funders":"","keywords":"Principle of compositionality; Computer science; Artificial intelligence; Natural language processing; Semantics (computer science); Metric (unit); Natural language; Question answering; Combinatory categorial grammar; Distributional semantics; Expression (computer science); Information extraction; Programming language; Semantic similarity; Link grammar; Rule-based machine translation","authors":[{"name":"Fan Bu","is_ca":false},{"name":"Xiaoyan Zhu","is_ca":false},{"name":"Ming Li","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0730969691908334,"gpt":0.3251466138805431,"spread":0.2520496446897096,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002456224,0.0007449691,0.0007226532,0.00282732,0.0008347637,0.001822397,0.0008609626,0.001127594,0.001325929],"category_scores_gemma":[0.02045633,0.0004179091,0.0006259092,0.001913542,0.0008864859,0.00481249,0.002667968,0.001033842,0.001017556],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0003802336,"about_ca_system_score_gemma":0.0005957267,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0004823039,"about_ca_topic_score_gemma":0.0008385936,"domain_scores_codex":[0.9962781,0.001206149,0.0005335329,0.0009544573,0.00087545,0.0001523398],"domain_scores_gemma":[0.9863704,0.008462582,0.001514085,0.001483108,0.001833315,0.0003364497],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0009589975,0.000299506,0.05271865,0.001413574,0.0003929724,0.0006661928,0.002299741,0.01237922,0.2991968,0.01681725,0.001576804,0.6112803],"study_design_scores_gemma":[0.0000940984,0.0012326,0.1083319,0.0002342626,0.000490425,0.003693757,0.003380146,0.5285692,0.2497405,0.07728264,0.02670508,0.0002453724],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.4840883,0.001593634,0.5077538,0.0002474152,0.0001570883,0.0001796649,0.0005141906,0.001157413,0.004308596],"genre_scores_gemma":[0.8126592,0.0005931136,0.1826106,0.00008799208,0.00009338513,0.0001696692,0.001826623,0.0003244131,0.001634966],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.00282732,"threshold_uncertainty_score":0.01298988,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2250571015","doi":"","title":"Towards Automatic Topical Question Generation","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":29,"is_retracted":false,"has_abstract":false,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Lethbridge","funders":"","keywords":"Computer science","authors":[{"name":"Yllias Chali","is_ca":true},{"name":"Sadid A. Hasan","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.09053207617422618,"gpt":0.3518676068766896,"spread":0.2613355307024634,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003492356,0.001250851,0.001346172,0.003915186,0.00154618,0.003569888,0.001711474,0.002230262,0.01708195],"category_scores_gemma":[0.010109,0.0009436179,0.001643951,0.001714645,0.0007190648,0.0047794,0.003882206,0.003253816,0.01283138],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009349969,"about_ca_system_score_gemma":0.001736217,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001472294,"about_ca_topic_score_gemma":0.002123127,"domain_scores_codex":[0.9967039,0.001495713,0.0002412061,0.0007631006,0.0005378001,0.0002583547],"domain_scores_gemma":[0.9938467,0.002906509,0.0001694842,0.0008881327,0.001960933,0.0002282439],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004678308,0.000340961,0.002082245,0.000614998,0.0001145296,0.0003783944,0.001080779,0.006785591,0.0718988,0.03615404,0.0807794,0.7993025],"study_design_scores_gemma":[0.0001486885,0.0001733665,0.001461661,0.0001218531,0.0001773074,0.0005316283,0.000894799,0.7415434,0.07270249,0.09500056,0.08717552,0.00006882106],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0231312,0.0007788274,0.9401515,0.001647788,0.000737522,0.0004541569,0.002308735,0.02319249,0.007597761],"genre_scores_gemma":[0.2196952,0.0004615805,0.7528078,0.000531251,0.0005123971,0.0005293142,0.0137951,0.001923562,0.009743728],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01708195,"threshold_uncertainty_score":0.05714482,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2816262648","doi":"","title":"Farewell Freebase: Migrating the SimpleQuestions Dataset to DBpedia.","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":23,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Waterloo","funders":"","keywords":"Computer science; Leverage (statistics); Question answering; Knowledge graph; Information retrieval; Benchmark (surveying); Entity linking; Task (project management); Simple (philosophy); Graph; World Wide Web; Knowledge base; Theoretical computer science; Artificial intelligence","authors":[{"name":"Michael Azmy","is_ca":false},{"name":"Peng Shi","is_ca":true},{"name":"Jimmy Lin","is_ca":true},{"name":"Ihab F. Ilyas","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.07890198444019045,"gpt":0.3637709757895359,"spread":0.2848689913493455,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002479936,0.002008047,0.0009878329,0.008066264,0.002313728,0.003215827,0.003679114,0.002622922,0.009437288],"category_scores_gemma":[0.01734786,0.0007483185,0.001494237,0.007305928,0.0009610266,0.005626526,0.003969439,0.002974337,0.008098693],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002605333,"about_ca_system_score_gemma":0.003213501,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.06907614,"about_ca_topic_score_gemma":0.09977357,"domain_scores_codex":[0.9965287,0.0007614117,0.000437979,0.001117057,0.0009376167,0.0002172429],"domain_scores_gemma":[0.9926026,0.002356673,0.0005049882,0.002167569,0.001708113,0.0006599839],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0003423569,0.0003884347,0.005121168,0.003121443,0.0002842534,0.000446297,0.0007706146,0.003901259,0.002944089,0.007295785,0.9403528,0.03503164],"study_design_scores_gemma":[0.0003500033,0.0001189201,0.01237073,0.00047614,0.0001247169,0.0005609918,0.001384819,0.02117838,0.007407871,0.01421641,0.9416279,0.0001831072],"study_design_candidate":"not_applicable","study_design_consensus":"not_applicable","genre_codex":"dataset","genre_gemma":"software","genre_scores_codex":[0.01740193,0.001695172,0.009498534,0.001175761,0.000477693,0.0004464388,0.9390919,0.01792284,0.01228977],"genre_scores_gemma":[0.01414907,0.0002806034,0.01578847,0.0003526076,0.0000341948,0.0002738733,0.9673031,0.0005717762,0.001246279],"genre_candidate":"software","genre_consensus":null,"teacher_disagreement_score":0.06907614,"threshold_uncertainty_score":0.1373481,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W169800830","doi":"","title":"Scaling up Analogical Learning","year":2008,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":22,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Scalability; Artificial intelligence; Simple (philosophy); Scaling; Identification (biology); Space (punctuation); Limit (mathematics); Machine learning; Theoretical computer science; Mathematics; Epistemology","authors":[{"name":"Philippe Langlais","is_ca":true},{"name":"François Yvon","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.06614344725555805,"gpt":0.337967393007861,"spread":0.271823945752303,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003200508,0.001591696,0.002058313,0.001707899,0.0008928993,0.002744257,0.004081717,0.002164096,0.02898109],"category_scores_gemma":[0.04074112,0.000781142,0.001084227,0.003074301,0.00166345,0.01235488,0.005032056,0.00309815,0.005724034],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001521036,"about_ca_system_score_gemma":0.00149213,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.0025916,"about_ca_topic_score_gemma":0.001955349,"domain_scores_codex":[0.9961792,0.001083134,0.0002672124,0.0009207384,0.001289187,0.0002605345],"domain_scores_gemma":[0.9659475,0.02124712,0.0006953474,0.008539135,0.002927282,0.0006436312],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.000652251,0.001137088,0.002763499,0.000910866,0.0002147877,0.0002085434,0.000457696,0.1114742,0.01589577,0.0581622,0.02137075,0.7867523],"study_design_scores_gemma":[0.0002721092,0.0003281154,0.0008404917,0.00006804761,0.000104513,0.0002538449,0.0003249384,0.7683631,0.00873557,0.2063669,0.01429818,0.00004411588],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.2751903,0.009223788,0.6361365,0.007733268,0.00107757,0.0007592945,0.001113222,0.015299,0.05346701],"genre_scores_gemma":[0.7238753,0.002400458,0.2629527,0.001203012,0.0004132334,0.000544758,0.001105017,0.0009057173,0.006599781],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.02898109,"threshold_uncertainty_score":0.09695143,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2759162401","doi":"","title":"Named Entity Recognition and Hashtag Decomposition to Improve the Classification of Tweets","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":20,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université du Québec; Université du Québec à Montréal","funders":"","keywords":"WordNet; Computer science; Artificial intelligence; Natural language processing; Preprocessor; Named-entity recognition; Segmentation; Field (mathematics); Task (project management); Information retrieval; Semantics (computer science)","authors":[{"name":"Billal Belainine","is_ca":true},{"name":"Alexsandro Fonseca","is_ca":false},{"name":"Fatiha Sadat","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0829395137147931,"gpt":0.3343411440966311,"spread":0.251401630381838,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001644848,0.001080936,0.0007878874,0.005513307,0.0007404467,0.001717609,0.0008847698,0.001057576,0.002555892],"category_scores_gemma":[0.004966146,0.0002776225,0.001068448,0.004363397,0.0003321591,0.004296219,0.0009368116,0.001062694,0.004900513],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005720077,"about_ca_system_score_gemma":0.0007847367,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003634545,"about_ca_topic_score_gemma":0.005031471,"domain_scores_codex":[0.9986726,0.0004027334,0.0001734953,0.0003228578,0.0002769213,0.0001515159],"domain_scores_gemma":[0.996586,0.001370486,0.0003719011,0.0005233038,0.001017746,0.0001306011],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.001135274,0.0007552558,0.02290549,0.0005453112,0.0003328935,0.0004969871,0.0008355009,0.01281829,0.07755624,0.00754678,0.02439554,0.8506764],"study_design_scores_gemma":[0.00007785489,0.0003662212,0.03187007,0.00009507802,0.0003697956,0.0006550906,0.001066421,0.7912063,0.1152668,0.01787667,0.04099265,0.0001571258],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.2415081,0.002172107,0.7153081,0.001322585,0.0007874202,0.0006183474,0.005462853,0.02413374,0.008686766],"genre_scores_gemma":[0.5241119,0.0007530591,0.4504972,0.0003234281,0.0003081037,0.0002749267,0.01544021,0.0004746833,0.007816541],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.005513307,"threshold_uncertainty_score":0.008698881,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2875408189","doi":"","title":"The APVA-TURBO Approach To Question Answering in Knowledge Base","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":19,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Ottawa","funders":"","keywords":"Question answering; Computer science; Correctness; Bottleneck; Knowledge base; Object (grammar); Artificial intelligence; Base (topology); Subject (documents); Information retrieval; Theoretical computer science; Machine learning; Programming language; World Wide Web","authors":[{"name":"Yue Wang","is_ca":false},{"name":"Richong Zhang","is_ca":false},{"name":"Xu Cheng","is_ca":false},{"name":"Yongyi Mao","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.06531592434647242,"gpt":0.3406081913215932,"spread":0.2752922669751208,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006162493,0.001026083,0.001369224,0.002046456,0.001024182,0.002548789,0.00547378,0.002737355,0.004450974],"category_scores_gemma":[0.0222159,0.001180696,0.001860925,0.002029258,0.002385426,0.007513308,0.004650464,0.004600807,0.001947391],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001464766,"about_ca_system_score_gemma":0.002045866,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.006590097,"about_ca_topic_score_gemma":0.007958282,"domain_scores_codex":[0.9952824,0.002525405,0.0001910215,0.00091667,0.0008286614,0.0002558912],"domain_scores_gemma":[0.9850085,0.009581022,0.0002882598,0.003463427,0.001357321,0.0003014758],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004778168,0.0002746218,0.002046547,0.000650591,0.0002370787,0.0002539613,0.001082996,0.2193874,0.007949492,0.1011293,0.01163917,0.654871],"study_design_scores_gemma":[0.00001197673,0.00004874466,0.0001704872,0.00002094222,0.00002477346,0.00008155451,0.00005014141,0.915821,0.002159396,0.07884306,0.002754541,0.00001346589],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.004315417,0.0002598247,0.9919853,0.0003740421,0.00002803161,0.00005862851,0.00008016502,0.001935428,0.0009630673],"genre_scores_gemma":[0.2406684,0.0004336787,0.7527431,0.0005820728,0.0001699146,0.0003528126,0.0008047688,0.0004632255,0.003782051],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.006590097,"threshold_uncertainty_score":0.03259075,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2962696263","doi":"","title":"Extracting Parallel Sentences with Bidirectional Recurrent Neural Networks to Improve Machine Translation","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":18,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université de Montréal","funders":"","keywords":"Computer science; Machine translation; Sentence; Artificial intelligence; Task (project management); Natural language processing; Translation (biology); Feature engineering; Recurrent neural network; Parallel corpora; Baseline (sea); Artificial neural network; Feature (linguistics); Feature extraction; Speech recognition; Deep learning","authors":[{"name":"Francis Grégoire","is_ca":false},{"name":"Philippe Langlais","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.03912998338648059,"gpt":0.3326886547462771,"spread":0.2935586713597965,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001058967,0.001651277,0.001080006,0.001273979,0.0006713939,0.0008562485,0.0009095888,0.0008769784,0.003375013],"category_scores_gemma":[0.004423319,0.0005207679,0.001009605,0.001644602,0.0003777567,0.002301431,0.001061894,0.00128753,0.003191439],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004087793,"about_ca_system_score_gemma":0.001091588,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002475366,"about_ca_topic_score_gemma":0.006049282,"domain_scores_codex":[0.9992028,0.000250522,0.0000871623,0.0001994176,0.00018923,0.00007087732],"domain_scores_gemma":[0.998403,0.0005940583,0.0001621127,0.0002409803,0.0005570083,0.00004281625],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0005703729,0.000479305,0.002073808,0.0007271477,0.0003263101,0.001100267,0.0005443455,0.06104802,0.1405773,0.007311102,0.01963778,0.7656043],"study_design_scores_gemma":[0.00007344663,0.0003390429,0.001361759,0.00004061663,0.0002555379,0.0003632751,0.0001688803,0.9005783,0.07185357,0.01399085,0.01091485,0.00005992521],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.09780084,0.001507756,0.8837446,0.0005762064,0.0004348702,0.0002187674,0.0009230502,0.009709714,0.005084129],"genre_scores_gemma":[0.3819616,0.0008031826,0.6018493,0.0004756417,0.0003404006,0.0003577566,0.006096927,0.0009261385,0.007189017],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.003375013,"threshold_uncertainty_score":0.01129055,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2848618769","doi":"","title":"Authorship Identification for Literary Book Recommendations","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Authorship Attribution and Profiling","field":"Computer Science","cited_by":13,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Ottawa","funders":"","keywords":"Reading (process); Recommender system; Pleasure; Computer science; Identification (biology); Style (visual arts); Factor (programming language); Writing style; Qualitative analysis; Information retrieval; World Wide Web; Natural language processing; Artificial intelligence; Qualitative research; Psychology; Linguistics; Literature; Art; Sociology","authors":[{"name":"Haifa Alharthi","is_ca":true},{"name":"Diana Inkpen","is_ca":true},{"name":"Stan Śzpakowicz","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1127695465733507,"gpt":0.3943209370164236,"spread":0.2815513904430729,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002096687,0.00096813,0.0009009565,0.005797121,0.001433809,0.002134788,0.001305084,0.001524805,0.006038892],"category_scores_gemma":[0.015784,0.0005235341,0.0009393167,0.003172705,0.0002998029,0.003510542,0.001121757,0.001401497,0.007183637],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008642376,"about_ca_system_score_gemma":0.001289081,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005264567,"about_ca_topic_score_gemma":0.01408992,"domain_scores_codex":[0.9978028,0.0005157184,0.0002075945,0.0007417834,0.0005605106,0.0001716538],"domain_scores_gemma":[0.9917223,0.0033831,0.0007185491,0.001696616,0.002046001,0.000433442],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004531674,0.0007295368,0.06016108,0.0005872764,0.000237983,0.0003246317,0.0009987069,0.01095793,0.007996833,0.003237958,0.0381837,0.8761312],"study_design_scores_gemma":[0.00009023163,0.0003132463,0.0373199,0.0003924035,0.0003441477,0.001249736,0.001099382,0.8548715,0.02523815,0.01777762,0.0611385,0.0001651387],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.2485244,0.007022094,0.6640635,0.003445747,0.001466942,0.001123033,0.01254236,0.03155316,0.03025878],"genre_scores_gemma":[0.6959285,0.001077746,0.2756044,0.0002921087,0.0004518348,0.0002178476,0.008792351,0.0003737731,0.01726159],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.006038892,"threshold_uncertainty_score":0.02020216,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2914752849","doi":"","title":"Cyberbullying Intervention Based on Convolutional Neural Networks","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Bullying, Victimization, and Aggression","field":"Psychology","cited_by":11,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Ottawa","funders":"","keywords":"Flagging; Computer science; Convolutional neural network; Intervention (counseling); Interface (matter); Process (computing); User interface; Mechanism (biology); Service (business); Human–computer interaction; Natural (archaeology); World Wide Web; Artificial intelligence; Multimedia; Psychology","authors":[{"name":"Qianjia Huang","is_ca":true},{"name":"Diana Inkpen","is_ca":true},{"name":"Jianhong Zhang","is_ca":false},{"name":"David Van Bruwaene","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.04830110351160088,"gpt":0.3533375149440351,"spread":0.3050364114324342,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0004057211,0.0006690941,0.0004087978,0.0005094378,0.0002938602,0.0005837632,0.00102819,0.0006013196,0.003160575],"category_scores_gemma":[0.001437005,0.0002684135,0.0004145719,0.00034006,0.0002734235,0.0007188174,0.0005053395,0.0008889324,0.0006595096],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001109282,"about_ca_system_score_gemma":0.0007186216,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01872916,"about_ca_topic_score_gemma":0.0213209,"domain_scores_codex":[0.9998136,0.00003266632,0.000009453841,0.00005927835,0.00004221333,0.0000428081],"domain_scores_gemma":[0.9996181,0.0002037677,0.00004250394,0.00002857209,0.00008352319,0.00002353591],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0004451446,0.0005454522,0.007246142,0.0001988017,0.0001701172,0.0003828306,0.0002738788,0.4503582,0.0263936,0.009567275,0.009507785,0.4949107],"study_design_scores_gemma":[0.000004338973,0.0000267064,0.0005638315,0.000008981612,0.00001414651,0.00001668568,0.000009667897,0.9941765,0.002828667,0.001779053,0.0005653747,0.000006035907],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1772418,0.001416037,0.7932702,0.001167798,0.0003366895,0.0002005281,0.0006000622,0.01197328,0.01379367],"genre_scores_gemma":[0.8977245,0.000432195,0.09080887,0.0002970762,0.00005577872,0.0001484411,0.0005824034,0.0001671136,0.009783629],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.01872916,"threshold_uncertainty_score":0.03724027,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2251031731","doi":"","title":"A System for Multilingual Sentiment Learning On Large Data Sets","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Sentiment Analysis and Opinion Mining","field":"Computer Science","cited_by":9,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Toronto","funders":"","keywords":"Computer science; Sentiment analysis; Artificial intelligence; Generalization; Natural language processing; Set (abstract data type); Empirical research; Machine learning; Mathematics","authors":[{"name":"Alex Cheng","is_ca":true},{"name":"Oles Zhulyn","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1209313347697141,"gpt":0.3898499392076454,"spread":0.2689186044379313,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002948735,0.001304644,0.0009033662,0.002234167,0.001375818,0.001379734,0.001065416,0.0009558889,0.008701602],"category_scores_gemma":[0.007303646,0.0005195507,0.0009276399,0.001797547,0.0003030433,0.004067365,0.002000093,0.001358712,0.007118938],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0008092294,"about_ca_system_score_gemma":0.0009870073,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003338074,"about_ca_topic_score_gemma":0.005254854,"domain_scores_codex":[0.9988078,0.0002683138,0.000157733,0.0004073776,0.0002823966,0.00007631964],"domain_scores_gemma":[0.9970751,0.000998505,0.0002062107,0.0004314449,0.001136914,0.0001518155],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0008020649,0.0004208439,0.006246318,0.0003947467,0.000229393,0.0003741218,0.0006288321,0.002872034,0.054233,0.003021212,0.06156396,0.8692135],"study_design_scores_gemma":[0.0002967622,0.0005559405,0.01019685,0.0001584349,0.0002661491,0.0007444558,0.0008635435,0.808055,0.0795351,0.01544913,0.08369618,0.0001824694],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.07903781,0.0004701489,0.7322468,0.001242514,0.0004761195,0.001371924,0.00760974,0.1705562,0.006988727],"genre_scores_gemma":[0.1569736,0.0002040867,0.8238915,0.0004345181,0.0001765182,0.0008512212,0.01056624,0.001005745,0.005896583],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.008701602,"threshold_uncertainty_score":0.02910972,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2573843450","doi":"","title":"Determining the Multiword Expression Inventory of a Surprise Language","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of New Brunswick","funders":"","keywords":"Treebank; Surprise; Computer science; Natural language processing; Identification (biology); Artificial intelligence; Language model; Natural language; Language identification; Parsing; Psychology","authors":[{"name":"Bahar Salehi","is_ca":false},{"name":"Paul Cook","is_ca":true},{"name":"Timothy Baldwin","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.03742044249948087,"gpt":0.3314787121336442,"spread":0.2940582696341633,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001075569,0.0006038567,0.0004031326,0.001601961,0.0004222373,0.001220022,0.0004661859,0.0005512749,0.003249325],"category_scores_gemma":[0.004425005,0.0005578109,0.0005832593,0.0006481379,0.0004697897,0.003107894,0.001085531,0.00132979,0.002637409],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0006053099,"about_ca_system_score_gemma":0.0008051221,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001279011,"about_ca_topic_score_gemma":0.001638681,"domain_scores_codex":[0.9993356,0.000163025,0.00006942808,0.0002813441,0.00008590304,0.00006463298],"domain_scores_gemma":[0.9978253,0.000924993,0.0002573266,0.0002459975,0.000640376,0.0001060792],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0009598696,0.0003121309,0.09591027,0.0008799175,0.0001263591,0.001935286,0.005184554,0.007465072,0.2991124,0.02092781,0.01764772,0.5495386],"study_design_scores_gemma":[0.0001181297,0.0008393952,0.201665,0.000356746,0.0003317608,0.006289718,0.008326554,0.5316431,0.1380048,0.03604274,0.07611299,0.0002690686],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8009186,0.0006407504,0.1811216,0.0004942706,0.0001568391,0.0001132827,0.002850934,0.002859955,0.01084365],"genre_scores_gemma":[0.917729,0.0003856149,0.07061923,0.0001354451,0.00006157615,0.0001235727,0.006808242,0.0005713076,0.003566073],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.003249325,"threshold_uncertainty_score":0.01087004,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2759598007","doi":"","title":"UQAM-NTL: Named entity recognition in Twitter messages.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":7,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université du Québec; Université du Québec à Montréal","funders":"","keywords":"Conditional random field; Named-entity recognition; Computer science; Task (project management); Conjunction (astronomy); Artificial intelligence; Natural language processing; Named entity; Entity linking; Information retrieval; Machine learning; Knowledge base; Engineering","authors":[{"name":"Ngoc Tan Le","is_ca":false},{"name":"Fatma Mallek","is_ca":true},{"name":"Fatiha Sadat","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0937972724176686,"gpt":0.3230069221583068,"spread":0.2292096497406382,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001753006,0.001605791,0.001046584,0.003528419,0.000901217,0.001548775,0.001795361,0.00142098,0.02592803],"category_scores_gemma":[0.008590053,0.0004967052,0.000573942,0.001839298,0.0003900431,0.004887812,0.002688827,0.0009231648,0.04295679],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009729012,"about_ca_system_score_gemma":0.001145538,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007786554,"about_ca_topic_score_gemma":0.008304489,"domain_scores_codex":[0.998669,0.0003294323,0.0001556986,0.0003653609,0.0003653076,0.0001151668],"domain_scores_gemma":[0.9978136,0.0006607587,0.0002652151,0.0006858767,0.000425902,0.0001485513],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"not_applicable","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001345118,0.0002790428,0.01029425,0.001516374,0.000162502,0.0009300897,0.000754544,0.004096952,0.03190209,0.004718136,0.6202483,0.3237525],"study_design_scores_gemma":[0.0003012942,0.000418358,0.0133215,0.0003326515,0.0001322472,0.001031351,0.0007085017,0.2813081,0.1282512,0.01159046,0.5623227,0.0002816513],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"software","genre_gemma":"empirical","genre_scores_codex":[0.02365916,0.0009198162,0.1407515,0.001257256,0.0007275176,0.0012287,0.2528756,0.5577907,0.02078977],"genre_scores_gemma":[0.1336122,0.0006527044,0.2727251,0.0008033835,0.0003130299,0.002219393,0.5464757,0.00931306,0.03388547],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.02592803,"threshold_uncertainty_score":0.08673787,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2574914175","doi":"","title":"An interactive system for exploring community question answering forums","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of British Columbia","funders":"","keywords":"Question answering; Computer science; World Wide Web; Interface (matter); Information retrieval; User interface; Graphical user interface; Human–computer interaction","authors":[{"name":"Enamul Hoque","is_ca":true},{"name":"Shafiq Joty","is_ca":false},{"name":"Lluı́s Màrquez","is_ca":false},{"name":"Alberto Barrón‐Cedeño","is_ca":false},{"name":"Giovanni Da San Martino","is_ca":false},{"name":"Alessandro Moschitti","is_ca":false},{"name":"Preslav Nakov","is_ca":false},{"name":"Salvatore Romeo","is_ca":false},{"name":"Giuseppe Carenini","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1382404640885254,"gpt":0.3595167125248636,"spread":0.2212762484363383,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.002468556,0.001314228,0.0008115504,0.00343247,0.001382634,0.00190836,0.002186694,0.001540647,0.03620894],"category_scores_gemma":[0.008234875,0.0006387189,0.0007924497,0.001928304,0.0004411732,0.004337752,0.004600656,0.0009481451,0.008040996],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005166664,"about_ca_system_score_gemma":0.0007386918,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.001617231,"about_ca_topic_score_gemma":0.002226074,"domain_scores_codex":[0.9985973,0.0004961696,0.0000888418,0.0003351272,0.0003651468,0.0001175126],"domain_scores_gemma":[0.9946569,0.003430071,0.0001695056,0.0005095545,0.0006763539,0.0005576286],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002829979,0.0008219493,0.005865743,0.001901898,0.0002307103,0.00123503,0.00804736,0.004549758,0.06855515,0.01604786,0.1918826,0.698032],"study_design_scores_gemma":[0.001345618,0.001202222,0.009638632,0.0004142562,0.0003415276,0.002082585,0.002772658,0.2240284,0.05293767,0.04358964,0.6611055,0.0005412346],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.0350836,0.0008051904,0.6966871,0.0005210279,0.0002026372,0.001128448,0.009824833,0.2406915,0.01505574],"genre_scores_gemma":[0.2068595,0.000464884,0.7338982,0.0004086492,0.0003004574,0.002882672,0.03027757,0.008683364,0.01622463],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.03620894,"threshold_uncertainty_score":0.1211309,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2579534470","doi":"","title":"Training Data Enrichment for Infrequent Discourse Relations","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":6,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of British Columbia","funders":"","keywords":"Parsing; Computer science; Training set; Relation (database); Natural language processing; Artificial intelligence; Training (meteorology); Confidence interval; Quality (philosophy); Machine learning; Data mining; Statistics","authors":[{"name":"Kailang Jiang","is_ca":false},{"name":"Giuseppe Carenini","is_ca":true},{"name":"Raymond T. Ng","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.1427233374709916,"gpt":0.4063764021825755,"spread":0.2636530647115839,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.009819563,0.001921571,0.001649217,0.002508579,0.001755386,0.001671356,0.002675527,0.003472533,0.003107717],"category_scores_gemma":[0.04363262,0.001035938,0.001568262,0.002113493,0.001382915,0.004070283,0.003579178,0.004153769,0.003210494],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009477233,"about_ca_system_score_gemma":0.002176586,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002563991,"about_ca_topic_score_gemma":0.005467535,"domain_scores_codex":[0.9920782,0.003442381,0.0006091158,0.002519438,0.0009441559,0.0004067457],"domain_scores_gemma":[0.9468291,0.03876746,0.00177678,0.007707406,0.004270123,0.0006490824],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.001506272,0.002378508,0.0494043,0.001073895,0.0002886457,0.001362294,0.002937096,0.05379548,0.06734043,0.006139383,0.02729565,0.786478],"study_design_scores_gemma":[0.0001582421,0.0005923553,0.009649676,0.0002388523,0.0002322437,0.0008444557,0.001198024,0.8404607,0.1054937,0.01479494,0.02623863,0.00009831866],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.3316331,0.001224178,0.6375189,0.001545961,0.0003439669,0.0007044203,0.002569364,0.01997371,0.004486403],"genre_scores_gemma":[0.4783269,0.0001930432,0.5049353,0.0006429923,0.00014967,0.0007217721,0.01042653,0.001070848,0.003532837],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.009819563,"threshold_uncertainty_score":0.05193144,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2889229100","doi":"","title":"NLP for Conversations: Sentiment, Summarization, and Group Dynamics","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Advanced Text Analysis Techniques","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":false,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of British Columbia; University of the Fraser Valley","funders":"","keywords":"Automatic summarization; Computer science; Natural language processing; Sentiment analysis; Artificial intelligence; Dynamics (music); Information retrieval; Group (periodic table); Psychology","authors":[{"name":"Gabriel Murray","is_ca":true},{"name":"Giuseppe Carenini","is_ca":true},{"name":"Shafiq Joty","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.02964105130895849,"gpt":0.3340051228688455,"spread":0.304364071559887,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.004400794,0.0009488459,0.0008258212,0.002536814,0.001494018,0.0025575,0.001070213,0.001078957,0.005163825],"category_scores_gemma":[0.02965687,0.0004442107,0.0008051083,0.002867084,0.0005975799,0.006166846,0.002051763,0.002027575,0.003252694],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009015207,"about_ca_system_score_gemma":0.001071412,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002286099,"about_ca_topic_score_gemma":0.002302055,"domain_scores_codex":[0.9958286,0.002307843,0.0003318266,0.0006984819,0.0006552272,0.0001780777],"domain_scores_gemma":[0.9825434,0.01296154,0.0009928436,0.001601217,0.001589223,0.0003117538],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0009171231,0.0004248792,0.007998055,0.001336918,0.0002451411,0.0003764234,0.00620324,0.01671482,0.02629807,0.03730067,0.04569811,0.8564866],"study_design_scores_gemma":[0.00008897978,0.000243619,0.009763062,0.000253839,0.000184298,0.0002915314,0.003759824,0.747309,0.01625489,0.1788331,0.04291267,0.0001052944],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.08452837,0.001638396,0.8854877,0.004536615,0.0004917681,0.0004463371,0.007724075,0.00525764,0.009889054],"genre_scores_gemma":[0.601488,0.001019247,0.3731925,0.0003271044,0.0009167192,0.0006941078,0.01658889,0.0007849846,0.004988445],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.005163825,"threshold_uncertainty_score":0.02327389,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2579773546","doi":"","title":"Reddit Temporal N-gram Corpus and its Applications on Paraphrase and Semantic Similarity in Social Media using a Topic-based Latent Semantic Analysis.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":5,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Dalhousie University","funders":"","keywords":"Paraphrase; Computer science; Latent semantic analysis; SemEval; Natural language processing; Artificial intelligence; Social media; Semantic similarity; Similarity (geometry); n-gram; Probabilistic latent semantic analysis; Text corpus; Information retrieval; Language model; World Wide Web","authors":[{"name":"Anh Duc Dang","is_ca":true},{"name":"Abidalrahman Moh’d","is_ca":true},{"name":"Aminul Islam","is_ca":true},{"name":"Rosane Minghim","is_ca":false},{"name":"Michael Smit","is_ca":true},{"name":"Evangelos Milios","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.08700156074749933,"gpt":0.3337642859727162,"spread":0.2467627252252169,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001464942,0.000792904,0.0005350296,0.004571416,0.001705727,0.001096527,0.001106884,0.0009626879,0.005859961],"category_scores_gemma":[0.01249801,0.0002878903,0.0005627694,0.004906395,0.0007282395,0.002960093,0.002502844,0.001221627,0.003157371],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007703181,"about_ca_system_score_gemma":0.001119867,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004380808,"about_ca_topic_score_gemma":0.009902752,"domain_scores_codex":[0.9981389,0.0008981363,0.0001549295,0.0003037885,0.0004250725,0.00007927178],"domain_scores_gemma":[0.9943094,0.002971546,0.0003299818,0.001066025,0.001126189,0.0001968366],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.002106715,0.001324917,0.01241132,0.003177935,0.0002507954,0.002362316,0.004167434,0.01667503,0.0612298,0.0300867,0.1402463,0.7259607],"study_design_scores_gemma":[0.0003820535,0.0007674061,0.04463049,0.0004854846,0.0001678691,0.003390463,0.004385207,0.6219219,0.06434038,0.04279044,0.2164014,0.000336938],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.441826,0.004779739,0.4074822,0.002483645,0.001251,0.001784258,0.09786966,0.01639821,0.02612524],"genre_scores_gemma":[0.5147017,0.0009896682,0.3390464,0.0002622823,0.0003332823,0.002542671,0.1337813,0.0009644694,0.007378113],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.005859961,"threshold_uncertainty_score":0.01960355,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2787587581","doi":"","title":"A Proposal for combining “general” and specialized frames","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université de Montréal","funders":"","keywords":"FrameNet; Merge (version control); Computer science; Domain (mathematical analysis); Resource (disambiguation); Semantics (computer science); Natural language processing; Representation (politics); Artificial intelligence; Information retrieval; Programming language; Parsing; Mathematics","authors":[{"name":"Marie-Claude L’Homme","is_ca":true},{"name":"Carlos Subirats","is_ca":false},{"name":"Benoît Robichaud","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.03474877466863473,"gpt":0.3396147986736514,"spread":0.3048660240050166,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006954605,0.001591776,0.00132757,0.004524835,0.003136967,0.006902109,0.004698055,0.003696007,0.01206495],"category_scores_gemma":[0.009133885,0.001372776,0.002319356,0.004154444,0.007219665,0.02608639,0.008903231,0.004055263,0.002958247],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002923522,"about_ca_system_score_gemma":0.003568393,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.004810745,"about_ca_topic_score_gemma":0.005071335,"domain_scores_codex":[0.9941583,0.002243522,0.0004880791,0.001772822,0.0007709668,0.0005664863],"domain_scores_gemma":[0.9951503,0.001239235,0.0003112397,0.002019096,0.0008269261,0.0004531631],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.00002543623,0.00001482525,0.0002305595,0.00005283504,0.00001445795,0.00006638606,0.000919647,0.0004975264,0.0007604071,0.9716713,0.001986024,0.02376064],"study_design_scores_gemma":[0.00003394626,0.00007348066,0.000363337,0.0001600484,0.00008745105,0.0004133041,0.001292133,0.01706614,0.002392986,0.8118659,0.1661884,0.00006298989],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.004866407,0.0002634782,0.9671198,0.001626284,0.0002763338,0.0001592307,0.000146613,0.0007406918,0.02480116],"genre_scores_gemma":[0.1319417,0.0003283537,0.8545499,0.000908018,0.000397722,0.0004296358,0.000552571,0.0005547649,0.01033742],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01206495,"threshold_uncertainty_score":0.04036123,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2785621828","doi":"","title":"Lexfom: a lexical functions ontology model","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":4,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université de Montréal; Université du Québec","funders":"","keywords":"Syntagmatic analysis; Computer science; Lexical item; Natural language processing; Lexical density; Lexical grammar; Lexical choice; Lexical functional grammar; Artificial intelligence; Relation (database); Function (biology); Ontology; Linguistics; Representation (politics); Perspective (graphical); Generative grammar; Phrase structure rules","authors":[{"name":"Alexsandro Fonseca","is_ca":true},{"name":"Fatiha Sadat","is_ca":true},{"name":"François Lareau","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.0599011662098063,"gpt":0.3441678799576087,"spread":0.2842667137478024,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001065565,0.001079211,0.000676723,0.003319776,0.001260037,0.003879717,0.002588302,0.001648324,0.01286951],"category_scores_gemma":[0.003029662,0.00088041,0.00236444,0.002167061,0.000955284,0.008334832,0.002246986,0.00207812,0.005042222],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.002780568,"about_ca_system_score_gemma":0.003213449,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.01473256,"about_ca_topic_score_gemma":0.0143052,"domain_scores_codex":[0.9992539,0.0001480308,0.0001172171,0.0001853626,0.0002196115,0.00007578787],"domain_scores_gemma":[0.9994142,0.0001551924,0.00004883145,0.0001754491,0.0001623586,0.00004393039],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.0002357197,0.0002604153,0.002901911,0.0007968387,0.0001489917,0.0006555555,0.001334822,0.02637389,0.006025353,0.6796489,0.05336642,0.2282511],"study_design_scores_gemma":[0.0000786392,0.00005781832,0.0009803739,0.00028051,0.0001205202,0.0005783506,0.0004234613,0.1992463,0.005758432,0.2976971,0.4946893,0.00008913242],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.008918064,0.0004546549,0.9370622,0.001449426,0.0001475585,0.0004883497,0.01438221,0.01754884,0.01954864],"genre_scores_gemma":[0.1282323,0.0009894954,0.8105661,0.0009022403,0.0001278224,0.001459053,0.03350372,0.003199821,0.02101947],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.01473256,"threshold_uncertainty_score":0.04305279,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2572452825","doi":"","title":"Selective Co-occurrences for Word-Emotion Association","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Sentiment Analysis and Opinion Mining","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"York University","funders":"","keywords":"Word (group theory); Computer science; Association (psychology); Natural language processing; Word Association; Artificial intelligence; Task (project management); Emotion classification; Semantic similarity; Psychology; Linguistics","authors":[{"name":"Ameeta Agrawal","is_ca":true},{"name":"Aijun An","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.05802061708972685,"gpt":0.3551458021427703,"spread":0.2971251850530434,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001401728,0.001132024,0.001085694,0.003874016,0.0007703872,0.00106595,0.0009654629,0.0008716413,0.002212543],"category_scores_gemma":[0.008338585,0.0003928622,0.001062668,0.003845321,0.0007743163,0.003064917,0.002146328,0.001511872,0.002306403],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0003682439,"about_ca_system_score_gemma":0.0009595037,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.00108969,"about_ca_topic_score_gemma":0.003549029,"domain_scores_codex":[0.9982865,0.0003995385,0.000163141,0.0007592672,0.000263502,0.0001279404],"domain_scores_gemma":[0.9954513,0.002263395,0.0005541928,0.00085825,0.0006903378,0.0001825945],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0009174516,0.0005448681,0.04328621,0.0005027491,0.0003420523,0.0002640867,0.001113242,0.0133431,0.03843475,0.005021889,0.008052931,0.8881767],"study_design_scores_gemma":[0.00008639182,0.0003984022,0.04320947,0.00009351523,0.0002173042,0.001056781,0.0009764948,0.8788162,0.03489234,0.02841609,0.01171436,0.0001227022],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.207225,0.001121644,0.7810143,0.0003061933,0.0001804044,0.0003273053,0.001646893,0.004310661,0.003867508],"genre_scores_gemma":[0.6869684,0.0005025734,0.3014522,0.0001298968,0.0002038446,0.0006173705,0.006273242,0.0005231504,0.003329313],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.003874016,"threshold_uncertainty_score":0.007413089,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2575224319","doi":"","title":"Capturing Pragmatic Knowledge in Article Usage Prediction using LSTMs.","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":3,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"McGill University","funders":"","keywords":"Interpretability; Computer science; Coreference; Artificial intelligence; Task (project management); Machine learning; Mechanism (biology); Recurrent neural network; Natural language processing; Long short term memory; Artificial neural network; Resolution (logic)","authors":[{"name":"Jad Kabbara","is_ca":true},{"name":"Yulan Feng","is_ca":false},{"name":"Jackie Chi Kit Cheung","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.04258849303752668,"gpt":0.3358901740860344,"spread":0.2933016810485077,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001327146,0.0008343553,0.0003919166,0.001470074,0.000326522,0.001250987,0.001020552,0.001073216,0.002183674],"category_scores_gemma":[0.01060308,0.0004315542,0.0004821419,0.001232107,0.0003816539,0.003463726,0.0007403807,0.001425695,0.001114286],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0007990486,"about_ca_system_score_gemma":0.0007785956,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005100396,"about_ca_topic_score_gemma":0.01278679,"domain_scores_codex":[0.9992555,0.000290092,0.00005659274,0.0002021961,0.0001318257,0.00006377717],"domain_scores_gemma":[0.9964406,0.002401772,0.0003995794,0.0002183794,0.0004609405,0.00007880598],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0007192743,0.0004430556,0.02543494,0.00139596,0.0003934377,0.0009449815,0.001369068,0.1380398,0.06548867,0.01106168,0.01344881,0.7412604],"study_design_scores_gemma":[0.00001454509,0.0000590001,0.004229909,0.00005581879,0.00005474245,0.0001203052,0.0001700342,0.971835,0.008592524,0.01172956,0.003112197,0.00002638152],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.378333,0.002924536,0.5892298,0.002421591,0.0004576166,0.0002385613,0.003862157,0.007980743,0.01455197],"genre_scores_gemma":[0.905984,0.0004858828,0.08819275,0.0001656738,0.0001201796,0.00009403758,0.002589275,0.0001677663,0.002200393],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.005100396,"threshold_uncertainty_score":0.01014143,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2577720462","doi":"","title":"Named Entity Disambiguation for little known referents: a topic-based approach","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":false,"ca_fund":false,"ca_venue":false,"about_ca":true},"ca_institutions":"","funders":"","keywords":"Computer science; Referent; Exploit; Task (project management); Property (philosophy); Named-entity recognition; Natural language processing; Entity linking; Artificial intelligence; Information retrieval; Training set; Named entity; Linguistics; Knowledge base","authors":[{"name":"Andrea Glaser","is_ca":false},{"name":"Jonas Kuhn","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.09194997485170264,"gpt":0.3341074863322806,"spread":0.242157511480578,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.006220178,0.00165675,0.002479359,0.009488828,0.002800216,0.003777918,0.004692754,0.002998338,0.003081372],"category_scores_gemma":[0.01442929,0.001229875,0.002707486,0.008121777,0.001288116,0.009456827,0.004922508,0.003011609,0.004467298],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0009128103,"about_ca_system_score_gemma":0.002376503,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003218338,"about_ca_topic_score_gemma":0.006890372,"domain_scores_codex":[0.9941949,0.001780954,0.0005309291,0.002076158,0.00110273,0.0003142836],"domain_scores_gemma":[0.9896684,0.005609176,0.0006876347,0.002018172,0.001711107,0.0003054929],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0006223609,0.0006668643,0.01123066,0.001115354,0.000655308,0.001382808,0.004280854,0.03827206,0.0302094,0.0416601,0.03399593,0.8359083],"study_design_scores_gemma":[0.0001147975,0.0001970994,0.00552349,0.0002906614,0.0006398302,0.001979173,0.00206486,0.7761338,0.03235357,0.09624717,0.0842006,0.0002550751],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.01131224,0.00132962,0.9804185,0.0007875542,0.000220141,0.0001286769,0.0007200566,0.00320654,0.001876716],"genre_scores_gemma":[0.249283,0.001673352,0.7327743,0.0006920127,0.001032024,0.0004247628,0.006177031,0.00104573,0.006897741],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.009488828,"threshold_uncertainty_score":0.0328958,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2888880755","doi":"","title":"Do Character-Level Neural Network Language Models Capture Knowledge of Multiword Expression Compositionality?","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of New Brunswick","funders":"","keywords":"Principle of compositionality; Computer science; Artificial intelligence; Natural language processing; Character (mathematics); Treebank; Expression (computer science); Artificial neural network; Programming language; Annotation","authors":[{"name":"Ali Hakimi Parizi","is_ca":false},{"name":"Paul Cook","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.06360377069358782,"gpt":0.3540654245744966,"spread":0.2904616538809088,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001071115,0.0007585135,0.0005765493,0.0005699777,0.0002523402,0.001521038,0.001095448,0.0008327164,0.00232336],"category_scores_gemma":[0.007353724,0.0004533743,0.0005262528,0.0005245949,0.0005139881,0.005612127,0.0006485967,0.00176541,0.00125806],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0005696893,"about_ca_system_score_gemma":0.0005861915,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002810754,"about_ca_topic_score_gemma":0.003994032,"domain_scores_codex":[0.9995651,0.0001496784,0.00002216632,0.000159759,0.00005730129,0.00004601953],"domain_scores_gemma":[0.9976704,0.001334024,0.0002895457,0.0002489868,0.0003803847,0.00007668759],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.0007655436,0.0004158066,0.02892595,0.0007555518,0.0006904267,0.0004415681,0.001020628,0.3404503,0.06436678,0.0629559,0.004677464,0.4945341],"study_design_scores_gemma":[0.000009184239,0.00004854711,0.002398079,0.00002954291,0.00004591726,0.0000543551,0.00007391087,0.9642832,0.003429945,0.02830921,0.001295674,0.0000223931],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.1854207,0.0008977634,0.8023847,0.001625228,0.0001801351,0.00009253507,0.0007118744,0.001180752,0.007506264],"genre_scores_gemma":[0.9132646,0.0006978319,0.08027994,0.0003793022,0.00008473181,0.0001671122,0.0009781008,0.0001535343,0.003994937],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.002810754,"threshold_uncertainty_score":0.007772386,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2573346038","doi":"","title":"Predicting sentential semantic compatibility for aggregation in text-to-text generation","year":2016,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"McGill University","funders":"","keywords":"Computer science; Sentence; Natural language processing; Artificial intelligence; Task (project management); Cluster analysis; Process (computing); Context (archaeology); Information retrieval","authors":[{"name":"Victor Chenal","is_ca":false},{"name":"Jackie Chi Kit Cheung","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.07688502906661919,"gpt":0.3287613563504564,"spread":0.2518763272838372,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.005651951,0.001122228,0.0007154646,0.003934626,0.001130775,0.001980493,0.0009407126,0.001396344,0.001993907],"category_scores_gemma":[0.0345922,0.0003714665,0.0009722483,0.002372116,0.0005118169,0.003942818,0.001404012,0.00136352,0.001274905],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.000913418,"about_ca_system_score_gemma":0.001288601,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003557357,"about_ca_topic_score_gemma":0.006092007,"domain_scores_codex":[0.9963157,0.001677354,0.0003383271,0.0008557483,0.0005903229,0.0002225627],"domain_scores_gemma":[0.9759448,0.01785277,0.001366986,0.001604183,0.002717842,0.0005133476],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.002287034,0.001042489,0.07798526,0.001145489,0.0003942559,0.000852232,0.002968399,0.06360371,0.05728237,0.01007382,0.02782758,0.7545374],"study_design_scores_gemma":[0.0001022418,0.0003683032,0.03686602,0.00006849034,0.0001934946,0.0003721606,0.0008095256,0.8931515,0.04463444,0.01722133,0.006120567,0.00009199715],"study_design_candidate":"simulation_or_modeling","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.6377878,0.001482883,0.3388717,0.0009890356,0.0002897285,0.0007476666,0.003938065,0.008942451,0.006950719],"genre_scores_gemma":[0.8106865,0.0001811969,0.1818631,0.00007538965,0.00009557742,0.0002179534,0.005738619,0.0003389973,0.0008026836],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.005651951,"threshold_uncertainty_score":0.02989072,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2849292054","doi":"","title":"Automatically Extracting Qualia Relations for the Rich Event Ontology","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Semantic Web and Ontologies","field":"Computer Science","cited_by":2,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Guelph","funders":"","keywords":"Qualia; Computer science; Ontology; Commonsense knowledge; Event (particle physics); Artificial intelligence; Focus (optics); Semantics (computer science); Structuring; Ontology learning; Natural language processing; Upper ontology; Suggested Upper Merged Ontology; Knowledge extraction; Semantic Web; Epistemology; Consciousness","authors":[{"name":"Ghazaleh Kazeminejad","is_ca":false},{"name":"Claire Bonial","is_ca":false},{"name":"Susan Windisch Brown","is_ca":true},{"name":"Martha Palmer","is_ca":false}],"retraction":null,"screen_n_in":null,"score":{"opus":0.08858363997313826,"gpt":0.3894872429872838,"spread":0.3009036030141455,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.001649339,0.0006776633,0.000509032,0.007603145,0.00155672,0.002598071,0.0009499292,0.0008025821,0.004540921],"category_scores_gemma":[0.007809014,0.0005198207,0.00149353,0.004382984,0.0008058121,0.006611177,0.002690648,0.001793984,0.001841816],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001793244,"about_ca_system_score_gemma":0.00281896,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.007987111,"about_ca_topic_score_gemma":0.01828645,"domain_scores_codex":[0.9987209,0.0002069171,0.000204064,0.0003287794,0.0004340443,0.000105235],"domain_scores_gemma":[0.9973648,0.001132324,0.0003667576,0.0004791164,0.0005501529,0.0001068043],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"bench_or_experimental","study_design_scores_codex":[0.0002642147,0.0004241494,0.01759065,0.001415265,0.0001627865,0.002451292,0.005259877,0.01124864,0.04634036,0.429867,0.04557849,0.4393973],"study_design_scores_gemma":[0.00007317478,0.00008068225,0.01565362,0.000563683,0.0002125907,0.001423128,0.005476701,0.2593204,0.0322213,0.3961098,0.2886807,0.0001841838],"study_design_candidate":"bench_or_experimental","study_design_consensus":null,"genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.07593526,0.0004658147,0.8794717,0.001622943,0.0002442275,0.0005537806,0.01537709,0.006445043,0.01988431],"genre_scores_gemma":[0.2793029,0.0006309386,0.6833987,0.0002688315,0.00007581562,0.0003455603,0.03255774,0.0008458591,0.002573627],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.007987111,"threshold_uncertainty_score":0.01588124,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2847160827","doi":"","title":"Reproducing and Regularizing the SCRN Model","year":2018,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Topic Modeling","field":"Computer Science","cited_by":1,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Toronto Metropolitan University","funders":"","keywords":"Dropout (neural networks); Tying; Computer science; Language model; Task (project management); Artificial intelligence; Data modeling; Machine learning; Algorithm; Engineering","authors":[{"name":"Olzhas Kabdolov","is_ca":false},{"name":"Zhenisbek Assylbekov","is_ca":true},{"name":"Rustem Takhanov","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.07632542245516838,"gpt":0.3238737291177107,"spread":0.2475483066625423,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.0008163284,0.0006515087,0.0004878051,0.0003170478,0.0002457387,0.0004594951,0.001437133,0.0007918947,0.002818501],"category_scores_gemma":[0.003523275,0.0003307712,0.0006977192,0.0003334154,0.0005455735,0.001220188,0.001031192,0.001438092,0.001487711],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.0004698087,"about_ca_system_score_gemma":0.001002901,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.005752943,"about_ca_topic_score_gemma":0.008228241,"domain_scores_codex":[0.9995832,0.00008757893,0.00001743123,0.0001613872,0.0000979853,0.00005250721],"domain_scores_gemma":[0.9993351,0.000197392,0.00006124152,0.0002353555,0.0001356305,0.00003521752],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"simulation_or_modeling","study_design_gemma":"simulation_or_modeling","study_design_scores_codex":[0.000111658,0.00009218569,0.001448537,0.000101955,0.0000713679,0.0002880895,0.000160342,0.8294194,0.0340237,0.04467933,0.004983946,0.08461945],"study_design_scores_gemma":[0.00000370911,0.00001421602,0.00009772358,0.000002505546,0.00000576455,0.00002792633,0.000004114226,0.9925249,0.001720052,0.005042295,0.0005513004,0.000005472298],"study_design_candidate":"simulation_or_modeling","study_design_consensus":"simulation_or_modeling","genre_codex":"methods","genre_gemma":"empirical","genre_scores_codex":[0.04570713,0.0001111929,0.9482759,0.0003198643,0.0001293001,0.00004446251,0.0003201588,0.001803589,0.003288454],"genre_scores_gemma":[0.7469057,0.0001987987,0.2396373,0.0003022396,0.0001008603,0.000190713,0.001417259,0.0007113001,0.01053566],"genre_candidate":"empirical","genre_consensus":null,"teacher_disagreement_score":0.005752943,"threshold_uncertainty_score":0.01143891,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2252045241","doi":"","title":"Flexible Structural Analysis of Near-Meet-Semilattices for Typed Unification-Based Grammar Design","year":2012,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"University of Toronto","funders":"","keywords":"Unification; Parsing; Computer science; Programming language; Head-driven phrase structure grammar; Grammar; Rule-based machine translation; Type (biology); Theoretical computer science; Algorithm; Natural language processing; Artificial intelligence; Generative grammar; Linguistics","authors":[{"name":"Rouzbeh Farahmand","is_ca":true},{"name":"Gerald Penn","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.07928709696875265,"gpt":0.3722236565658698,"spread":0.2929365595971172,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.003564658,0.0006699596,0.000636664,0.00119832,0.00131444,0.002395982,0.001861315,0.0008448368,0.00502908],"category_scores_gemma":[0.006492186,0.001015898,0.001899544,0.0009352793,0.002487119,0.004014394,0.002550199,0.002240656,0.001364181],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.001554354,"about_ca_system_score_gemma":0.002388401,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.002010535,"about_ca_topic_score_gemma":0.003660196,"domain_scores_codex":[0.9970809,0.0008593098,0.0003103642,0.0003754041,0.001175141,0.0001988753],"domain_scores_gemma":[0.9971706,0.001312862,0.0001967516,0.0006031904,0.0005515571,0.0001650349],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"theoretical_or_conceptual","study_design_gemma":"theoretical_or_conceptual","study_design_scores_codex":[0.00009931222,0.00008744186,0.000889023,0.0002017205,0.00003777099,0.0004024509,0.001029485,0.06096716,0.01775195,0.8418993,0.001821237,0.07481316],"study_design_scores_gemma":[0.00002387254,0.00004004886,0.0001024162,0.00004658033,0.00002316756,0.0001198424,0.0001593255,0.3330038,0.01120943,0.6459767,0.009261044,0.00003388837],"study_design_candidate":"theoretical_or_conceptual","study_design_consensus":"theoretical_or_conceptual","genre_codex":"methods","genre_gemma":"methods","genre_scores_codex":[0.005357228,0.00003409848,0.9925493,0.00008194221,0.00002196752,0.00004446135,0.00004784645,0.0005551083,0.001308134],"genre_scores_gemma":[0.12889,0.00007448968,0.8677093,0.0000639966,0.00003238261,0.0001886231,0.000261161,0.0004884747,0.002291573],"genre_candidate":"methods","genre_consensus":"methods","teacher_disagreement_score":0.00502908,"threshold_uncertainty_score":0.01885194,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null},{"id":"W2251004534","doi":"","title":"Fourteen Light Tasks for comparing Analogical and Phrase-based Machine Translation","year":2014,"lang":"en","type":"article","venue":"International Conference on Computational Linguistics","topic":"Natural Language Processing Techniques","field":"Computer Science","cited_by":0,"is_retracted":false,"has_abstract":true,"routes":{"ca_aff":true,"ca_fund":false,"ca_venue":false,"about_ca":false},"ca_institutions":"Université de Montréal","funders":"","keywords":"Transliteration; Machine translation; Phrase; Computer science; Scripting language; Natural language processing; Artificial intelligence; Rule-based machine translation; Translation (biology); Example-based machine translation; Machine translation software usability; Programming language","authors":[{"name":"Rafik Rhouma","is_ca":false},{"name":"Phillippe Langlais","is_ca":true}],"retraction":null,"screen_n_in":null,"score":{"opus":0.05860160290824588,"gpt":0.332169161512734,"spread":0.2735675586044881,"validation_status":"score_only:v0-immature-baseline"},"prediction":{"model_version":"metacan-v3-hybrid-931329e0061c","candidate_categories":[],"consensus_categories":[],"category_scores_codex":[0.00574626,0.002137405,0.00154397,0.00407831,0.001755496,0.002021737,0.002026021,0.002567154,0.006774161],"category_scores_gemma":[0.0207995,0.0004760612,0.001618525,0.003039628,0.001087523,0.002966921,0.003273259,0.002044247,0.003872226],"about_ca_system_candidate":false,"about_ca_system_consensus":false,"about_ca_system_score_codex":0.00126384,"about_ca_system_score_gemma":0.001475768,"about_ca_topic_candidate":false,"about_ca_topic_consensus":false,"about_ca_topic_score_codex":0.003027789,"about_ca_topic_score_gemma":0.005815649,"domain_scores_codex":[0.9933809,0.002359647,0.001229694,0.001306768,0.001380184,0.0003427815],"domain_scores_gemma":[0.9853948,0.009156359,0.0007127495,0.002378872,0.001659258,0.0006979499],"domain_codex":null,"domain_gemma":null,"domain_candidate":null,"domain_consensus":null,"study_design_codex":"design_other","study_design_gemma":"not_applicable","study_design_scores_codex":[0.0145471,0.005805582,0.03156887,0.008883675,0.002668289,0.0007248577,0.0009148674,0.03290587,0.04182458,0.005537467,0.0746709,0.779948],"study_design_scores_gemma":[0.008642872,0.02139004,0.258911,0.001749273,0.004104425,0.004068796,0.003842986,0.2903233,0.1704368,0.03642344,0.1989383,0.001169001],"study_design_candidate":"not_applicable","study_design_consensus":null,"genre_codex":"empirical","genre_gemma":"empirical","genre_scores_codex":[0.8451771,0.0157751,0.07087881,0.001204854,0.002664494,0.003355478,0.02341769,0.008212156,0.02931442],"genre_scores_gemma":[0.7610103,0.002359075,0.1320142,0.001277772,0.0006137888,0.00367432,0.08832413,0.001346146,0.009380348],"genre_candidate":"empirical","genre_consensus":"empirical","teacher_disagreement_score":0.006774161,"threshold_uncertainty_score":0.03038949,"prediction_status":"machine_predicted_unvalidated"},"labels":[],"label_agreement":null}]}